{"as_of":"2026-08-13T04:00:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:57d5a4d6b4a4242532471c9604fa4758809556653737f05194594338331321d6","coverage":[{"denominator":31,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":31,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-12T11:40:09.167187Z","state":"measured"},{"denominator":32,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":32,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-12T06:34:41.77262+00:00","state":"measured"},{"denominator":1,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":1,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-12T11:40:09.024777Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"pith","source_observed_at":"2026-08-12T11:40:09.282733Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2411.17995","last_updated":"2024-11-27T02:24:51Z","snapshot_observed_at":"2026-08-12T11:35:14.236202Z","submitted_at":"2024-11-27T02:24:51Z","title":"Revisiting Misalignment in Multispectral Pedestrian Detection: A Language-Driven Approach for Cross-modal Alignment Fusion","version":1},"cited_work":{"arxiv_id":"2411.17995","doi":null,"metadata_source":"pith","pith_arxiv_id":"2411.17995","snapshot_observed_at":"2026-08-12T11:40:09.282733Z","title":"Revisiting Misalignment in Multispectral Pedestrian Detection: A Language-Driven Approach for Cross-modal Alignment Fusion","venue":"cs.CV","work_id":"904593b0-45c2-4324-80c1-2410a9c1c944","year":2024},"citing_paper":{"arxiv_id":"2411.17995","last_updated":"2024-11-27T02:24:51Z","snapshot_observed_at":"2026-08-12T11:35:14.236202Z","submitted_at":"2024-11-27T02:24:51Z","title":"Revisiting Misalignment in Multispectral Pedestrian Detection: A Language-Driven Approach for Cross-modal Alignment Fusion","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-12T11:40:09.024777Z"},"links":{"cited_paper":"/paper/2411.17995","citing_paper":"/paper/2411.17995"},"observation_digest":"sha256:64e3e7606a0e51a30eeee2741f84bba0d5fec83c82afc402e9cb20a6ac2a1e20","observation_id":"b1648490-e76e-4a4c-ab03-af243d25564e","resolution":{"observed_at":"2026-08-12T11:40:09.287941Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2411.17995/citation-record","integrity":"/paper/2411.17995/integrity","json":"/paper/2411.17995/citation-record.json","paper":"/paper/2411.17995"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2411.17995","last_updated":"2024-11-27T02:24:51Z","snapshot_observed_at":"2026-08-12T11:35:14.236202Z","submitted_at":"2024-11-27T02:24:51Z","title":"Revisiting Misalignment in Multispectral Pedestrian Detection: A Language-Driven Approach for Cross-modal Alignment Fusion","version":1},"cited_work":{"arxiv_id":"2411.17995","doi":null,"metadata_source":"pith","pith_arxiv_id":"2411.17995","snapshot_observed_at":"2026-08-12T11:40:09.282733Z","title":"Revisiting Misalignment in Multispectral Pedestrian Detection: A Language-Driven Approach for Cross-modal Alignment Fusion","venue":"cs.CV","work_id":"904593b0-45c2-4324-80c1-2410a9c1c944","year":2024},"citing_paper":{"arxiv_id":"2411.17995","last_updated":"2024-11-27T02:24:51Z","snapshot_observed_at":"2026-08-12T11:35:14.236202Z","submitted_at":"2024-11-27T02:24:51Z","title":"Revisiting Misalignment in Multispectral Pedestrian Detection: A Language-Driven Approach for Cross-modal Alignment Fusion","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-12T11:40:09.024777Z"},"links":{"cited_paper":"/paper/2411.17995","citing_paper":"/paper/2411.17995"},"observation_digest":"sha256:64e3e7606a0e51a30eeee2741f84bba0d5fec83c82afc402e9cb20a6ac2a1e20","observation_id":"b1648490-e76e-4a4c-ab03-af243d25564e","resolution":{"observed_at":"2026-08-12T11:40:09.287941Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:40:09.679803Z","title":"To solve the misalignment problems, camera calibration tech- niques and image registration algorithms have been devel- oped","venue":null,"work_id":"7bb8db29-d650-4ca7-9d9d-cd0e80cf7153","year":null},"citing_paper":{"arxiv_id":"2411.17995","last_updated":"2024-11-27T02:24:51Z","snapshot_observed_at":"2026-08-12T11:35:14.236202Z","submitted_at":"2024-11-27T02:24:51Z","title":"Revisiting Misalignment in Multispectral Pedestrian Detection: A Language-Driven Approach for Cross-modal Alignment Fusion","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-12T11:40:09.030580Z"},"links":{"citing_paper":"/paper/2411.17995"},"observation_digest":"sha256:f924e93cf4301c8587d5933ddb87eef67a1db0de7f943f34d187490f7b800809","observation_id":"fa3c65d1-4ded-491e-80a1-d88a21516245","resolution":{"observed_at":"2026-08-12T11:40:09.684750Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:40:09.664638Z","title":"front\"} Clothes: {clothes: [","venue":null,"work_id":"85da9c66-9cfe-4b5b-9baf-faf82d4df8a6","year":null},"citing_paper":{"arxiv_id":"2411.17995","last_updated":"2024-11-27T02:24:51Z","snapshot_observed_at":"2026-08-12T11:35:14.236202Z","submitted_at":"2024-11-27T02:24:51Z","title":"Revisiting Misalignment in Multispectral Pedestrian Detection: A Language-Driven Approach for Cross-modal Alignment Fusion","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-12T11:40:09.035532Z"},"links":{"citing_paper":"/paper/2411.17995"},"observation_digest":"sha256:52eee2312c3bf48e9cb5e0d43e789b9eead296c4413004cdc255fa4cf726dd27","observation_id":"59f5cf56-1a57-4aeb-9ca6-ae1a892b5eb3","resolution":{"observed_at":"2026-08-12T11:40:09.669540Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:40:09.649325Z","title":"Implementation Details We conduct experiments on two different heavily misaligned RGB-thermal multispectral dataset","venue":null,"work_id":"b0f76751-128d-4180-9657-1a928b25b481","year":null},"citing_paper":{"arxiv_id":"2411.17995","last_updated":"2024-11-27T02:24:51Z","snapshot_observed_at":"2026-08-12T11:35:14.236202Z","submitted_at":"2024-11-27T02:24:51Z","title":"Revisiting Misalignment in Multispectral Pedestrian Detection: A Language-Driven Approach for Cross-modal Alignment Fusion","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-12T11:40:09.040813Z"},"links":{"citing_paper":"/paper/2411.17995"},"observation_digest":"sha256:e8955b72657d7fcd38f6da5707b1e0a8c3482c07268de55bbbe63c1c1be1f509","observation_id":"4c9a6126-c464-4bc8-94e2-4d70f4a1ee3b","resolution":{"observed_at":"2026-08-12T11:40:09.654241Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:40:09.634475Z","title":null,"venue":null,"work_id":"5adb090b-203b-439c-86d8-42fbb7e79e54","year":null},"citing_paper":{"arxiv_id":"2411.17995","last_updated":"2024-11-27T02:24:51Z","snapshot_observed_at":"2026-08-12T11:35:14.236202Z","submitted_at":"2024-11-27T02:24:51Z","title":"Revisiting Misalignment in Multispectral Pedestrian Detection: A Language-Driven Approach for Cross-modal Alignment Fusion","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-12T11:40:09.045475Z"},"links":{"citing_paper":"/paper/2411.17995"},"observation_digest":"sha256:ddb9b6a3d1ab4e4d24af891a864258bbce04e9e78ef4707180ec2fc40c721742","observation_id":"4c083780-3d37-4d88-966d-b977c9fb4a20","resolution":{"observed_at":"2026-08-12T11:40:09.639500Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:40:09.619554Z","title":"Multispectral pedestrian de- tection: Benchmark dataset and baseline,","venue":null,"work_id":"c9785419-7d9d-4696-976e-1b59da0743df","year":2015},"citing_paper":{"arxiv_id":"2411.17995","last_updated":"2024-11-27T02:24:51Z","snapshot_observed_at":"2026-08-12T11:35:14.236202Z","submitted_at":"2024-11-27T02:24:51Z","title":"Revisiting Misalignment in Multispectral Pedestrian Detection: A Language-Driven Approach for Cross-modal Alignment Fusion","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-12T11:40:09.050069Z"},"links":{"citing_paper":"/paper/2411.17995"},"observation_digest":"sha256:fe9231f7c3929163b039bfc77766e11917709fc8ddf0638add16b1938708416f","observation_id":"68200406-309a-4f42-be09-26671a4baf4c","resolution":{"observed_at":"2026-08-12T11:40:09.624230Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:40:09.604213Z","title":"Uncertainty-guided cross-modal learning for robust multispectral pedestrian detection,","venue":null,"work_id":"90f6c5f7-1b62-4afe-a984-1f869a22bfd8","year":2021},"citing_paper":{"arxiv_id":"2411.17995","last_updated":"2024-11-27T02:24:51Z","snapshot_observed_at":"2026-08-12T11:35:14.236202Z","submitted_at":"2024-11-27T02:24:51Z","title":"Revisiting Misalignment in Multispectral Pedestrian Detection: A Language-Driven Approach for Cross-modal Alignment Fusion","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-12T11:40:09.054868Z"},"links":{"citing_paper":"/paper/2411.17995"},"observation_digest":"sha256:7bed2d8d797b165b9868fcd38a1bab15609d7ea514fc9711e748f1f0edd14a8e","observation_id":"00a3f00f-2bb6-4f26-ba4e-7e02e9e98267","resolution":{"observed_at":"2026-08-12T11:40:09.609259Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2005.10987","last_updated":"2020-05-22T03:45:06Z","snapshot_observed_at":"2026-08-08T11:18:20.826611Z","submitted_at":"2020-05-22T03:45:06Z","title":"Investigating Vulnerability to Adversarial Examples on Multimodal Data Fusion in Deep Learning","version":1},"cited_work":{"arxiv_id":"2005.10987","doi":null,"metadata_source":"pith","pith_arxiv_id":"2005.10987","snapshot_observed_at":"2026-08-12T11:40:09.261088Z","title":"Investigating Vulnerability to Adversarial Examples on Multimodal Data Fusion in Deep Learning","venue":"cs.CV","work_id":"dae2df82-5d67-44dc-aaf0-36cc2152acf4","year":2020},"citing_paper":{"arxiv_id":"2411.17995","last_updated":"2024-11-27T02:24:51Z","snapshot_observed_at":"2026-08-12T11:35:14.236202Z","submitted_at":"2024-11-27T02:24:51Z","title":"Revisiting Misalignment in Multispectral Pedestrian Detection: A Language-Driven Approach for Cross-modal Alignment Fusion","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-12T11:40:09.059584Z"},"links":{"cited_paper":"/paper/2005.10987","citing_paper":"/paper/2411.17995"},"observation_digest":"sha256:5700baa0d0b18db1bae753e64f2438f9a81b780ecec7ff142f59e4208d073a0d","observation_id":"f6928708-69ea-401d-a16b-7f21e1de3b77","resolution":{"observed_at":"2026-08-12T11:40:09.266539Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:40:09.588688Z","title":"Mul- tispectral invisible coating: laminated visible-thermal physical attack against multispectral object detectors us- ing transparent low-e films,","venue":null,"work_id":"55476ed7-10dd-427f-ae81-9b45ef056103","year":2023},"citing_paper":{"arxiv_id":"2411.17995","last_updated":"2024-11-27T02:24:51Z","snapshot_observed_at":"2026-08-12T11:35:14.236202Z","submitted_at":"2024-11-27T02:24:51Z","title":"Revisiting Misalignment in Multispectral Pedestrian Detection: A Language-Driven Approach for Cross-modal Alignment Fusion","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-12T11:40:09.064580Z"},"links":{"citing_paper":"/paper/2411.17995"},"observation_digest":"sha256:68a921df2cd222af3ef5c05123ae50abf5a950be30cc1abd4751f9bf948b4624","observation_id":"0d01b9a8-509d-49c1-a7a9-adb17615dd1b","resolution":{"observed_at":"2026-08-12T11:40:09.593836Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:40:09.573865Z","title":"Towards robust train- ing of multi-sensor data fusion network against adver- sarial examples in semantic segmentation,","venue":null,"work_id":"d45d434c-ccca-4af3-ba1c-5aba7b7eb1e8","year":2021},"citing_paper":{"arxiv_id":"2411.17995","last_updated":"2024-11-27T02:24:51Z","snapshot_observed_at":"2026-08-12T11:35:14.236202Z","submitted_at":"2024-11-27T02:24:51Z","title":"Revisiting Misalignment in Multispectral Pedestrian Detection: A Language-Driven Approach for Cross-modal Alignment Fusion","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-12T11:40:09.069411Z"},"links":{"citing_paper":"/paper/2411.17995"},"observation_digest":"sha256:90729fffd25fb1125c7a95010260f336655ee1a7c1be10306a1e749a1193c555","observation_id":"326249a2-6ec2-46ed-873a-840f7f2da7a0","resolution":{"observed_at":"2026-08-12T11:40:09.578915Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:40:09.558329Z","title":"Robust multispectral pedestrian de- tection via spectral position-free feature mapping,","venue":null,"work_id":"ecc6e516-f3f7-46c7-b062-876cd2f317cb","year":2023},"citing_paper":{"arxiv_id":"2411.17995","last_updated":"2024-11-27T02:24:51Z","snapshot_observed_at":"2026-08-12T11:35:14.236202Z","submitted_at":"2024-11-27T02:24:51Z","title":"Revisiting Misalignment in Multispectral Pedestrian Detection: A Language-Driven Approach for Cross-modal Alignment Fusion","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-12T11:40:09.074699Z"},"links":{"citing_paper":"/paper/2411.17995"},"observation_digest":"sha256:340486f281d0afd8d8dbebca7843eb996b9e3639c4486bec4ff6133350d120c0","observation_id":"475e3ab5-ef0e-47e8-810d-77f62afc9f1a","resolution":{"observed_at":"2026-08-12T11:40:09.563157Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:40:09.543281Z","title":"Robust multispectral pedestrian detection via uncertainty-aware cross-modal learning,","venue":null,"work_id":"ec3ad1d8-c2b0-47f8-a10f-62562b0e75de","year":2021},"citing_paper":{"arxiv_id":"2411.17995","last_updated":"2024-11-27T02:24:51Z","snapshot_observed_at":"2026-08-12T11:35:14.236202Z","submitted_at":"2024-11-27T02:24:51Z","title":"Revisiting Misalignment in Multispectral Pedestrian Detection: A Language-Driven Approach for Cross-modal Alignment Fusion","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-12T11:40:09.079672Z"},"links":{"citing_paper":"/paper/2411.17995"},"observation_digest":"sha256:e45da1c4099894b570f1bdab8d0e9733ad1972c70bbcbd4837c7c3f07e42f4e8","observation_id":"a9cf74e8-d0d8-4a1c-8903-166943515bdf","resolution":{"observed_at":"2026-08-12T11:40:09.548261Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:40:09.528173Z","title":"Detrs with collaborative hybrid assignments training,","venue":null,"work_id":"dc8d874d-19b4-472f-b45e-f1119a84ef50","year":2023},"citing_paper":{"arxiv_id":"2411.17995","last_updated":"2024-11-27T02:24:51Z","snapshot_observed_at":"2026-08-12T11:35:14.236202Z","submitted_at":"2024-11-27T02:24:51Z","title":"Revisiting Misalignment in Multispectral Pedestrian Detection: A Language-Driven Approach for Cross-modal Alignment Fusion","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-12T11:40:09.084314Z"},"links":{"citing_paper":"/paper/2411.17995"},"observation_digest":"sha256:4c35bc6e2e6741b79b1f455a2e2b8c0f1173a20a594f5bb236711259b0cec0be","observation_id":"8d41737a-1537-404b-b176-cd85e21aee9c","resolution":{"observed_at":"2026-08-12T11:40:09.533268Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:40:09.513782Z","title":"Defending person detection against adversarial patch attack by using universal defensive frame,","venue":null,"work_id":"5b19fd3f-1c05-4d52-8551-196abf3dc767","year":2022},"citing_paper":{"arxiv_id":"2411.17995","last_updated":"2024-11-27T02:24:51Z","snapshot_observed_at":"2026-08-12T11:35:14.236202Z","submitted_at":"2024-11-27T02:24:51Z","title":"Revisiting Misalignment in Multispectral Pedestrian Detection: A Language-Driven Approach for Cross-modal Alignment Fusion","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-12T11:40:09.088816Z"},"links":{"citing_paper":"/paper/2411.17995"},"observation_digest":"sha256:9fdf6f88f50e8ec1514d3e0660611db1bf1dbdd519a74d379b64d90489bc3394","observation_id":"ab21b13b-4fc3-4668-b328-484566cc810a","resolution":{"observed_at":"2026-08-12T11:40:09.518542Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:40:09.499063Z","title":"De- fending physical adversarial attack on object detection via adversarial patch-feature energy,","venue":null,"work_id":"cf745432-9218-45da-8218-5d235f11693c","year":2022},"citing_paper":{"arxiv_id":"2411.17995","last_updated":"2024-11-27T02:24:51Z","snapshot_observed_at":"2026-08-12T11:35:14.236202Z","submitted_at":"2024-11-27T02:24:51Z","title":"Revisiting Misalignment in Multispectral Pedestrian Detection: A Language-Driven Approach for Cross-modal Alignment Fusion","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-12T11:40:09.093456Z"},"links":{"citing_paper":"/paper/2411.17995"},"observation_digest":"sha256:a0acae312dce877e838c31ea41dd121bae74bc5bf6517d7e2ebc3bc9eed44655","observation_id":"f1cbb4a3-a076-4b1b-9a81-a45a38e6abe4","resolution":{"observed_at":"2026-08-12T11:40:09.504007Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:40:09.484081Z","title":"Inte- grating language-derived appearance elements with vi- sual cues in pedestrian detection,","venue":null,"work_id":"e69ed233-973c-44aa-8719-87a62632769d","year":2024},"citing_paper":{"arxiv_id":"2411.17995","last_updated":"2024-11-27T02:24:51Z","snapshot_observed_at":"2026-08-12T11:35:14.236202Z","submitted_at":"2024-11-27T02:24:51Z","title":"Revisiting Misalignment in Multispectral Pedestrian Detection: A Language-Driven Approach for Cross-modal Alignment Fusion","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-12T11:40:09.097808Z"},"links":{"citing_paper":"/paper/2411.17995"},"observation_digest":"sha256:499df1700b9a202c6441eb82e557f2533bf71746f69632cdec2a24f3435a6df2","observation_id":"dab158b2-8aa2-4292-b270-4460ec768701","resolution":{"observed_at":"2026-08-12T11:40:09.488918Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:40:09.469219Z","title":"Robust pedestrian detection via constructing versatile pedestrian knowledge bank,","venue":null,"work_id":"5d9d0edb-f9ea-4907-af2e-f83724e6cfc8","year":2024},"citing_paper":{"arxiv_id":"2411.17995","last_updated":"2024-11-27T02:24:51Z","snapshot_observed_at":"2026-08-12T11:35:14.236202Z","submitted_at":"2024-11-27T02:24:51Z","title":"Revisiting Misalignment in Multispectral Pedestrian Detection: A Language-Driven Approach for Cross-modal Alignment Fusion","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-12T11:40:09.102220Z"},"links":{"citing_paper":"/paper/2411.17995"},"observation_digest":"sha256:6448f540ff818e8f395fc0340e779ccd55c9ba1fc9306b3ae19d1c6f1e70c0bd","observation_id":"6e85df86-458e-4db8-99d8-18cab345f31c","resolution":{"observed_at":"2026-08-12T11:40:09.473993Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:40:09.454395Z","title":"Map: Multispectral adversarial patch to attack person detec- tion,","venue":null,"work_id":"0022f510-59b8-4eed-a43f-73105cee3ef5","year":2022},"citing_paper":{"arxiv_id":"2411.17995","last_updated":"2024-11-27T02:24:51Z","snapshot_observed_at":"2026-08-12T11:35:14.236202Z","submitted_at":"2024-11-27T02:24:51Z","title":"Revisiting Misalignment in Multispectral Pedestrian Detection: A Language-Driven Approach for Cross-modal Alignment Fusion","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-12T11:40:09.106814Z"},"links":{"citing_paper":"/paper/2411.17995"},"observation_digest":"sha256:bbdeda15afa2a708516ba739d7553b5218f43e8d6dd7260c16685f5bbd0fc10d","observation_id":"8a0a4258-ff65-4a8f-a7b4-e764eef941e5","resolution":{"observed_at":"2026-08-12T11:40:09.459237Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:40:09.439267Z","title":"Causal mode multiplexer: A novel framework for unbiased multispectral pedestrian detec- tion,","venue":null,"work_id":"7f088e5b-cb75-4153-9414-3edde8f62615","year":2024},"citing_paper":{"arxiv_id":"2411.17995","last_updated":"2024-11-27T02:24:51Z","snapshot_observed_at":"2026-08-12T11:35:14.236202Z","submitted_at":"2024-11-27T02:24:51Z","title":"Revisiting Misalignment in Multispectral Pedestrian Detection: A Language-Driven Approach for Cross-modal Alignment Fusion","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-12T11:40:09.111360Z"},"links":{"citing_paper":"/paper/2411.17995"},"observation_digest":"sha256:e9b3dbb8ae3e9892457264aa73c6009b67a4db304859be77de29bbdf575511e7","observation_id":"81700dbc-6f3e-4e2e-82bb-5ad26fb5801d","resolution":{"observed_at":"2026-08-12T11:40:09.444256Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2403.15209","last_updated":"2025-01-08T09:29:10Z","snapshot_observed_at":"2026-08-13T00:47:52.966815Z","submitted_at":"2024-03-22T13:50:27Z","title":"MSCoTDet: Language-driven Multi-modal Fusion for Improved Multispectral Pedestrian Detection","version":3},"cited_work":{"arxiv_id":"2403.15209","doi":null,"metadata_source":"pith","pith_arxiv_id":"2403.15209","snapshot_observed_at":"2026-08-12T11:40:09.236323Z","title":"MSCoTDet: Language-driven Multi-modal Fusion for Improved Multispectral Pedestrian Detection","venue":"cs.CV","work_id":"d2366e8f-689b-41f0-9308-aa7ce1e48fda","year":2024},"citing_paper":{"arxiv_id":"2411.17995","last_updated":"2024-11-27T02:24:51Z","snapshot_observed_at":"2026-08-12T11:35:14.236202Z","submitted_at":"2024-11-27T02:24:51Z","title":"Revisiting Misalignment in Multispectral Pedestrian Detection: A Language-Driven Approach for Cross-modal Alignment Fusion","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-12T11:40:09.115908Z"},"links":{"cited_paper":"/paper/2403.15209","citing_paper":"/paper/2411.17995"},"observation_digest":"sha256:1a082bddfd7d2c4398f54d14c530a02202fc1cea033d821564cc5759091b3331","observation_id":"000fc0e8-f46a-42b6-a421-43db9dce4b82","resolution":{"observed_at":"2026-08-12T11:40:09.243651Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:40:09.424325Z","title":"A survey of image registration techniques,","venue":null,"work_id":"02478fc0-9d5c-4e10-9912-45a8bc3ec949","year":1992},"citing_paper":{"arxiv_id":"2411.17995","last_updated":"2024-11-27T02:24:51Z","snapshot_observed_at":"2026-08-12T11:35:14.236202Z","submitted_at":"2024-11-27T02:24:51Z","title":"Revisiting Misalignment in Multispectral Pedestrian Detection: A Language-Driven Approach for Cross-modal Alignment Fusion","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-12T11:40:09.120641Z"},"links":{"citing_paper":"/paper/2411.17995"},"observation_digest":"sha256:3b717b00bddf471975d99e331bbd4807d8b8782a91cd65a65789d8f6e9529ec7","observation_id":"e4f114e3-4996-4aca-834e-b3a4b17a634c","resolution":{"observed_at":"2026-08-12T11:40:09.429092Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:40:09.409370Z","title":"Re- mote sensing image registration techniques: A survey,","venue":null,"work_id":"eb6d17c5-1216-45a9-8591-8351288ee5f1","year":2010},"citing_paper":{"arxiv_id":"2411.17995","last_updated":"2024-11-27T02:24:51Z","snapshot_observed_at":"2026-08-12T11:35:14.236202Z","submitted_at":"2024-11-27T02:24:51Z","title":"Revisiting Misalignment in Multispectral Pedestrian Detection: A Language-Driven Approach for Cross-modal Alignment Fusion","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-12T11:40:09.125723Z"},"links":{"citing_paper":"/paper/2411.17995"},"observation_digest":"sha256:96e2336a50652649b54a6171348ce0fb5f38a7652dc3a86e84719d0dd0f404c9","observation_id":"b0eef678-d6b5-494d-bc99-b30f42ae9a66","resolution":{"observed_at":"2026-08-12T11:40:09.414436Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:40:09.394282Z","title":"A survey of medical image registration,","venue":null,"work_id":"5847071a-6096-4b67-97fa-6fe1c15e2b18","year":1998},"citing_paper":{"arxiv_id":"2411.17995","last_updated":"2024-11-27T02:24:51Z","snapshot_observed_at":"2026-08-12T11:35:14.236202Z","submitted_at":"2024-11-27T02:24:51Z","title":"Revisiting Misalignment in Multispectral Pedestrian Detection: A Language-Driven Approach for Cross-modal Alignment Fusion","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-12T11:40:09.130240Z"},"links":{"citing_paper":"/paper/2411.17995"},"observation_digest":"sha256:f254ba1f71fe3813ec0eebc9f12bf7aa7824b0d32fa7018e139ee0dbd80ed253","observation_id":"bfabbda6-3aed-4f1b-954c-229341a7dfc9","resolution":{"observed_at":"2026-08-12T11:40:09.399259Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:40:09.378899Z","title":"An iterative integrated frame- work for thermal–visible image registration, sensor fu- sion, and people tracking for video surveillance applica- tions,","venue":null,"work_id":"5da70119-5ab9-4a7f-9a1a-ec5c34feb238","year":2012},"citing_paper":{"arxiv_id":"2411.17995","last_updated":"2024-11-27T02:24:51Z","snapshot_observed_at":"2026-08-12T11:35:14.236202Z","submitted_at":"2024-11-27T02:24:51Z","title":"Revisiting Misalignment in Multispectral Pedestrian Detection: A Language-Driven Approach for Cross-modal Alignment Fusion","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-12T11:40:09.134695Z"},"links":{"citing_paper":"/paper/2411.17995"},"observation_digest":"sha256:c3565917ebae88ff4a39cf57c24d4e4251985a81fc0a442d754427bc0b6b929c","observation_id":"2262fa51-aab1-4406-b127-7c90d879af8b","resolution":{"observed_at":"2026-08-12T11:40:09.384258Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:40:09.362770Z","title":"Attentive alignment network for multispectral pedestrian detection,","venue":null,"work_id":"f4f50c2e-3715-4049-b8c6-1752ede353d7","year":2023},"citing_paper":{"arxiv_id":"2411.17995","last_updated":"2024-11-27T02:24:51Z","snapshot_observed_at":"2026-08-12T11:35:14.236202Z","submitted_at":"2024-11-27T02:24:51Z","title":"Revisiting Misalignment in Multispectral Pedestrian Detection: A Language-Driven Approach for Cross-modal Alignment Fusion","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-12T11:40:09.139322Z"},"links":{"citing_paper":"/paper/2411.17995"},"observation_digest":"sha256:a6d418f113b392c6d675a43bc479b2df7f72c9ad0c607da2f219a0f8d6a44218","observation_id":"3231b0d4-ec05-4fa3-ab07-87b6bdadeaba","resolution":{"observed_at":"2026-08-12T11:40:09.367574Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:40:09.347297Z","title":"Weakly aligned cross-modal learning for multispectral pedestrian detection,","venue":null,"work_id":"7a1a62e5-ff81-43c2-825c-afd41be431ce","year":2019},"citing_paper":{"arxiv_id":"2411.17995","last_updated":"2024-11-27T02:24:51Z","snapshot_observed_at":"2026-08-12T11:35:14.236202Z","submitted_at":"2024-11-27T02:24:51Z","title":"Revisiting Misalignment in Multispectral Pedestrian Detection: A Language-Driven Approach for Cross-modal Alignment Fusion","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-12T11:40:09.143722Z"},"links":{"citing_paper":"/paper/2411.17995"},"observation_digest":"sha256:854dab550806f839b6f9bdba4cc253e61605a13cb480973f6c335f061f52d366","observation_id":"fca2de4d-cb9a-44ef-8804-2b93e8dab9d2","resolution":{"observed_at":"2026-08-12T11:40:09.352262Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:40:09.330136Z","title":"Illumination-aware faster r-cnn for robust multispectral pedestrian detection,","venue":null,"work_id":"000634a8-620b-4647-b5b6-91004607f381","year":2019},"citing_paper":{"arxiv_id":"2411.17995","last_updated":"2024-11-27T02:24:51Z","snapshot_observed_at":"2026-08-12T11:35:14.236202Z","submitted_at":"2024-11-27T02:24:51Z","title":"Revisiting Misalignment in Multispectral Pedestrian Detection: A Language-Driven Approach for Cross-modal Alignment Fusion","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-12T11:40:09.148398Z"},"links":{"citing_paper":"/paper/2411.17995"},"observation_digest":"sha256:2db64d5ad73ca12fbbf8da870398726955ae1b1861de07eb0291bab4e2ac7853","observation_id":"e7bd2dc4-df73-4897-8d66-e6f186963bf0","resolution":{"observed_at":"2026-08-12T11:40:09.336564Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:40:09.314207Z","title":"Improving multispectral pedestrian detection by addressing modal- ity imbalance problems,","venue":null,"work_id":"9d037baa-fbc2-468e-9a77-8247780b1321","year":2020},"citing_paper":{"arxiv_id":"2411.17995","last_updated":"2024-11-27T02:24:51Z","snapshot_observed_at":"2026-08-12T11:35:14.236202Z","submitted_at":"2024-11-27T02:24:51Z","title":"Revisiting Misalignment in Multispectral Pedestrian Detection: A Language-Driven Approach for Cross-modal Alignment Fusion","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-12T11:40:09.153149Z"},"links":{"citing_paper":"/paper/2411.17995"},"observation_digest":"sha256:775aa125da3e716f06fd7d85805a487b6dd20c6034cd9310e5d9ee7368d26324","observation_id":"80b6d849-9154-4c12-a9c1-a8e948a02ae5","resolution":{"observed_at":"2026-08-12T11:40:09.319275Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T11:40:09.298795Z","title":"Multimodal object de- tection via probabilistic ensembling,","venue":null,"work_id":"ce6cad3e-cf17-4791-a03b-34612ec7e3e7","year":2022},"citing_paper":{"arxiv_id":"2411.17995","last_updated":"2024-11-27T02:24:51Z","snapshot_observed_at":"2026-08-12T11:35:14.236202Z","submitted_at":"2024-11-27T02:24:51Z","title":"Revisiting Misalignment in Multispectral Pedestrian Detection: A Language-Driven Approach for Cross-modal Alignment Fusion","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-12T11:40:09.157779Z"},"links":{"citing_paper":"/paper/2411.17995"},"observation_digest":"sha256:4667c5b60fe88b46f39329ee445ada008063b738174d124d5bc8f4e45798c29e","observation_id":"4beac263-610c-4121-995a-3ec78e0c9262","resolution":{"observed_at":"2026-08-12T11:40:09.303946Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2303.08774","last_updated":"2024-03-04T06:01:33Z","snapshot_observed_at":"2026-08-07T07:30:12.213965Z","submitted_at":"2023-03-15T17:15:04Z","title":"GPT-4 Technical Report","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2303.08774","snapshot_observed_at":"2026-08-12T11:40:09.162305Z","title":"Gpt-4 technical report,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2411.17995","last_updated":"2024-11-27T02:24:51Z","snapshot_observed_at":"2026-08-12T11:35:14.236202Z","submitted_at":"2024-11-27T02:24:51Z","title":"Revisiting Misalignment in Multispectral Pedestrian Detection: A Language-Driven Approach for Cross-modal Alignment Fusion","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-12T11:40:09.162305Z"},"links":{"cited_paper":"/paper/2303.08774","citing_paper":"/paper/2411.17995"},"observation_digest":"sha256:6153dcb2448fc1a46bb1d5b2045c5d72cee5ff5858a91bd8f3e501ddfcad20c1","observation_id":"84e6e9b0-689f-4e1f-a1f0-e08bcb400f9e","resolution":{"observed_at":"2026-08-12T11:40:09.162305Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2403.05530","last_updated":"2024-12-16T17:39:39Z","snapshot_observed_at":"2026-07-06T17:41:42.995949Z","submitted_at":"2024-03-08T18:54:20Z","title":"Gemini 1.5: Unlocking multimodal understanding across millions of tokens of context","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.05530","snapshot_observed_at":"2026-08-12T11:40:09.167187Z","title":"Gemini 1.5: Unlocking multimodal understanding across millions of tokens of context,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2411.17995","last_updated":"2024-11-27T02:24:51Z","snapshot_observed_at":"2026-08-12T11:35:14.236202Z","submitted_at":"2024-11-27T02:24:51Z","title":"Revisiting Misalignment in Multispectral Pedestrian Detection: A Language-Driven Approach for Cross-modal Alignment Fusion","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-12T11:40:09.167187Z"},"links":{"cited_paper":"/paper/2403.05530","citing_paper":"/paper/2411.17995"},"observation_digest":"sha256:caba25cf263ec5674351a4cdbcbac145629d1c9600b5c782c9655205b1858c31","observation_id":"2d557494-09a6-4060-868d-38249b4b54bf","resolution":{"observed_at":"2026-08-12T11:40:09.167187Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2411.17995","last_updated":"2024-11-27T02:24:51Z","latest_version":1,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-12T11:35:14.236202Z","submitted_at":"2024-11-27T02:24:51Z","title":"Revisiting Misalignment in Multispectral Pedestrian Detection: A Language-Driven Approach for Cross-modal Alignment Fusion"},"reference_resolution":{"displayed":31,"state_counts":{"malformed_identifier":0,"metadata_mismatch":1,"parse_uncertain":0,"unresolved":3,"verified_exact":2,"verified_fuzzy":25},"total_outbound_references":31},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"thesis":"As of 13 August 2026, this Paper Citation Record lists 31 of 31 outbound references and 1 inbound Pith citation observation for arXiv:2411.17995."}