{"as_of":"2026-08-19T19:49:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:57b6716fc62264f122d2a7ca5d897fa7ef4f05e5f49b54bee667d5aaca49308b","coverage":[{"denominator":80,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":80,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-16T04:10:45.280927Z","state":"measured"},{"denominator":80,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":80,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-19T06:32:44.657259+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2505.01950/citation-record","integrity":"/paper/2505.01950/integrity","json":"/paper/2505.01950/citation-record.json","paper":"/paper/2505.01950"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2302.08890","last_updated":"2024-04-11T15:34:46Z","snapshot_observed_at":"2026-08-17T09:05:47.162911Z","submitted_at":"2023-02-17T14:19:28Z","title":"Deep Learning for Event-based Vision: A Comprehensive Survey and Benchmarks","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2302.08890","snapshot_observed_at":"2026-08-16T04:10:44.668699Z","title":"Deep learning for event-based vision: A comprehensive survey and benchmarks,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:44.668699Z"},"links":{"cited_paper":"/paper/2302.08890","citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:b8f1b28994201a3c9243a9e61ec4a279f525f5a199eec04ca3d93ffc8e91f18e","observation_id":"cd5ae109-5273-4f94-9920-b915a8ecb5ad","resolution":{"observed_at":"2026-08-16T04:10:44.668699Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T04:10:47.504107Z","title":"Eventdance: Unsupervised source-free cross- modal adaptation for event-based object recognition,","venue":null,"work_id":"bc616488-0091-4516-a644-5f08dd256790","year":2024},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:44.675647Z"},"links":{"citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:f85123e544bd2880e53eccfd4e414b3f2d1fed44876ce5777288885c799a9b17","observation_id":"d1b17d69-3687-4ce0-bdd0-7b5ba1eaddcc","resolution":{"observed_at":"2026-08-16T04:10:47.509637Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T04:10:47.477256Z","title":"Exact: Language-guided conceptual reasoning and uncertainty estimation for event-based action recognition and more,","venue":null,"work_id":"3a312c5f-6200-47be-a25e-5e7b99cc9520","year":2024},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:44.681899Z"},"links":{"citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:9a612d9d7428c77af5ec5f9f1e1b45adbc6d9c158a1d351a8e5b401567a5fd46","observation_id":"33b14ac3-05e4-42c3-bf5c-3dbdbc9ebe37","resolution":{"observed_at":"2026-08-16T04:10:47.483643Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T04:10:47.456311Z","title":"Dehazed image quality evaluation: From partial discrepancy to blind perception,","venue":null,"work_id":"319e655b-3a89-4624-924f-afe1299dd7a7","year":2024},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:44.690389Z"},"links":{"citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:0d28c04c6e00397f2d39a28cc88ac6e5d91a1ed56b7b13ce5cf00676e8c75821","observation_id":"cf23a4e5-78ef-434c-8a9e-344e89b486ee","resolution":{"observed_at":"2026-08-16T04:10:47.462115Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T04:10:47.435850Z","title":"Context-aware interaction network for rgb-t semantic segmentation,","venue":null,"work_id":"254a25ea-6c7e-4266-a5d9-ec0324863f62","year":2024},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:44.698278Z"},"links":{"citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:7a9ca288e253afd57157951991f681b9c39f0571558a7aea405b024802959a57","observation_id":"76744527-0e48-4c21-828d-c6ee65446c14","resolution":{"observed_at":"2026-08-16T04:10:47.442330Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T04:10:47.409724Z","title":"Mfnet: Towards real-time semantic segmentation for autonomous vehicles with multi-spectral scenes,","venue":null,"work_id":"0f48f5a2-e979-4792-a6a1-19d9bcbce1b9","year":2017},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:44.704045Z"},"links":{"citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:39acf7556d2f205fb624919899fb1a88afd25b5c1e310bcee94dc730bc1e48b2","observation_id":"156cf127-41e5-4e0f-8363-ef433538eea5","resolution":{"observed_at":"2026-08-16T04:10:47.417218Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T04:10:47.388407Z","title":"Multi-interactive feature learning and a full-time multi-modality benchmark for image fusion and segmentation,","venue":null,"work_id":"98d76b30-1f28-4498-9aad-1c976e7e0f8a","year":2023},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:44.710628Z"},"links":{"citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:5e7a7abbe4960723f759574fea3dc7ada51a68268b77b351319d4ef8b42807ab","observation_id":"999af57c-36cf-42ac-a753-2252ca41ab3e","resolution":{"observed_at":"2026-08-16T04:10:47.396003Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T04:10:47.363909Z","title":"Pst900: Rgb-thermal calibration, dataset and segmentation network,","venue":null,"work_id":"13312285-19f8-4aa8-99d2-b6446b4c0214","year":2020},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:44.715987Z"},"links":{"citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:fca38c54074e1bccac11d87fa7d309603dd34deeb4670a19cbeaa52c63f1d4ab","observation_id":"9a1db663-1238-428a-b12f-e3364fe8a756","resolution":{"observed_at":"2026-08-16T04:10:47.371853Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T04:10:47.335239Z","title":"Mffenet: Multiscale feature fusion and enhancement network for rgb–thermal urban road scene parsing,","venue":null,"work_id":"767c77b4-73da-4267-bad2-8b891825d949","year":2022},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:44.721585Z"},"links":{"citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:2a14e87c5a33a7622a47d0300d5ce2dee0ba28ca0b18fdd401df8d5aa69c67a8","observation_id":"b65d3908-c29c-4cec-a7e8-90201d94368f","resolution":{"observed_at":"2026-08-16T04:10:47.343057Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T04:10:47.311001Z","title":"Rtfnet: Rgb-thermal fusion network for semantic segmentation of urban scenes,","venue":null,"work_id":"1935323d-b068-4523-9656-035345c55cae","year":2019},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:44.729022Z"},"links":{"citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:df6865ec289094cc6a5edf9a336134c96dfb706dcecff290bb500bd4842a4bb0","observation_id":"ae09e60f-dda6-4687-9d55-46d11de98ad5","resolution":{"observed_at":"2026-08-16T04:10:47.317634Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T04:10:47.287075Z","title":"Fuseseg: Semantic segmentation of urban scenes based on rgb and thermal data fusion,","venue":null,"work_id":"0f577438-c80b-4eb8-85af-c3e77f42d693","year":2020},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:44.735290Z"},"links":{"citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:e99df5b5180f8a276adb0fc9716c7f8053353003e9b5cc3633e20adf4c297ae8","observation_id":"dd9bf0fd-592a-4fa8-9541-2b876e490cf9","resolution":{"observed_at":"2026-08-16T04:10:47.294010Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T04:10:47.262779Z","title":"Context-aware interaction network for rgb-t semantic segmentation,","venue":null,"work_id":"037540a2-8d02-4290-ab7a-48e4b238b01f","year":2024},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:44.752963Z"},"links":{"citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:9c41c33032a461da967be5236195067a11174d159a27674abd4bd993cac924fe","observation_id":"aec84846-dc2e-4699-aee9-ac966fc150b9","resolution":{"observed_at":"2026-08-16T04:10:47.271496Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T04:10:47.242317Z","title":"Segment anything,","venue":null,"work_id":"f1a22948-40d5-43c7-983f-672567f568d1","year":2023},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:44.759494Z"},"links":{"citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:e0c2efe283a20d288c1a8448f92aa6704f4c1a9341454ee8c435e91934f31532","observation_id":"2ec2ae7d-edeb-4ec4-ba4b-5ea9fb8a7e75","resolution":{"observed_at":"2026-08-16T04:10:47.248235Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2408.00714","last_updated":"2024-10-28T16:37:57Z","snapshot_observed_at":"2026-07-06T18:55:41.459417Z","submitted_at":"2024-08-01T17:00:08Z","title":"SAM 2: Segment Anything in Images and Videos","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2408.00714","snapshot_observed_at":"2026-08-16T04:10:44.765801Z","title":"Sam 2: Segment anything in images and videos,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:44.765801Z"},"links":{"cited_paper":"/paper/2408.00714","citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:73347b3d965f47effbfabb24cac707930a6d38289b48fcf70a2aeca36baf0932","observation_id":"fdc3dc25-05bc-4ed0-9b91-e7e21ed3727e","resolution":{"observed_at":"2026-08-16T04:10:44.765801Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T04:10:47.223038Z","title":"Msgfusion: Medical semantic guided two-branch network for multi- modal brain image fusion,","venue":null,"work_id":"2cf5b098-86ab-4b49-a28a-067ae4557eba","year":2024},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:44.775091Z"},"links":{"citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:40c3ac1a557ae3838e04732aaf686b96fd62ce7116651c8ae65c733c1fed9c81","observation_id":"ceda88f4-5854-4e3c-9cdd-0d50b9b7dc87","resolution":{"observed_at":"2026-08-16T04:10:47.229772Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T04:10:47.203282Z","title":"Cddfuse: Correlation-driven dual-branch feature decomposition for multi-modality image fusion,","venue":null,"work_id":"b4307f0a-c5ea-45fb-82e4-afe0d11802a4","year":2023},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:44.781528Z"},"links":{"citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:26d5f2d573c9d5e1afd925c11289e9e909bb969bf5d7b7d8f0a3f3a6f26aab0f","observation_id":"dfcba1df-3d51-421e-b3fd-232e6b2b547a","resolution":{"observed_at":"2026-08-16T04:10:47.209921Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T04:10:47.182162Z","title":"Multi-focus image fusion based on multi-scale gradients and image matting,","venue":null,"work_id":"9cf66a33-c6b3-41de-a285-c7c441263089","year":2022},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:44.791693Z"},"links":{"citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:49e32729f827698064b43965bc49efd18466dd83f67590ef1cd5beb81d46aa2c","observation_id":"e16b82f5-a863-487a-a2b8-080ef9d12bed","resolution":{"observed_at":"2026-08-16T04:10:47.188394Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T04:10:47.155276Z","title":"Ifsepr: A general framework for image fusion based on separate representation learning,","venue":null,"work_id":"f45329c0-bd70-4b1e-98af-1aa3a5e0ceee","year":2023},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:44.799023Z"},"links":{"citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:dc96a9f2e9b61912c33bbcc6004e2cd3c94df134dd2df7ca031f149353e00765","observation_id":"73288662-e5a4-468d-b598-c07d9853d489","resolution":{"observed_at":"2026-08-16T04:10:47.162568Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2309.03905","last_updated":"2023-09-11T20:25:16Z","snapshot_observed_at":"2026-08-16T15:02:27.041378Z","submitted_at":"2023-09-07T17:59:45Z","title":"ImageBind-LLM: Multi-modality Instruction Tuning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2309.03905","snapshot_observed_at":"2026-08-16T04:10:44.804697Z","title":"Imagebind-llm: Multi-modality instruction tuning,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:44.804697Z"},"links":{"cited_paper":"/paper/2309.03905","citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:4cfa1c0de8d2eff85fa2c2a600e596ac9595510517243d30a310e38f35e5c5da","observation_id":"731f1e66-3426-4d4a-be69-1883635ce662","resolution":{"observed_at":"2026-08-16T04:10:44.804697Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2304.15010","last_updated":"2023-04-28T17:59:25Z","snapshot_observed_at":"2026-08-12T23:40:42.885633Z","submitted_at":"2023-04-28T17:59:25Z","title":"LLaMA-Adapter V2: Parameter-Efficient Visual Instruction Model","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2304.15010","snapshot_observed_at":"2026-08-16T04:10:44.813317Z","title":"Llama-adapter v2: Parameter-efficient visual instruction model,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:44.813317Z"},"links":{"cited_paper":"/paper/2304.15010","citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:03f86b075b52348803af3da4ad89195871e5fc36960647c46ccba9c3a32817ec","observation_id":"8dbecf08-2a7e-4193-b4f9-b4bee52a0846","resolution":{"observed_at":"2026-08-16T04:10:44.813317Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.06500","last_updated":"2023-06-15T08:00:18Z","snapshot_observed_at":"2026-08-13T18:58:34.541884Z","submitted_at":"2023-05-11T00:38:10Z","title":"InstructBLIP: Towards General-purpose Vision-Language Models with Instruction Tuning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.06500","snapshot_observed_at":"2026-08-16T04:10:44.825725Z","title":"Instructblip: Towards general-purpose vision-language models with instruction tuning. arxiv 2023,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:44.825725Z"},"links":{"cited_paper":"/paper/2305.06500","citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:0bdcc8430610ec111780577262e80d6c006b2b4202eb56e243dbe51a762d7b7b","observation_id":"5fd60997-05ac-43e8-9e9b-04f771452e23","resolution":{"observed_at":"2026-08-16T04:10:44.825725Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T04:10:47.131230Z","title":"Learning transferable visual models from natural language supervision,","venue":null,"work_id":"b6f61b95-8be5-4764-8ef0-1e4475800fe5","year":2021},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:44.835374Z"},"links":{"citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:96445df17fb171f31405c1b01e09b3cafc2986110e4826edb796838067b6b713","observation_id":"3a416865-f1f1-4b47-a41c-30e578072793","resolution":{"observed_at":"2026-08-16T04:10:47.140046Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T04:10:47.108187Z","title":"Unleash the power of vision-language models by visual attention prompt and multi-modal interaction,","venue":null,"work_id":"ee20e2c2-4e98-4f9b-a0c0-e26b4c80a5cd","year":2024},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:44.841669Z"},"links":{"citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:d1400fe4cad14b4d83abecd16bb1b5fa7c2029f70370d5e24120eca5f07cd402","observation_id":"f1a6b2d7-82d3-459f-9ae5-976bbdb3158c","resolution":{"observed_at":"2026-08-16T04:10:47.115912Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T04:10:47.079811Z","title":"Multi-task paired masking with alignment modeling for medical vision- language pre-training,","venue":null,"work_id":"55c7ac4c-c93a-47e1-8e61-40d68d233e16","year":2024},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:44.847280Z"},"links":{"citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:ad24390da64bd88dd15d59bd0d67a34e2b008d1d8e1538fdaf49f418cfd7783b","observation_id":"fce145e4-1969-4b4a-86f8-9a841ff988eb","resolution":{"observed_at":"2026-08-16T04:10:47.088193Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T04:10:47.049048Z","title":"Joint bilateral upsampling,","venue":null,"work_id":"40ccc062-7aa4-406f-8e3e-ec25a0f61c22","year":2007},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:44.853547Z"},"links":{"citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:647db2a5ce80b791fdf948e7c5f2bd6034d75abb43b2aa28b25209087554a80c","observation_id":"9d4a712d-6ea7-47c8-bf1e-1b9d4a6f5f69","resolution":{"observed_at":"2026-08-16T04:10:47.061786Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T04:10:47.029963Z","title":"Distilling efficient vision transformers from cnns for semantic segmentation,","venue":null,"work_id":"dec021c8-c22f-4fad-af56-293f58622458","year":2025},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:44.860252Z"},"links":{"citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:2d632b3118ab2fb9590c30236c3e6f64971c2aa6439c3842e780858d68f052da","observation_id":"81f8f65e-8faa-4a8c-a2c4-56d67f6656eb","resolution":{"observed_at":"2026-08-16T04:10:47.035513Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T04:10:47.007342Z","title":"Eventbind: Learning a unified representation to bind them all for event-based open-world understanding,","venue":null,"work_id":"6e631f4c-ab9d-4f1a-b46f-a492dc888c78","year":2024},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:44.867302Z"},"links":{"citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:182f24ce1a051f19cc80e39aabd7ebc405a857aabc8740e95e7c46c21e2fb3a1","observation_id":"bf4c2ab3-a5ec-44b0-a36d-4115067449d6","resolution":{"observed_at":"2026-08-16T04:10:47.016147Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2412.16876","last_updated":"2024-12-22T06:12:03Z","snapshot_observed_at":"2026-08-15T05:47:29.055177Z","submitted_at":"2024-12-22T06:12:03Z","title":"MAGIC++: Efficient and Resilient Modality-Agnostic Semantic Segmentation via Hierarchical Modality Selection","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.16876","snapshot_observed_at":"2026-08-16T04:10:44.873281Z","title":"Magic++: Efficient and resilient modality-agnostic semantic segmentation via hierarchical modality selection,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:44.873281Z"},"links":{"cited_paper":"/paper/2412.16876","citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:f4840a5cff299bd029a9af1e437d92b6b18b9375328075de2dbcc25ef1744000","observation_id":"350c20ab-9b83-4ef9-bc2e-612d9528f7a4","resolution":{"observed_at":"2026-08-16T04:10:44.873281Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2411.17141","last_updated":"2025-05-15T18:54:21Z","snapshot_observed_at":"2026-08-12T12:25:58.124112Z","submitted_at":"2024-11-26T06:15:27Z","title":"Learning Robust Anymodal Segmentor with Unimodal and Cross-modal Distillation","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.17141","snapshot_observed_at":"2026-08-16T04:10:44.883691Z","title":"Learning robust anymodal segmentor with unimodal and cross-modal distillation,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:44.883691Z"},"links":{"cited_paper":"/paper/2411.17141","citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:d89db67b68b4cd2e8f7f4fa536321c353d9ee0accadc4b15e9cdd34462f58cc1","observation_id":"a4ef6b52-f6e4-4ca7-a035-926c2f6bb63f","resolution":{"observed_at":"2026-08-16T04:10:44.883691Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T04:10:46.983168Z","title":"Mrfs: Mutually rein- forcing image fusion and segmentation,","venue":null,"work_id":"c9be97b3-c83c-4293-b547-ad5f62d449d4","year":2024},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:44.890541Z"},"links":{"citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:8066a00b9a7ada79f063406a499ef49bd3f9e6d3871b7b225f0a64b718d7b0d7","observation_id":"74f15663-4a85-419c-b5b1-75ebe13771c5","resolution":{"observed_at":"2026-08-16T04:10:46.989654Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2412.04220","last_updated":"2024-12-05T14:54:31Z","snapshot_observed_at":"2026-08-17T23:29:24.004954Z","submitted_at":"2024-12-05T14:54:31Z","title":"Customize Segment Anything Model for Multi-Modal Semantic Segmentation with Mixture of LoRA Experts","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.04220","snapshot_observed_at":"2026-08-16T04:10:44.896585Z","title":"Customize segment anything model for multi-modal semantic segmentation with mixture of lora experts,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:44.896585Z"},"links":{"cited_paper":"/paper/2412.04220","citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:8e61718173de70fb59c622d6fd019b7d0ef4a70cae7d50c14555397b42e5b5b6","observation_id":"21449838-4645-4b43-a7b2-fdb26820cc7f","resolution":{"observed_at":"2026-08-16T04:10:44.896585Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2503.02581","last_updated":"2025-07-22T13:20:58Z","snapshot_observed_at":"2026-08-18T08:58:31.049671Z","submitted_at":"2025-03-04T13:04:46Z","title":"Unveiling the Potential of Segment Anything Model 2 for RGB-Thermal Semantic Segmentation with Language Guidance","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.02581","snapshot_observed_at":"2026-08-16T04:10:44.902798Z","title":"Unveiling the potential of segment anything model 2 for rgb- thermal semantic segmentation with language guidance,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:44.902798Z"},"links":{"cited_paper":"/paper/2503.02581","citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:ad3ead03cdf4e7df035c0033877a03b70984f7f0a5922ee7eb2bf5739b1adbd6","observation_id":"4b815fa3-55c0-4008-a977-d1d92c47aac7","resolution":{"observed_at":"2026-08-16T04:10:44.902798Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T04:10:44.909735Z","title":"Omnisam: Omnidirectional segment anything model for uda in panoramic semantic segmentation,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:44.909735Z"},"links":{"citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:4c85e90afeb9adbb5da10f211319c6dcbb1efb9255633b286bf457ffcf444fd9","observation_id":"4a6a9f0e-48ff-4f9f-a18e-cb27fdc23b63","resolution":{"observed_at":"2026-08-16T04:10:44.909735Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T04:10:46.954274Z","title":"Eviprompt: A training-free evidential prompt generation method for adapting segment anything model in medical images,","venue":null,"work_id":"580a960a-2815-486f-abf6-749f5db578d1","year":2024},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:44.920798Z"},"links":{"citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:8848a01b462c0e1579b58a8159dbb8022cdced8bb6cb6a5bcfa5e88013f1ad4d","observation_id":"5ea4386a-25bf-41be-a35c-98920179577a","resolution":{"observed_at":"2026-08-16T04:10:46.961630Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2408.12889","last_updated":"2024-08-23T07:51:10Z","snapshot_observed_at":"2026-08-18T18:23:36.231137Z","submitted_at":"2024-08-23T07:51:10Z","title":"Unleashing the Potential of SAM2 for Biomedical Images and Videos: A Survey","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2408.12889","snapshot_observed_at":"2026-08-16T04:10:44.927618Z","title":"Unleashing the potential of sam2 for biomedical images and videos: A survey,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:44.927618Z"},"links":{"cited_paper":"/paper/2408.12889","citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:82cdec67f5f98a8cfdaced416886c035b24ffb4ea9dc5f693203cb4ccb2230d7","observation_id":"0be9c375-a0f3-43c6-89cb-ac4580fc3eda","resolution":{"observed_at":"2026-08-16T04:10:44.927618Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.03678","last_updated":"2023-08-11T04:23:29Z","snapshot_observed_at":"2026-08-18T18:28:54.267462Z","submitted_at":"2023-05-05T16:48:45Z","title":"Towards Segment Anything Model (SAM) for Medical Image Segmentation: A Survey","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.03678","snapshot_observed_at":"2026-08-16T04:10:44.937378Z","title":"Towards segment anything model (sam) for med- ical image segmentation: a survey,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:44.937378Z"},"links":{"cited_paper":"/paper/2305.03678","citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:a14e4fc777fb94c2d45750dce77458cf31c7c9631280153ebfc7876e283e30d2","observation_id":"9ce79b5a-ce24-4993-8ec2-6e3ae06dec80","resolution":{"observed_at":"2026-08-16T04:10:44.937378Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T04:10:46.934495Z","title":"Segment anything model for medical image segmentation: Current applications and future directions,","venue":null,"work_id":"703d5ecc-59a2-4d28-8064-fa5f26be8e14","year":2024},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:44.944922Z"},"links":{"citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:5013c67f42f3f3c7082f06106a96aa7e5e80439a300ae91a026d4f3d69d6e698","observation_id":"7f6a3f2e-f86f-4a5e-b1db-e3a53995639a","resolution":{"observed_at":"2026-08-16T04:10:46.940700Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T04:10:46.897987Z","title":"Samrs: Scaling-up remote sensing segmentation dataset with segment anything model,","venue":null,"work_id":"c874a55d-6125-4256-ae73-9b95ccdd6af8","year":null},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:44.951304Z"},"links":{"citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:c26ef17ed5a1eda27c82dcc1c707a83f0e537a8215a7aca80c91c59825cd2381","observation_id":"83e130ce-095d-4789-9a7e-ead58424b6a0","resolution":{"observed_at":"2026-08-16T04:10:46.905296Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T04:10:46.872842Z","title":"Ringmo-sam: A foundation model for segment anything in multimodal remote-sensing images,","venue":null,"work_id":"3d11a57a-4369-42a4-9a89-27ff8ead89b4","year":2023},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:44.958433Z"},"links":{"citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:6cc9128aa410e0cdaf0a1f9f9330551ac1ea9809adef9d49d17753e1194f70e1","observation_id":"0444c3ac-a0ef-4644-9252-e1618ff46dc4","resolution":{"observed_at":"2026-08-16T04:10:46.878877Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T04:10:46.850052Z","title":"Rsprompter: Learning to prompt for remote sensing instance seg- mentation based on visual foundation model,","venue":null,"work_id":"879f62c8-925b-4f0e-82da-d9d67d363770","year":2024},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:44.969456Z"},"links":{"citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:70881ba52b5274887cbb0411c5248f20fb6488ee785b06304d27fd3c6cce183c","observation_id":"5ae25380-a782-4140-b221-fdba7964d390","resolution":{"observed_at":"2026-08-16T04:10:46.856152Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.12659","last_updated":"2025-07-08T12:38:02Z","snapshot_observed_at":"2026-08-19T05:01:47.443864Z","submitted_at":"2023-05-22T03:03:29Z","title":"UVOSAM: A Mask-free Paradigm for Unsupervised Video Object Segmentation via Segment Anything Model","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.12659","snapshot_observed_at":"2026-08-16T04:10:44.977868Z","title":"Uvosam: A mask-free paradigm for unsupervised video object segmentation via segment anything model,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:44.977868Z"},"links":{"cited_paper":"/paper/2305.12659","citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:da88f68e81f21e0945659257a6fa136451edaeee2eb25848ae488242c689ad38","observation_id":"fa59431d-835d-467a-976c-ef46ff752fe6","resolution":{"observed_at":"2026-08-16T04:10:44.977868Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T04:10:46.828027Z","title":"Foodsam: Any food segmentation,","venue":null,"work_id":"eea50043-a860-4902-89ef-649cc80a777a","year":2023},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:44.984313Z"},"links":{"citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:1d3c2ca16209b90c035aaf13c641c3f06ca29e7a301cea82a41ef2384f4a44c6","observation_id":"df389029-e716-47a1-81bd-44a13406eda9","resolution":{"observed_at":"2026-08-16T04:10:46.834958Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T04:10:46.795908Z","title":"Recalling unknowns without losing precision: An effective solution to large model-guided open world object detection,","venue":null,"work_id":"dc882d33-1559-4963-9aaa-98dcd6e74e32","year":2025},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:44.991749Z"},"links":{"citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:8dcf9ecce8ddb7dfeb4e1f34f9a80b4dba65952c6e9b089e590a6444204c3b33","observation_id":"8c1f2c6e-63a5-46b7-b2be-8acd4c221633","resolution":{"observed_at":"2026-08-16T04:10:46.811786Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T04:10:46.770153Z","title":"Segmenting anything in the dark via depth perception,","venue":null,"work_id":"d032dba8-bd55-4479-ad79-f25326612293","year":2025},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:44.999172Z"},"links":{"citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:efaaddb0637b41405e35c4322d63c28cbfb0f6e3a34b7f70b6aee3c38dd87034","observation_id":"8f9640cd-08e2-488f-bf8b-6a7714f9aaeb","resolution":{"observed_at":"2026-08-16T04:10:46.778613Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.06558","last_updated":"2023-05-11T04:33:08Z","snapshot_observed_at":"2026-08-16T15:33:58.868647Z","submitted_at":"2023-05-11T04:33:08Z","title":"Segment and Track Anything","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.06558","snapshot_observed_at":"2026-08-16T04:10:45.004955Z","title":"Segment and track anything,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:45.004955Z"},"links":{"cited_paper":"/paper/2305.06558","citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:84187d01e632355db0ac882e2c9766cc457db3c1fc5c40a2ac7f67d7ca1a3845","observation_id":"300e7357-7e4a-482e-b58e-f1d173197387","resolution":{"observed_at":"2026-08-16T04:10:45.004955Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T04:10:46.742367Z","title":"Rog- sam: A language-driven framework for instance-level robotic grasping detection,","venue":null,"work_id":"5f5c81f2-7a89-4802-8c51-976563c4a126","year":2025},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:45.019038Z"},"links":{"citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:f657c8efb48b43d078a5647b0548dcd0796ee776fded12c97478bc292991744b","observation_id":"2ff02e8a-17fd-4ab7-82e3-3edaedf91b8a","resolution":{"observed_at":"2026-08-16T04:10:46.748989Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T04:10:46.714306Z","title":"Frequency-guided spatial adaptation for camouflaged object detection,","venue":null,"work_id":"0b14f5d8-6c36-4736-b38f-282a754405c3","year":2025},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:45.025522Z"},"links":{"citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:618b82806a7f5a12301111505be0772c6e117daebbbb104841a92718980906ad","observation_id":"24aa7df3-7c38-4ab2-bb59-db33063538c6","resolution":{"observed_at":"2026-08-16T04:10:46.725616Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T04:10:46.691313Z","title":"Nto3d: Neural target object 3d reconstruction with segment anything,","venue":null,"work_id":"2f9e4339-9860-4cf1-8696-1d9832dfe8e8","year":2024},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:45.032049Z"},"links":{"citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:dec86534516cd465f9ec658df5572529bae57de25fcd7272d6edcb01f6f804b8","observation_id":"a7a4d24f-1ed1-48f8-9550-5298e147b312","resolution":{"observed_at":"2026-08-16T04:10:46.698958Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T04:10:46.661957Z","title":"Learning transferable visual models from natural language supervision,","venue":null,"work_id":"99ecf531-88d8-4578-b78d-bf80ea2a3798","year":2021},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:45.040952Z"},"links":{"citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:61681f2a6c06a225433fc91f8e022be80a7c8e2e4a724a09937cc759be02afb6","observation_id":"5a93683c-dcf1-4ebb-a259-4b0bcfb8d7af","resolution":{"observed_at":"2026-08-16T04:10:46.673402Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T04:10:46.640018Z","title":"Show, attend and tell: Neural image caption generation with visual attention,","venue":null,"work_id":"033e3e87-de63-44e9-8472-7cf87d6a5700","year":2015},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:45.047451Z"},"links":{"citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:b7f1d8c09b92ed82f280aa675e1285082c652fd2c350cff4b1ec92ffa3fd875e","observation_id":"5961195d-7b88-4f80-b729-0bdb188e417d","resolution":{"observed_at":"2026-08-16T04:10:46.647194Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T04:10:46.616516Z","title":"Instructpix2pix: Learning to follow image editing instructions,","venue":null,"work_id":"884c4584-a892-47b5-a1b9-59342859af5c","year":2023},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:45.058145Z"},"links":{"citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:1f532143ef5ed56a661b43f5f36918c94597e88000c92df89ed23511dac395c8","observation_id":"03093492-b724-463e-8a85-ab5e1057bcf4","resolution":{"observed_at":"2026-08-16T04:10:46.623120Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T04:10:46.586259Z","title":"Unibind: Llm-augmented unified and balanced representation space to bind them all,","venue":null,"work_id":"4517df64-d96f-4c06-8e61-769cc5c29449","year":2024},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:45.067080Z"},"links":{"citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:e236618b7111a94f3cea540319d8327e28c1523c36b904987b1ccba83dacbfe1","observation_id":"27e8dbcd-7c02-4436-a422-6cb1a7d4d8bf","resolution":{"observed_at":"2026-08-16T04:10:46.596108Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2405.16108","last_updated":"2024-05-25T07:45:41Z","snapshot_observed_at":"2026-08-18T18:24:53.484623Z","submitted_at":"2024-05-25T07:45:41Z","title":"OmniBind: Teach to Build Unequal-Scale Modality Interaction for Omni-Bind of All","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.16108","snapshot_observed_at":"2026-08-16T04:10:45.075767Z","title":"Omnibind: Teach to build unequal-scale modality interaction for omni-bind of all,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":53,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:45.075767Z"},"links":{"cited_paper":"/paper/2405.16108","citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:170accce0fb41d0dc65f3049f2ffed3eac5e63623db625f15940dd0549f261a4","observation_id":"dfec9af9-ea0b-4115-8c14-4bf50fd80abf","resolution":{"observed_at":"2026-08-16T04:10:45.075767Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T04:10:46.561507Z","title":"Vision-language consistency guided multi-modal prompt learning for blind ai generated image quality assessment,","venue":null,"work_id":"1d5299c4-1e07-4e60-b0a5-75725431efaf","year":2024},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":54,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:45.082159Z"},"links":{"citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:7701e0e31da1e7cc71ca701fa8bf94b8bd2fd6b6f1ac188f0efb164113b61cff","observation_id":"e186b863-2362-48c0-bc87-8291957eb8da","resolution":{"observed_at":"2026-08-16T04:10:46.567438Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T04:10:46.539899Z","title":"Dall-e: Creating images from text,","venue":null,"work_id":"d8f921b3-f284-4937-a7c1-da2f5d756526","year":2021},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:45.090141Z"},"links":{"citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:0d66a053e229ff7d68edd89995e23bb07478e31e3b7047f0c71bc81e08e87396","observation_id":"aac53fae-cb72-4f65-95e7-8e7562b12f68","resolution":{"observed_at":"2026-08-16T04:10:46.545532Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2303.08774","last_updated":"2024-03-04T06:01:33Z","snapshot_observed_at":"2026-08-17T09:58:46.058102Z","submitted_at":"2023-03-15T17:15:04Z","title":"GPT-4 Technical Report","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2303.08774","snapshot_observed_at":"2026-08-16T04:10:45.096934Z","title":"Gpt-4 technical report,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":56,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:45.096934Z"},"links":{"cited_paper":"/paper/2303.08774","citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:7f69383fdeaea85379f17bf7504cba69248f26209af724ae8b867d25a34e288d","observation_id":"d3beabec-f782-4312-a7cd-d5ce5248086d","resolution":{"observed_at":"2026-08-16T04:10:45.096934Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T04:10:46.517966Z","title":"Multi-modal interaction graph convolutional network for temporal language localization in videos,","venue":null,"work_id":"5404aea5-9dc4-49a7-af89-6dbfad15d571","year":2021},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":57,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:45.104732Z"},"links":{"citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:61c9182236df074d3eeae9aec8b56e55d6b7265d42413a486f248715f38820b9","observation_id":"e7ac2a3b-2747-4d28-a897-5625e9f4c26e","resolution":{"observed_at":"2026-08-16T04:10:46.524962Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T04:10:46.488217Z","title":"Prompt-driven referring image segmentation with instance contrasting,","venue":null,"work_id":"62b3582b-7dad-4542-8903-281a70f58fbe","year":2024},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":58,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:45.110919Z"},"links":{"citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:65d7ee11cafd220c7313d86390b5996e0a0e6bd19465dced9b2681d4d285b45c","observation_id":"0b397359-6bd5-4977-9a93-b5703f127415","resolution":{"observed_at":"2026-08-16T04:10:46.499615Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T04:10:46.466057Z","title":"Cris: Clip- driven referring image segmentation,","venue":null,"work_id":"c5ca43c1-bf5c-4ac8-b62d-f82b71d98da2","year":2022},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":59,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:45.116751Z"},"links":{"citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:d8f2b1c710d7714c05667574a0d9da6f0b91ba62642d8942f7965aa17fddffe3","observation_id":"47ce4181-571c-401a-83c2-d2f2a9f445f4","resolution":{"observed_at":"2026-08-16T04:10:46.472334Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T04:10:46.439365Z","title":"Egfnet: Edge-aware guidance fusion network for rgb–thermal urban scene parsing,","venue":null,"work_id":"ef3ddded-1671-4d2c-a793-c35a2f3b4942","year":2024},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":60,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:45.123125Z"},"links":{"citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:d02aec4ea54a695f42304f25247aeb183c81a2f02883a5098b93d5b812edf95c","observation_id":"12c77383-f737-4932-a7f4-312a9129c48e","resolution":{"observed_at":"2026-08-16T04:10:46.445348Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T04:10:46.419392Z","title":"Abmdrnet: Adaptive-weighted bi-directional modality difference reduction network for rgb-t semantic segmentation,","venue":null,"work_id":"b1247003-629c-4e9c-83a0-7bdadde0ec5c","year":2021},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":61,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:45.128960Z"},"links":{"citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:f0545622b7a2959e4a2bf8da9a6a673323300ea2af2e8eb0a7b1367c902d9c07","observation_id":"174337df-9f03-4a3b-99a7-81e2ff94cbfd","resolution":{"observed_at":"2026-08-16T04:10:46.426106Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T04:10:46.397466Z","title":"Feanet: Feature-enhanced attention network for rgb-thermal real-time semantic segmentation,","venue":null,"work_id":"8800fd4d-04f7-46f7-a2df-09b480edc204","year":2021},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":62,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:45.134934Z"},"links":{"citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:10edfa1cdd6474595e5decc6929c473b6c604dadebd9f2b775599571a3423460","observation_id":"19296c66-b8ea-4705-8b9b-cd1284386503","resolution":{"observed_at":"2026-08-16T04:10:46.406061Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T04:10:46.364942Z","title":"Dbcnet: Dynamic bilateral cross- fusion network for rgb-t urban scene understanding in intelligent vehi- cles,","venue":null,"work_id":"726d4f79-5986-44e0-a4e0-a0724b65d412","year":2023},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":63,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:45.144142Z"},"links":{"citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:f36b92f5bfc90464aab75ee5d2989b8064c6de2037e5925605ae5a76193a42d7","observation_id":"8cd05749-82e2-478c-ae88-343df44cb905","resolution":{"observed_at":"2026-08-16T04:10:46.376814Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T04:10:46.341657Z","title":"Ex- plicit attention-enhanced fusion for rgb-thermal perception tasks,","venue":null,"work_id":"32f27dc1-5de8-45c2-b060-5c25ba63956f","year":2023},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":64,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:45.150118Z"},"links":{"citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:c52ed279200134ae26b2e54edd245ba58ab6a830fead3b545d3f3431148ee3ae","observation_id":"77f5fdc5-44bf-4fb4-9cb7-a50109bcc401","resolution":{"observed_at":"2026-08-16T04:10:46.347638Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T04:10:46.320726Z","title":"Gmnet: Graded- feature multilabel-learning network for rgb-thermal urban scene seman- tic segmentation,","venue":null,"work_id":"0069d591-a8dd-4e28-a4f5-8e719fded279","year":2021},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":65,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:45.157135Z"},"links":{"citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:d37bc05a258efa8aae9fae3e85fc3a926ae705b9ed903c3cbb7dcc87955cf50d","observation_id":"0d0955f8-1318-477b-8a1f-9ab6892bff52","resolution":{"observed_at":"2026-08-16T04:10:46.326328Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T04:10:46.294805Z","title":"Mmsformer: Multi- modal transformer for material and semantic segmentation,","venue":null,"work_id":"64a2c1d3-82dc-4440-9e50-7b849fb9cc41","year":2024},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":66,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:45.166797Z"},"links":{"citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:925796e1b133611dd0bdd9cce2c848360a4d195368ff01d4dab7da4e2b5bebed","observation_id":"239142f4-a9b2-4e58-8b97-3ef6c9559cf5","resolution":{"observed_at":"2026-08-16T04:10:46.305905Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T04:10:46.263828Z","title":"Complementary random masking for rgb-thermal semantic segmentation,","venue":null,"work_id":"3abda350-fb99-43d7-a549-ad7d8ebaca64","year":2024},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":67,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:45.174987Z"},"links":{"citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:a7b2b3e7535b7bed204e1c26ef5466600ddfe93dba2953715ff32fdb3583ec70","observation_id":"72409809-0207-4672-bc19-b74a5525f9e8","resolution":{"observed_at":"2026-08-16T04:10:46.271937Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T04:10:45.181238Z","title":"Complementary random masking for rgb-thermal semantic segmentation,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":68,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:45.181238Z"},"links":{"citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:7d405cce76af1435821dace9f5cd5f6a8984e73a270d113915e0dc71bc584f52","observation_id":"d3df1989-d650-4e9e-813a-475e5fd81df3","resolution":{"observed_at":"2026-08-16T04:10:45.181238Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T04:10:46.239599Z","title":"Decoupled weight decay regularization,","venue":null,"work_id":"dc20e671-72c3-4c37-b257-691d4a0500de","year":2019},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":69,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:45.188504Z"},"links":{"citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:eabaa152b403ad4bba4f66e58e1bdd7f482cb543499c4a1bd33dfb081f30b0d9","observation_id":"ee90a38b-0176-4cc3-899c-dcf9ef9a57fb","resolution":{"observed_at":"2026-08-16T04:10:46.245965Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T04:10:46.219043Z","title":"Multi-interactive feature learning and a full-time multi-modality bench- mark for image fusion and segmentation,","venue":null,"work_id":"0fd4dc93-1316-4f8a-8cd9-94dc2dec1b64","year":2023},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":70,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:45.194680Z"},"links":{"citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:db58d02ad7be3a8c7a0a32c131615ae8f8bf29850ca8f893ca721626a73aaeef","observation_id":"e7839e16-9f43-42f5-9cee-93fb1e104fe9","resolution":{"observed_at":"2026-08-16T04:10:46.225420Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T04:10:46.197771Z","title":"Cmx: Cross-modal fusion for rgb-x semantic segmentation with transformers,","venue":null,"work_id":"6a0bc2d3-f0d8-41ef-89a0-b781b7ee7426","year":2023},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":71,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:45.201208Z"},"links":{"citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:52225d87daad2c5236a4b0c33081a03b6ec92b7c037ab21be0a88cdfa521cb86","observation_id":"82aaa764-a5af-4f02-8689-73fe75d0bf53","resolution":{"observed_at":"2026-08-16T04:10:46.203550Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T04:10:46.178969Z","title":"Delivering arbitrary-modal semantic segmentation,","venue":null,"work_id":"ce3e4fff-b3cb-4e20-946e-b0c544f56bde","year":2023},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":72,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:45.213773Z"},"links":{"citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:d1741ec3076c2657600204e21175dc753bd9a6b59536a2d9dcaa5545cff0569e","observation_id":"e06ab98c-3b43-4aa5-aa1c-d327e8fdd0cf","resolution":{"observed_at":"2026-08-16T04:10:46.185286Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T04:10:46.154534Z","title":"Gmnet: graded-feature multilabel-learning network for rgb-thermal urban scene semantic seg- mentation,","venue":null,"work_id":"ad60c8bf-15cb-4e11-a2a0-b0cebbdbae79","year":2021},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":73,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:45.222621Z"},"links":{"citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:a84e1471e40f4b25b9bf3d0db06d6bba00a68709c4fb7f7bcb586fe8ae2d9c76","observation_id":"37a61209-4085-4ecb-b5eb-0ab4c6d7f917","resolution":{"observed_at":"2026-08-16T04:10:46.161549Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T04:10:46.132102Z","title":"Rgb-t semantic segmentation with location, activation, and sharpening,","venue":null,"work_id":"1e2d88fc-5a80-4f35-a3a6-1823e907a32e","year":null},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":74,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:45.229938Z"},"links":{"citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:8af8bfe287d6c3e2015fd9363dc5be2be88569785736a69ba612d24736a5b08c","observation_id":"dd36547d-fecd-44d4-af15-f799377dcaaa","resolution":{"observed_at":"2026-08-16T04:10:46.138101Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T04:10:46.111296Z","title":"Edge-aware guidance fusion network for rgb thermal scene parsing,","venue":null,"work_id":"e9203df0-eb46-4fb6-9797-97450f94573d","year":2022},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":75,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:45.236280Z"},"links":{"citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:fbb42e877c5959a17ecce8461a6ea4bdbfaf04962940ad48d7eed5d71c130133","observation_id":"ef0b2269-d5a8-4707-afde-352b50ba5344","resolution":{"observed_at":"2026-08-16T04:10:46.118173Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T04:10:46.087636Z","title":"Didfuse: Deep image decomposition for infrared and visible image fusion,","venue":null,"work_id":"16c12017-a022-4b37-8548-92483d1b74bf","year":null},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":76,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:45.244586Z"},"links":{"citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:f224768f2a7ca9b4dc3d02bf3fe961c8d1424634bd9a772c9c1261f1bb5cb518","observation_id":"d80a0fc8-349a-4c16-b8ac-192c3892471c","resolution":{"observed_at":"2026-08-16T04:10:46.095017Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T04:10:46.066003Z","title":"Reconet: Recurrent correction network for fast and efficient multi-modality image fusion,","venue":null,"work_id":"1838dcb9-aefe-4d59-b02d-1034cb348b44","year":2022},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":77,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:45.253611Z"},"links":{"citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:bd2d8f74a73d3ce5947049f2892597830e71e0ef62db7b5b535cc907b00256f4","observation_id":"1ba369f0-8312-441c-bdfc-036cd6882ac6","resolution":{"observed_at":"2026-08-16T04:10:46.073290Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T04:10:46.038627Z","title":"U2fusion: A unified unsupervised image fusion network,","venue":null,"work_id":"9a5a0699-02aa-44ec-a515-db405584a2da","year":2020},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":78,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:45.264173Z"},"links":{"citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:71a1d6e987cd55cdeb9d174508d5bcf52de98f1e4626381ce72aa1d5445f6314","observation_id":"7b51b893-f08a-4d9d-a7a4-c0ad91730ec6","resolution":{"observed_at":"2026-08-16T04:10:46.045722Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-16T04:10:46.005505Z","title":"Target-aware dual adversarial learning and a multi-scenario multi- modality benchmark to fuse infrared and visible for object detection,","venue":null,"work_id":"3f4c297f-b56a-4a35-9377-59325df3709e","year":2022},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":79,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:45.273385Z"},"links":{"citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:3b1ceefa5c9871aff95ec5907ed231302fa44ec264ad706dfbe5f0965dbd7170","observation_id":"a3555f3b-f58f-4f9f-883c-f79e477159e5","resolution":{"observed_at":"2026-08-16T04:10:46.016428Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-19T06:32:44.657259+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2405.15365","last_updated":"2024-05-24T08:58:48Z","snapshot_observed_at":"2026-08-18T18:28:54.916810Z","submitted_at":"2024-05-24T08:58:48Z","title":"U3M: Unbiased Multiscale Modal Fusion Model for Multimodal Semantic Segmentation","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.15365","snapshot_observed_at":"2026-08-16T04:10:45.280927Z","title":"U3m: Unbiased multiscale modal fusion model for multimodal semantic segmentation,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation","version":1},"reference_index":80,"source":"pdf_text","source_observed_at":"2026-08-16T04:10:45.280927Z"},"links":{"cited_paper":"/paper/2405.15365","citing_paper":"/paper/2505.01950"},"observation_digest":"sha256:8cdb3a39bea6d17d9beace80857826706dadbbcfb4db7ec13cd4df62871c1cfb","observation_id":"93c0199c-b247-4990-945e-bcf7f923d1c9","resolution":{"observed_at":"2026-08-16T04:10:45.280927Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2505.01950","last_updated":"2025-05-04T00:24:17Z","latest_version":1,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-19T15:10:05.707696Z","submitted_at":"2025-05-04T00:24:17Z","title":"Segment Any RGB-Thermal Model with Language-aided Distillation"},"reference_resolution":{"displayed":80,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":18,"verified_exact":0,"verified_fuzzy":62},"total_outbound_references":80},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-19T06:32:44.657259+00:00","source":"crossref"},{"observed_at":"2026-08-19T06:32:39.956319+00:00","source":"retraction_watch"}],"thesis":"As of 19 August 2026, this Paper Citation Record lists 80 of 80 outbound references and 0 inbound Pith citation observations for arXiv:2505.01950."}