{"as_of":"2026-08-13T21:35:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:5dcda40baf80918e87f52a3db62548e1e9160a94fadf8cad56c0f11d8ed24b4f","coverage":[{"denominator":49,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":49,"source":"paper_references, paper_reference_links","source_observed_at":"2026-05-17T05:23:11.261860Z","state":"measured"},{"denominator":53,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":53,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-13T06:32:02.005865+00:00","state":"measured"},{"denominator":4,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":4,"source":"paper_references, paper_reference_links","source_observed_at":"2026-07-13T14:26:43.963945Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":1,"source":"pith","source_observed_at":"2026-08-05T02:28:24.338817Z","state":"measured"}],"external_citation_measurements":[{"count":0,"observed_at":"2026-08-05T02:28:24.338817Z","source":"pith"}],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2511.19704","last_updated":"2026-04-10T00:50:03Z","snapshot_observed_at":"2026-08-11T14:26:55.408955Z","submitted_at":"2025-11-24T21:15:01Z","title":"RADSeg: Unleashing Parameter and Compute Efficient Zero-Shot Open-Vocabulary Segmentation Using Agglomerative Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2511.19704","snapshot_observed_at":"2026-07-13T14:26:43.963945Z","title":"arXiv:2511.19704 (2025)","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2604.01388","last_updated":"2026-07-09T16:05:00Z","snapshot_observed_at":"2026-08-09T23:26:10.253620Z","submitted_at":"2026-04-01T20:48:06Z","title":"LESV: Language Embedded Sparse Voxel Fusion for Open-Vocabulary 3D Scene Understanding","version":2},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-07-13T14:26:43.963945Z"},"links":{"cited_paper":"/paper/2511.19704","citing_paper":"/paper/2604.01388"},"observation_digest":"sha256:bc578ccf92056de6293633c6b4770aa0e52d17b7b60459b6f01fc425c83870fd","observation_id":"0d7d0866-bcea-4831-8804-65cf45051247","resolution":{"observed_at":"2026-07-13T14:26:43.963945Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2511.19704","last_updated":"2026-04-10T00:50:03Z","snapshot_observed_at":"2026-08-11T14:26:55.408955Z","submitted_at":"2025-11-24T21:15:01Z","title":"RADSeg: Unleashing Parameter and Compute Efficient Zero-Shot Open-Vocabulary Segmentation Using Agglomerative Models","version":2},"cited_work":{"arxiv_id":"2511.19704","doi":"10.48550/arxiv.2511.19704","metadata_source":"pith","pith_arxiv_id":"2511.19704","snapshot_observed_at":"2026-08-05T02:49:54.815029Z","title":"RADSeg: Unleashing Parameter and Compute Efficient Zero-Shot Open-Vocabulary Segmentation Using Agglomerative Models","venue":"cs.CV","work_id":"9e5eb98d-1b2d-48f4-8c40-8444afe6c7bc","year":2025},"citing_paper":{"arxiv_id":"2604.26067","last_updated":"2026-04-28T19:09:09Z","snapshot_observed_at":"2026-08-10T22:42:46.696246Z","submitted_at":"2026-04-28T19:09:09Z","title":"RADIO-ViPE: Online Tightly Coupled Multi-Modal Fusion for Open-Vocabulary Semantic SLAM in Dynamic Environments","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-05-07T16:40:25.188693Z"},"links":{"cited_paper":"/paper/2511.19704","citing_paper":"/paper/2604.26067"},"observation_digest":"sha256:24aaacb4f2be82ec3adc4e24b7925baccfff17f15f78afec6a65bb0de172af13","observation_id":"cc6c75d3-cd97-487a-996b-3c52961202a2","resolution":{"observed_at":"2026-05-11T23:36:28.729353Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.19704","last_updated":"2026-04-10T00:50:03Z","snapshot_observed_at":"2026-08-11T14:26:55.408955Z","submitted_at":"2025-11-24T21:15:01Z","title":"RADSeg: Unleashing Parameter and Compute Efficient Zero-Shot Open-Vocabulary Segmentation Using Agglomerative Models","version":2},"cited_work":{"arxiv_id":"2511.19704","doi":"10.48550/arxiv.2511.19704","metadata_source":"pith","pith_arxiv_id":"2511.19704","snapshot_observed_at":"2026-08-05T02:49:54.815029Z","title":"RADSeg: Unleashing Parameter and Compute Efficient Zero-Shot Open-Vocabulary Segmentation Using Agglomerative Models","venue":"cs.CV","work_id":"9e5eb98d-1b2d-48f4-8c40-8444afe6c7bc","year":2025},"citing_paper":{"arxiv_id":"2605.03669","last_updated":"2026-05-05T12:08:16Z","snapshot_observed_at":"2026-08-12T16:42:13.233251Z","submitted_at":"2026-05-05T12:08:16Z","title":"FUS3DMaps: Scalable and Accurate Open-Vocabulary Semantic Mapping by 3D Fusion of Voxel- and Instance-Level Layers","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-05-07T15:47:11.780364Z"},"links":{"cited_paper":"/paper/2511.19704","citing_paper":"/paper/2605.03669"},"observation_digest":"sha256:59e85d69f2037ba3aa506e142bcae861c3a4c228f3f7366fbca731e8b96801cb","observation_id":"e0f082e7-9fb8-4bf4-890c-b7a35796bcda","resolution":{"observed_at":"2026-05-12T10:56:30.070157Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.19704","last_updated":"2026-04-10T00:50:03Z","snapshot_observed_at":"2026-08-11T14:26:55.408955Z","submitted_at":"2025-11-24T21:15:01Z","title":"RADSeg: Unleashing Parameter and Compute Efficient Zero-Shot Open-Vocabulary Segmentation Using Agglomerative Models","version":2},"cited_work":{"arxiv_id":"2511.19704","doi":"10.48550/arxiv.2511.19704","metadata_source":"pith","pith_arxiv_id":"2511.19704","snapshot_observed_at":"2026-08-05T02:49:54.815029Z","title":"RADSeg: Unleashing Parameter and Compute Efficient Zero-Shot Open-Vocabulary Segmentation Using Agglomerative Models","venue":"cs.CV","work_id":"9e5eb98d-1b2d-48f4-8c40-8444afe6c7bc","year":2025},"citing_paper":{"arxiv_id":"2606.19088","last_updated":"2026-06-17T13:58:06Z","snapshot_observed_at":"2026-08-03T06:49:13.095613Z","submitted_at":"2026-06-17T13:58:06Z","title":"ReSiReg: Towards Spatially Consistent Semantics in Language-Conditioned Robotic Tasks","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-06-26T20:53:01.922037Z"},"links":{"cited_paper":"/paper/2511.19704","citing_paper":"/paper/2606.19088"},"observation_digest":"sha256:b0217f2e61c697efbb208ffb510788590b0926dc823f709a8d117e68b4261c7b","observation_id":"bcd13e96-8847-4f55-a02d-dbb4ee1fad0a","resolution":{"observed_at":"2026-06-26T20:59:58.021464Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2511.19704/citation-record","integrity":"/paper/2511.19704/integrity","json":"/paper/2511.19704/citation-record.json","paper":"/paper/2511.19704"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2504.06994","last_updated":"2025-04-09T16:06:58Z","snapshot_observed_at":"2026-08-07T16:07:08.599268Z","submitted_at":"2025-04-09T16:06:58Z","title":"RayFronts: Open-Set Semantic Ray Frontiers for Online Scene Understanding and Exploration","version":1},"cited_work":{"arxiv_id":"2504.06994","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2504.06994","snapshot_observed_at":"2026-07-01T19:36:09.401596Z","title":"Rayfronts: Open-set semantic ray frontiers for online scene understanding and ex- ploration","venue":null,"work_id":"f4da4aac-9dc3-4d63-a190-4b2584f0022b","year":2025},"citing_paper":{"arxiv_id":"2511.19704","last_updated":"2026-04-10T00:50:03Z","snapshot_observed_at":"2026-08-11T14:26:55.408955Z","submitted_at":"2025-11-24T21:15:01Z","title":"RADSeg: Unleashing Parameter and Compute Efficient Zero-Shot Open-Vocabulary Segmentation Using Agglomerative Models","version":2},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-05-17T05:23:11.261860Z"},"links":{"cited_paper":"/paper/2504.06994","citing_paper":"/paper/2511.19704"},"observation_digest":"sha256:729829506c94c394906dae4bb68e1d81566ea8bddf28f1653a5083c9a32055fd","observation_id":"c8f71eb9-6fdb-4c51-8cff-439fb0995b06","resolution":{"observed_at":"2026-05-17T05:24:04.755631Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Single-stage seman- tic segmentation from image labels","venue":null,"work_id":"7e6d18a0-df5b-4041-b719-485996d667e9","year":2020},"citing_paper":{"arxiv_id":"2511.19704","last_updated":"2026-04-10T00:50:03Z","snapshot_observed_at":"2026-08-11T14:26:55.408955Z","submitted_at":"2025-11-24T21:15:01Z","title":"RADSeg: Unleashing Parameter and Compute Efficient Zero-Shot Open-Vocabulary Segmentation Using Agglomerative Models","version":2},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-05-17T05:23:11.261860Z"},"links":{"citing_paper":"/paper/2511.19704"},"observation_digest":"sha256:423a4e7f00254b5c82c472a982c13d8260210143e2bb60e5718b8e5d71939e4b","observation_id":"e2e7def3-c29a-44c1-85dd-dfd761a66b5b","resolution":{"observed_at":"2026-05-17T05:24:05.285841Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2411.15869","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"IEEE Transactions on Image Processing34, 8271–8284 (2025) arXiv:2411.15869 [cs.CV]","venue":null,"work_id":"eb3b0b97-4f81-45f7-a75b-048963a5fbe9","year":2025},"citing_paper":{"arxiv_id":"2511.19704","last_updated":"2026-04-10T00:50:03Z","snapshot_observed_at":"2026-08-11T14:26:55.408955Z","submitted_at":"2025-11-24T21:15:01Z","title":"RADSeg: Unleashing Parameter and Compute Efficient Zero-Shot Open-Vocabulary Segmentation Using Agglomerative Models","version":2},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-05-17T05:23:11.261860Z"},"links":{"citing_paper":"/paper/2511.19704"},"observation_digest":"sha256:1cfc5140b936da41829df8bc03c263cf55d2451d7e3d7c66d7c167ab196c8527","observation_id":"647c05f1-97fd-43b5-8962-7acf1d4d6637","resolution":{"observed_at":"2026-05-17T05:24:04.760277Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2411.19331","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Talking to dino: Bridging self- supervised vision backbones with language for open- vocabulary segmentation.arXiv preprint arXiv:2411.19331","venue":null,"work_id":"472a321f-06dc-4408-b48d-88785ff4eda0","year":null},"citing_paper":{"arxiv_id":"2511.19704","last_updated":"2026-04-10T00:50:03Z","snapshot_observed_at":"2026-08-11T14:26:55.408955Z","submitted_at":"2025-11-24T21:15:01Z","title":"RADSeg: Unleashing Parameter and Compute Efficient Zero-Shot Open-Vocabulary Segmentation Using Agglomerative Models","version":2},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-05-17T05:23:11.261860Z"},"links":{"citing_paper":"/paper/2511.19704"},"observation_digest":"sha256:6501cca0425fcb15b80159b4286cd51562e4ba0daaa906d126a6f1e64b369925","observation_id":"63c21d26-ea17-43fc-93a2-626b39e86003","resolution":{"observed_at":"2026-05-17T05:24:04.708236Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Coco- stuff: Thing and stuff classes in context","venue":null,"work_id":"7fab3395-58a6-4a16-951f-1ada1d1fd424","year":2018},"citing_paper":{"arxiv_id":"2511.19704","last_updated":"2026-04-10T00:50:03Z","snapshot_observed_at":"2026-08-11T14:26:55.408955Z","submitted_at":"2025-11-24T21:15:01Z","title":"RADSeg: Unleashing Parameter and Compute Efficient Zero-Shot Open-Vocabulary Segmentation Using Agglomerative Models","version":2},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-05-17T05:23:11.261860Z"},"links":{"citing_paper":"/paper/2511.19704"},"observation_digest":"sha256:03f6d503271b13a067d5c9fa491ad37bc0cf46242544e2eccb83c0c0d97873fe","observation_id":"df815701-8a5e-4cd9-b5c2-00976c2c05ee","resolution":{"observed_at":"2026-05-17T05:24:05.327197Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Emerg- ing properties in self-supervised vision transformers","venue":null,"work_id":"8a5a47d5-320f-431e-953a-e774104b82a0","year":2021},"citing_paper":{"arxiv_id":"2511.19704","last_updated":"2026-04-10T00:50:03Z","snapshot_observed_at":"2026-08-11T14:26:55.408955Z","submitted_at":"2025-11-24T21:15:01Z","title":"RADSeg: Unleashing Parameter and Compute Efficient Zero-Shot Open-Vocabulary Segmentation Using Agglomerative Models","version":2},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-05-17T05:23:11.261860Z"},"links":{"citing_paper":"/paper/2511.19704"},"observation_digest":"sha256:c260a6694f9f4011c1c2924d7f22284ee8e255804c8929b5c67be3124397047a","observation_id":"e6f67200-9ba2-49b9-93a6-b44b478aa6fe","resolution":{"observed_at":"2026-05-17T05:24:05.329673Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"The cityscapes dataset for semantic urban scene understanding","venue":null,"work_id":"afbb61e5-38bb-4397-9a7e-7aa6ec17122e","year":2016},"citing_paper":{"arxiv_id":"2511.19704","last_updated":"2026-04-10T00:50:03Z","snapshot_observed_at":"2026-08-11T14:26:55.408955Z","submitted_at":"2025-11-24T21:15:01Z","title":"RADSeg: Unleashing Parameter and Compute Efficient Zero-Shot Open-Vocabulary Segmentation Using Agglomerative Models","version":2},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-05-17T05:23:11.261860Z"},"links":{"citing_paper":"/paper/2511.19704"},"observation_digest":"sha256:8dfbf1817eea985166dd2789c039a38854e9336d38303837cb42dc101445571c","observation_id":"b57ba03f-6c22-4923-9438-a2487e0ed62b","resolution":{"observed_at":"2026-05-17T05:24:05.278893Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Scannet: Richly-annotated 3d reconstructions of indoor scenes","venue":null,"work_id":"35fdd3b3-df49-4e3b-8e5d-d2bbede0ba89","year":2017},"citing_paper":{"arxiv_id":"2511.19704","last_updated":"2026-04-10T00:50:03Z","snapshot_observed_at":"2026-08-11T14:26:55.408955Z","submitted_at":"2025-11-24T21:15:01Z","title":"RADSeg: Unleashing Parameter and Compute Efficient Zero-Shot Open-Vocabulary Segmentation Using Agglomerative Models","version":2},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-05-17T05:23:11.261860Z"},"links":{"citing_paper":"/paper/2511.19704"},"observation_digest":"sha256:58d9e947298b0b67517d24e13af67f03ee8724d08fba87bb1446495655cf1b6a","observation_id":"d06a3af7-4e68-427c-b982-04e7ccb36938","resolution":{"observed_at":"2026-05-17T05:24:05.325049Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2309.16588","last_updated":"2024-04-12T09:38:33Z","snapshot_observed_at":"2026-08-07T09:40:32.614733Z","submitted_at":"2023-09-28T16:45:46Z","title":"Vision Transformers Need Registers","version":2},"cited_work":{"arxiv_id":"2309.16588","doi":"10.48550/arxiv.2309.16588","metadata_source":"pith","pith_arxiv_id":"2309.16588","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Vision Transformers Need Registers","venue":"cs.CV","work_id":"57106da4-5420-4778-94eb-e821589aa7a0","year":2023},"citing_paper":{"arxiv_id":"2511.19704","last_updated":"2026-04-10T00:50:03Z","snapshot_observed_at":"2026-08-11T14:26:55.408955Z","submitted_at":"2025-11-24T21:15:01Z","title":"RADSeg: Unleashing Parameter and Compute Efficient Zero-Shot Open-Vocabulary Segmentation Using Agglomerative Models","version":2},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-05-17T05:23:11.261860Z"},"links":{"cited_paper":"/paper/2309.16588","citing_paper":"/paper/2511.19704"},"observation_digest":"sha256:0485cb940f981b4275b57e6cb9652777ef013f8b035dbaec17c5a3ef6a8c8ba8","observation_id":"2489931e-25b4-4a7a-821b-03dd6d39e456","resolution":{"observed_at":"2026-05-17T05:24:04.749921Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2010.11929","last_updated":"2021-06-03T13:08:56Z","snapshot_observed_at":"2026-08-13T14:19:26.598265Z","submitted_at":"2020-10-22T17:55:59Z","title":"An Image is Worth 16x16 Words: Transformers for Image Recognition at Scale","version":2},"cited_work":{"arxiv_id":"2010.11929","doi":"10.1175/jcli-d-22-0357.1","metadata_source":"pith","pith_arxiv_id":"2010.11929","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"An Image is Worth 16x16 Words: Transformers for Image Recognition at Scale","venue":"cs.CV","work_id":"e96730e3-129b-4db6-b981-15ab7932e297","year":2020},"citing_paper":{"arxiv_id":"2511.19704","last_updated":"2026-04-10T00:50:03Z","snapshot_observed_at":"2026-08-11T14:26:55.408955Z","submitted_at":"2025-11-24T21:15:01Z","title":"RADSeg: Unleashing Parameter and Compute Efficient Zero-Shot Open-Vocabulary Segmentation Using Agglomerative Models","version":2},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-05-17T05:23:11.261860Z"},"links":{"cited_paper":"/paper/2010.11929","citing_paper":"/paper/2511.19704"},"observation_digest":"sha256:0f765f47c1a352674692a2de473d4aa468e6c2f1a8d86213317628d975d1fa64","observation_id":"1dd36494-719d-4497-93fd-4d0fcc561601","resolution":{"observed_at":"2026-05-17T05:24:04.770455Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Williams, John Winn, and Andrew Zisserman","venue":null,"work_id":"1c5f80b0-c3b4-4062-b9ad-e261b22057af","year":2010},"citing_paper":{"arxiv_id":"2511.19704","last_updated":"2026-04-10T00:50:03Z","snapshot_observed_at":"2026-08-11T14:26:55.408955Z","submitted_at":"2025-11-24T21:15:01Z","title":"RADSeg: Unleashing Parameter and Compute Efficient Zero-Shot Open-Vocabulary Segmentation Using Agglomerative Models","version":2},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-05-17T05:23:11.261860Z"},"links":{"citing_paper":"/paper/2511.19704"},"observation_digest":"sha256:2d25e76663dd150941b4c22189c010c2057cec73f87194c9cb2b44a65ed8d928","observation_id":"dbe09bf0-4f86-4680-b4a2-b8920020fd08","resolution":{"observed_at":"2026-05-17T05:24:05.274029Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Conceptgraphs: Open-vocabulary 3d scene graphs for per- ception and planning","venue":null,"work_id":"47518a0b-228c-4d63-b03f-c8ba3bda40b4","year":2024},"citing_paper":{"arxiv_id":"2511.19704","last_updated":"2026-04-10T00:50:03Z","snapshot_observed_at":"2026-08-11T14:26:55.408955Z","submitted_at":"2025-11-24T21:15:01Z","title":"RADSeg: Unleashing Parameter and Compute Efficient Zero-Shot Open-Vocabulary Segmentation Using Agglomerative Models","version":2},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-05-17T05:23:11.261860Z"},"links":{"citing_paper":"/paper/2511.19704"},"observation_digest":"sha256:102881e1881985ad5313d77a40880020fced3e95227c7156185f00be42b08e09","observation_id":"9737e210-af77-46cf-a1f4-104f23d09cfe","resolution":{"observed_at":"2026-05-17T05:24:05.342348Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Pay attention to your neighbours: Training-free open-vocabulary semantic segmentation","venue":null,"work_id":"78ad93c8-be36-4327-889d-2bdc070c7e91","year":2025},"citing_paper":{"arxiv_id":"2511.19704","last_updated":"2026-04-10T00:50:03Z","snapshot_observed_at":"2026-08-11T14:26:55.408955Z","submitted_at":"2025-11-24T21:15:01Z","title":"RADSeg: Unleashing Parameter and Compute Efficient Zero-Shot Open-Vocabulary Segmentation Using Agglomerative Models","version":2},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-05-17T05:23:11.261860Z"},"links":{"citing_paper":"/paper/2511.19704"},"observation_digest":"sha256:2f810fb3f3063e0498420d96c8dce7dc15319f5d4ff8d95bd4ee9c92bc65d343","observation_id":"1873077d-6929-4cbf-9a67-cc46a241b22f","resolution":{"observed_at":"2026-05-17T05:24:05.320913Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Radiov2.5: Improved baselines for agglomerative vision foundation models","venue":null,"work_id":"e23a795c-e270-46b4-92e5-d6e99596aa07","year":2025},"citing_paper":{"arxiv_id":"2511.19704","last_updated":"2026-04-10T00:50:03Z","snapshot_observed_at":"2026-08-11T14:26:55.408955Z","submitted_at":"2025-11-24T21:15:01Z","title":"RADSeg: Unleashing Parameter and Compute Efficient Zero-Shot Open-Vocabulary Segmentation Using Agglomerative Models","version":2},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-05-17T05:23:11.261860Z"},"links":{"citing_paper":"/paper/2511.19704"},"observation_digest":"sha256:fcd7f3d03f52e8bb61dcc5c49a7720cc6cb8c01551406d5243bfb2aeb2e7b3a6","observation_id":"ecb71844-f89f-4c7a-982c-d0afb0c776d9","resolution":{"observed_at":"2026-05-17T05:24:05.316501Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Tenenbaum, Celso Miguel de Melo, Madhava Krishna, Liam Paull, Florian Shkurti, and Antonio Torralba","venue":null,"work_id":"7f11860e-ee2d-42dc-93b9-a11177c931be","year":2023},"citing_paper":{"arxiv_id":"2511.19704","last_updated":"2026-04-10T00:50:03Z","snapshot_observed_at":"2026-08-11T14:26:55.408955Z","submitted_at":"2025-11-24T21:15:01Z","title":"RADSeg: Unleashing Parameter and Compute Efficient Zero-Shot Open-Vocabulary Segmentation Using Agglomerative Models","version":2},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-05-17T05:23:11.261860Z"},"links":{"citing_paper":"/paper/2511.19704"},"observation_digest":"sha256:c4c74c627b3e412eac792f3ab2ba7be649ee08129ca8c8232449194e7cdaf122","observation_id":"50ba426c-5c12-4cf0-9f88-299895497732","resolution":{"observed_at":"2026-05-17T05:24:05.331948Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"14490449-53de-485b-9dc3-374d6c4260c1","year":2025},"citing_paper":{"arxiv_id":"2511.19704","last_updated":"2026-04-10T00:50:03Z","snapshot_observed_at":"2026-08-11T14:26:55.408955Z","submitted_at":"2025-11-24T21:15:01Z","title":"RADSeg: Unleashing Parameter and Compute Efficient Zero-Shot Open-Vocabulary Segmentation Using Agglomerative Models","version":2},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-05-17T05:23:11.261860Z"},"links":{"citing_paper":"/paper/2511.19704"},"observation_digest":"sha256:3e8a5964db5082dd1889114bb8c6c17c7df4eec2e4fe8eef102d1b2e770f12a8","observation_id":"f2a9fd16-da53-4c76-9f88-396e1d373239","resolution":{"observed_at":"2026-05-17T05:24:05.322934Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Garfield: Group anything with radiance fields","venue":null,"work_id":"40ec6b86-7b62-464c-9575-26051734d898","year":2024},"citing_paper":{"arxiv_id":"2511.19704","last_updated":"2026-04-10T00:50:03Z","snapshot_observed_at":"2026-08-11T14:26:55.408955Z","submitted_at":"2025-11-24T21:15:01Z","title":"RADSeg: Unleashing Parameter and Compute Efficient Zero-Shot Open-Vocabulary Segmentation Using Agglomerative Models","version":2},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-05-17T05:23:11.261860Z"},"links":{"citing_paper":"/paper/2511.19704"},"observation_digest":"sha256:3c577d93008bb0952c661119128e2e3a577e89650f2c0b18dc0daa047d6e5520","observation_id":"c7734f8e-30c4-439b-ac27-db99c6dcdd09","resolution":{"observed_at":"2026-05-17T05:24:05.339712Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2509.23563","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-29T21:53:59.372372Z","title":"RA VEN: Resilient Aerial Navigation via Open-Set Semantic Memory and Behavior Adaptation","venue":null,"work_id":"49229d16-ccca-4e68-92c2-f3dadfa2cfca","year":2025},"citing_paper":{"arxiv_id":"2511.19704","last_updated":"2026-04-10T00:50:03Z","snapshot_observed_at":"2026-08-11T14:26:55.408955Z","submitted_at":"2025-11-24T21:15:01Z","title":"RADSeg: Unleashing Parameter and Compute Efficient Zero-Shot Open-Vocabulary Segmentation Using Agglomerative Models","version":2},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-05-17T05:23:11.261860Z"},"links":{"citing_paper":"/paper/2511.19704"},"observation_digest":"sha256:d528eedefc363da64f181041a3067f2de11d379be15c5574ee1c0f4221649632","observation_id":"06655dac-e7b6-4f8e-b55d-2cc7dbcddb77","resolution":{"observed_at":"2026-05-17T05:24:04.736435Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Segment any- thing","venue":null,"work_id":"c6080054-4644-4c79-bbdd-751fa6bd3433","year":2023},"citing_paper":{"arxiv_id":"2511.19704","last_updated":"2026-04-10T00:50:03Z","snapshot_observed_at":"2026-08-11T14:26:55.408955Z","submitted_at":"2025-11-24T21:15:01Z","title":"RADSeg: Unleashing Parameter and Compute Efficient Zero-Shot Open-Vocabulary Segmentation Using Agglomerative Models","version":2},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-05-17T05:23:11.261860Z"},"links":{"citing_paper":"/paper/2511.19704"},"observation_digest":"sha256:e536b2630a46f883219bf71b050d268f656abca6c7d6c3540ab7c0681cb87170","observation_id":"8ed39f26-0b4b-49ca-82e8-2144408caafc","resolution":{"observed_at":"2026-05-17T05:24:05.311523Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Proxyclip: Proxy attention improves clip for open-vocabulary segmentation","venue":null,"work_id":"50e80f19-f58f-479e-94be-107fb825a00e","year":2024},"citing_paper":{"arxiv_id":"2511.19704","last_updated":"2026-04-10T00:50:03Z","snapshot_observed_at":"2026-08-11T14:26:55.408955Z","submitted_at":"2025-11-24T21:15:01Z","title":"RADSeg: Unleashing Parameter and Compute Efficient Zero-Shot Open-Vocabulary Segmentation Using Agglomerative Models","version":2},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-05-17T05:23:11.261860Z"},"links":{"citing_paper":"/paper/2511.19704"},"observation_digest":"sha256:bb742e7d34ff97b622f4e177cf79d6e714adc9afb7d50b69b3d32145bbd8522e","observation_id":"d3ca958b-be3a-4368-b49c-704114f3aac6","resolution":{"observed_at":"2026-05-17T05:24:05.283268Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2201.03546","last_updated":"2022-04-03T03:33:43Z","snapshot_observed_at":"2026-08-01T06:14:31.660540Z","submitted_at":"2022-01-10T18:59:10Z","title":"Language-driven Semantic Segmentation","version":2},"cited_work":{"arxiv_id":"2201.03546","doi":null,"metadata_source":"pith","pith_arxiv_id":"2201.03546","snapshot_observed_at":"2026-07-07T21:34:09.314172Z","title":"Language-driven Semantic Segmentation","venue":"cs.CV","work_id":"20f0fe18-203a-461e-b860-9b93636dd77b","year":2022},"citing_paper":{"arxiv_id":"2511.19704","last_updated":"2026-04-10T00:50:03Z","snapshot_observed_at":"2026-08-11T14:26:55.408955Z","submitted_at":"2025-11-24T21:15:01Z","title":"RADSeg: Unleashing Parameter and Compute Efficient Zero-Shot Open-Vocabulary Segmentation Using Agglomerative Models","version":2},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-05-17T05:23:11.261860Z"},"links":{"cited_paper":"/paper/2201.03546","citing_paper":"/paper/2511.19704"},"observation_digest":"sha256:20a75c38e8d7e1adbef68c233984d31dc5c26c03701ce94a5e76cdc190ef157a","observation_id":"41138eb8-5730-49b2-b608-1991b4bfbc29","resolution":{"observed_at":"2026-05-17T05:24:04.731932Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.06756","last_updated":"2025-03-16T10:12:23Z","snapshot_observed_at":"2026-08-09T04:04:20.740573Z","submitted_at":"2025-02-10T18:33:15Z","title":"SAMRefiner: Taming Segment Anything Model for Universal Mask Refinement","version":2},"cited_work":{"arxiv_id":"2502.06756","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2502.06756","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Samrefiner: Taming segment anything model for universal mask refinement","venue":null,"work_id":"cf55e49e-8630-4de0-ac7c-0d785dcf71ae","year":2025},"citing_paper":{"arxiv_id":"2511.19704","last_updated":"2026-04-10T00:50:03Z","snapshot_observed_at":"2026-08-11T14:26:55.408955Z","submitted_at":"2025-11-24T21:15:01Z","title":"RADSeg: Unleashing Parameter and Compute Efficient Zero-Shot Open-Vocabulary Segmentation Using Agglomerative Models","version":2},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-05-17T05:23:11.261860Z"},"links":{"cited_paper":"/paper/2502.06756","citing_paper":"/paper/2511.19704"},"observation_digest":"sha256:7e27cd83390b841c0b5561c7ebedcda30684e597a8c54668c3379ad11b73168a","observation_id":"14c782f9-22d9-4c13-b2eb-08019d702439","resolution":{"observed_at":"2026-05-17T05:24:04.741146Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2303.05499","last_updated":"2024-07-19T06:00:41Z","snapshot_observed_at":"2026-07-06T15:00:58.804337Z","submitted_at":"2023-03-09T18:52:16Z","title":"Grounding DINO: Marrying DINO with Grounded Pre-Training for Open-Set Object Detection","version":5},"cited_work":{"arxiv_id":"2303.05499","doi":null,"metadata_source":"pith","pith_arxiv_id":"2303.05499","snapshot_observed_at":"2026-07-10T23:17:45.410080Z","title":"Grounding DINO: Marrying DINO with Grounded Pre-Training for Open-Set Object Detection","venue":"cs.CV","work_id":"3757dc8f-79d5-4beb-a03b-eb4c9a33427d","year":2023},"citing_paper":{"arxiv_id":"2511.19704","last_updated":"2026-04-10T00:50:03Z","snapshot_observed_at":"2026-08-11T14:26:55.408955Z","submitted_at":"2025-11-24T21:15:01Z","title":"RADSeg: Unleashing Parameter and Compute Efficient Zero-Shot Open-Vocabulary Segmentation Using Agglomerative Models","version":2},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-05-17T05:23:11.261860Z"},"links":{"cited_paper":"/paper/2303.05499","citing_paper":"/paper/2511.19704"},"observation_digest":"sha256:bda0d4f6beae3778682e0dbb64f95a872ad6206350039e832d8cbcfd9ac8dfd4","observation_id":"d9a88dca-b886-48fc-8cd8-85015092622c","resolution":{"observed_at":"2026-05-17T05:24:04.726651Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.05266","last_updated":"2025-06-24T05:12:24Z","snapshot_observed_at":"2026-08-12T22:28:54.701173Z","submitted_at":"2024-10-07T17:59:45Z","title":"Brain Mapping with Dense Features: Grounding Cortical Semantic Selectivity in Natural Images With Vision Transformers","version":2},"cited_work":{"arxiv_id":"2410.05266","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2410.05266","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Brain mapping with dense features: Grounding cortical semantic selectivity in natural images with vision transformers","venue":null,"work_id":"0eb43f6e-c361-4d6b-a7ef-2187a78e7813","year":2024},"citing_paper":{"arxiv_id":"2511.19704","last_updated":"2026-04-10T00:50:03Z","snapshot_observed_at":"2026-08-11T14:26:55.408955Z","submitted_at":"2025-11-24T21:15:01Z","title":"RADSeg: Unleashing Parameter and Compute Efficient Zero-Shot Open-Vocabulary Segmentation Using Agglomerative Models","version":2},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-05-17T05:23:11.261860Z"},"links":{"cited_paper":"/paper/2410.05266","citing_paper":"/paper/2511.19704"},"observation_digest":"sha256:a0ff2e9766602a46742e2b15404ca079bc92bf83f46373d24d98c1de65152792","observation_id":"4b68dd6a-d1c0-426e-8a29-66c246383d40","resolution":{"observed_at":"2026-05-17T05:24:04.698400Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"The role of context for object detection and semantic segmentation in the wild","venue":null,"work_id":"449d65f6-dbb3-4e3c-86cc-2899b36024b8","year":2014},"citing_paper":{"arxiv_id":"2511.19704","last_updated":"2026-04-10T00:50:03Z","snapshot_observed_at":"2026-08-11T14:26:55.408955Z","submitted_at":"2025-11-24T21:15:01Z","title":"RADSeg: Unleashing Parameter and Compute Efficient Zero-Shot Open-Vocabulary Segmentation Using Agglomerative Models","version":2},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-05-17T05:23:11.261860Z"},"links":{"citing_paper":"/paper/2511.19704"},"observation_digest":"sha256:e4b9c54612d0b7840a1e3d0aed463f93c176ffbf486b364f09a62bf29ec84785","observation_id":"b3a51852-d01a-4cbc-95bf-dd610471772f","resolution":{"observed_at":"2026-05-17T05:24:05.318724Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"24b932a5-9825-45e0-ab7e-6130be44f105","year":2023},"citing_paper":{"arxiv_id":"2511.19704","last_updated":"2026-04-10T00:50:03Z","snapshot_observed_at":"2026-08-11T14:26:55.408955Z","submitted_at":"2025-11-24T21:15:01Z","title":"RADSeg: Unleashing Parameter and Compute Efficient Zero-Shot Open-Vocabulary Segmentation Using Agglomerative Models","version":2},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-05-17T05:23:11.261860Z"},"links":{"citing_paper":"/paper/2511.19704"},"observation_digest":"sha256:0094dad427342245d8f7d708da946ec6fdcbbb07b9b0ef5a7880564b14b9edb1","observation_id":"4e240f31-c5af-4463-8833-784446a640e4","resolution":{"observed_at":"2026-05-17T05:24:05.337099Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Learning transferable visual models from natural language supervi- sion","venue":null,"work_id":"d3d19f4f-8d2b-414d-b18a-38944fdd36ca","year":2021},"citing_paper":{"arxiv_id":"2511.19704","last_updated":"2026-04-10T00:50:03Z","snapshot_observed_at":"2026-08-11T14:26:55.408955Z","submitted_at":"2025-11-24T21:15:01Z","title":"RADSeg: Unleashing Parameter and Compute Efficient Zero-Shot Open-Vocabulary Segmentation Using Agglomerative Models","version":2},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-05-17T05:23:11.261860Z"},"links":{"citing_paper":"/paper/2511.19704"},"observation_digest":"sha256:fe2255f1de56d1da3f1358f694173164166eb4fe8db2aab3d9d1abdfb3d61ac8","observation_id":"d0c7cca4-6e21-4363-9878-1f83efa33a18","resolution":{"observed_at":"2026-05-17T05:24:05.334707Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Am-radio: Agglomerative vision foundation model reduce all domains into one","venue":null,"work_id":"6af56de3-7b31-4d93-b97d-936bdb6d3393","year":2024},"citing_paper":{"arxiv_id":"2511.19704","last_updated":"2026-04-10T00:50:03Z","snapshot_observed_at":"2026-08-11T14:26:55.408955Z","submitted_at":"2025-11-24T21:15:01Z","title":"RADSeg: Unleashing Parameter and Compute Efficient Zero-Shot Open-Vocabulary Segmentation Using Agglomerative Models","version":2},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-05-17T05:23:11.261860Z"},"links":{"citing_paper":"/paper/2511.19704"},"observation_digest":"sha256:89c92e12dce565cd070917f4c33ad7acb3445df749e9f2148eb41df1c20b05d1","observation_id":"89a4b9dc-3944-4d32-b42b-41a70a8e426a","resolution":{"observed_at":"2026-05-17T05:24:05.306762Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Denseclip: Language-guided dense prediction with context- aware prompting","venue":null,"work_id":"f3e0c77a-0574-4d7c-ac88-f58e96b9cc28","year":2022},"citing_paper":{"arxiv_id":"2511.19704","last_updated":"2026-04-10T00:50:03Z","snapshot_observed_at":"2026-08-11T14:26:55.408955Z","submitted_at":"2025-11-24T21:15:01Z","title":"RADSeg: Unleashing Parameter and Compute Efficient Zero-Shot Open-Vocabulary Segmentation Using Agglomerative Models","version":2},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-05-17T05:23:11.261860Z"},"links":{"citing_paper":"/paper/2511.19704"},"observation_digest":"sha256:69fa63ffbfb8afa2600693abbfcfb85ad8a182e7ff285f9bc1eec726d2838958","observation_id":"40dc5aeb-659c-46d1-891c-61e7d2dbaee2","resolution":{"observed_at":"2026-05-17T05:24:05.304232Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Language embedded radiance fields for zero-shot task-oriented grasping","venue":null,"work_id":"ced1dadc-1524-4103-a0e7-61ca6ab913d6","year":2023},"citing_paper":{"arxiv_id":"2511.19704","last_updated":"2026-04-10T00:50:03Z","snapshot_observed_at":"2026-08-11T14:26:55.408955Z","submitted_at":"2025-11-24T21:15:01Z","title":"RADSeg: Unleashing Parameter and Compute Efficient Zero-Shot Open-Vocabulary Segmentation Using Agglomerative Models","version":2},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-05-17T05:23:11.261860Z"},"links":{"citing_paper":"/paper/2511.19704"},"observation_digest":"sha256:87e452bc3f9db1d6013dacd70fb52de5414dd00491405cd506048562848666e8","observation_id":"5f794781-2899-44ca-97c5-3683a53ad251","resolution":{"observed_at":"2026-05-17T05:24:05.296560Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2408.00714","last_updated":"2024-10-28T16:37:57Z","snapshot_observed_at":"2026-07-06T18:55:41.459417Z","submitted_at":"2024-08-01T17:00:08Z","title":"SAM 2: Segment Anything in Images and Videos","version":2},"cited_work":{"arxiv_id":"2408.00714","doi":"10.1038/s41598-025-97590-3","metadata_source":"pith","pith_arxiv_id":"2408.00714","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 2: Segment Anything in Images and Videos","venue":"cs.CV","work_id":"acc13f66-d814-44f9-9688-375688bf2d4a","year":2024},"citing_paper":{"arxiv_id":"2511.19704","last_updated":"2026-04-10T00:50:03Z","snapshot_observed_at":"2026-08-11T14:26:55.408955Z","submitted_at":"2025-11-24T21:15:01Z","title":"RADSeg: Unleashing Parameter and Compute Efficient Zero-Shot Open-Vocabulary Segmentation Using Agglomerative Models","version":2},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-05-17T05:23:11.261860Z"},"links":{"cited_paper":"/paper/2408.00714","citing_paper":"/paper/2511.19704"},"observation_digest":"sha256:a1610e0473b371f5cfc6c25085dc4380983c24e27b45468354e58f0da092b56c","observation_id":"0ba69b5a-5f04-49d8-aa7a-f183f5e819c6","resolution":{"observed_at":"2026-05-17T05:24:04.717739Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-05-24T04:24:23.885301+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-24T04:24:23.885301+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Laion-5b: An open large-scale dataset for training next generation image-text models.Advances in neural in- formation processing systems, 35:25278–25294","venue":null,"work_id":"7188568c-a234-4fc3-9abb-ccf52299b0a1","year":2022},"citing_paper":{"arxiv_id":"2511.19704","last_updated":"2026-04-10T00:50:03Z","snapshot_observed_at":"2026-08-11T14:26:55.408955Z","submitted_at":"2025-11-24T21:15:01Z","title":"RADSeg: Unleashing Parameter and Compute Efficient Zero-Shot Open-Vocabulary Segmentation Using Agglomerative Models","version":2},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-05-17T05:23:11.261860Z"},"links":{"citing_paper":"/paper/2511.19704"},"observation_digest":"sha256:dbc1dfe9ffa8a585a448c9d3d7dbcf5d06fb88aed7fb076e96d1f9c1a04eb501","observation_id":"5025fdad-d7a2-4213-8f29-30a1425da114","resolution":{"observed_at":"2026-05-17T05:24:05.268860Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2210.05663","last_updated":"2023-05-22T22:34:29Z","snapshot_observed_at":"2026-08-13T14:08:05.014142Z","submitted_at":"2022-10-11T17:57:10Z","title":"CLIP-Fields: Weakly Supervised Semantic Fields for Robotic Memory","version":3},"cited_work":{"arxiv_id":"2210.05663","doi":null,"metadata_source":"pith","pith_arxiv_id":"2210.05663","snapshot_observed_at":"2026-07-10T11:47:02.958622Z","title":"arXiv preprint arXiv:2210.05663 (2022)","venue":"cs.RO","work_id":"54a0bbf5-4695-4324-97fd-157d052ec297","year":2022},"citing_paper":{"arxiv_id":"2511.19704","last_updated":"2026-04-10T00:50:03Z","snapshot_observed_at":"2026-08-11T14:26:55.408955Z","submitted_at":"2025-11-24T21:15:01Z","title":"RADSeg: Unleashing Parameter and Compute Efficient Zero-Shot Open-Vocabulary Segmentation Using Agglomerative Models","version":2},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-05-17T05:23:11.261860Z"},"links":{"cited_paper":"/paper/2210.05663","citing_paper":"/paper/2511.19704"},"observation_digest":"sha256:27f3efb6281f94bf8f44f2145bd9a2073b431b5422dfd096390ebfb54c089f88","observation_id":"9ec4e3d1-e614-4a74-b2b0-b4c5f954dfb8","resolution":{"observed_at":"2026-05-17T05:24:04.745930Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2407.20179","last_updated":"2024-10-10T17:27:46Z","snapshot_observed_at":"2026-08-12T23:12:40.143372Z","submitted_at":"2024-07-29T17:08:21Z","title":"Theia: Distilling Diverse Vision Foundation Models for Robot Learning","version":2},"cited_work":{"arxiv_id":"2407.20179","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2407.20179","snapshot_observed_at":"2026-07-04T03:09:29.580056Z","title":"Theia: Distilling diverse vision foundation models for robot learning","venue":null,"work_id":"ee430a90-3ded-47db-bdf1-bbc20527e97a","year":2024},"citing_paper":{"arxiv_id":"2511.19704","last_updated":"2026-04-10T00:50:03Z","snapshot_observed_at":"2026-08-11T14:26:55.408955Z","submitted_at":"2025-11-24T21:15:01Z","title":"RADSeg: Unleashing Parameter and Compute Efficient Zero-Shot Open-Vocabulary Segmentation Using Agglomerative Models","version":2},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-05-17T05:23:11.261860Z"},"links":{"cited_paper":"/paper/2407.20179","citing_paper":"/paper/2511.19704"},"observation_digest":"sha256:c17fef70db02c786bae9ed800febb5e4825dfa5870a4fc01e81b8651ad019786","observation_id":"890b4ed2-661c-42f5-88be-4dcaf068aabd","resolution":{"observed_at":"2026-05-17T05:24:04.765662Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Language embedded 3d gaussians for open- vocabulary scene understanding","venue":null,"work_id":"96b72511-5ddd-4c75-9f54-44f18aa3b21e","year":2024},"citing_paper":{"arxiv_id":"2511.19704","last_updated":"2026-04-10T00:50:03Z","snapshot_observed_at":"2026-08-11T14:26:55.408955Z","submitted_at":"2025-11-24T21:15:01Z","title":"RADSeg: Unleashing Parameter and Compute Efficient Zero-Shot Open-Vocabulary Segmentation Using Agglomerative Models","version":2},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-05-17T05:23:11.261860Z"},"links":{"citing_paper":"/paper/2511.19704"},"observation_digest":"sha256:e6a88fa5afde7c2bc3a4f730b4a23d41bdb7aeead32e8c54da87279c954c98b0","observation_id":"143d6142-55a7-4cb3-ab23-8352cfc86f99","resolution":{"observed_at":"2026-05-17T05:24:05.293969Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2411.09219","last_updated":"2024-11-14T06:31:20Z","snapshot_observed_at":"2026-08-12T20:51:32.088773Z","submitted_at":"2024-11-14T06:31:20Z","title":"Harnessing Vision Foundation Models for High-Performance, Training-Free Open Vocabulary Segmentation","version":1},"cited_work":{"arxiv_id":"2411.09219","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2411.09219","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Har- nessing vision foundation models for high-performance, training-free open vocabulary segmentation","venue":null,"work_id":"ee2e65e1-3d74-45f2-ae9e-308411c03d93","year":2024},"citing_paper":{"arxiv_id":"2511.19704","last_updated":"2026-04-10T00:50:03Z","snapshot_observed_at":"2026-08-11T14:26:55.408955Z","submitted_at":"2025-11-24T21:15:01Z","title":"RADSeg: Unleashing Parameter and Compute Efficient Zero-Shot Open-Vocabulary Segmentation Using Agglomerative Models","version":2},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-05-17T05:23:11.261860Z"},"links":{"cited_paper":"/paper/2411.09219","citing_paper":"/paper/2511.19704"},"observation_digest":"sha256:b73986cffc9f4a4c450da51d528bdb700677c6990c48b15860fd0b20c5e82a2c","observation_id":"cba7ff5d-3d7b-4b5e-afee-f79262bb3edd","resolution":{"observed_at":"2026-05-17T05:24:04.703470Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1906.05797","last_updated":"2019-06-13T16:29:58Z","snapshot_observed_at":"2026-08-01T13:51:16.469557Z","submitted_at":"2019-06-13T16:29:58Z","title":"The Replica Dataset: A Digital Replica of Indoor Spaces","version":1},"cited_work":{"arxiv_id":"1906.05797","doi":"10.48550/arxiv.1906.05797","metadata_source":"pith","pith_arxiv_id":"1906.05797","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"The Replica Dataset: A Digital Replica of Indoor Spaces","venue":"cs.CV","work_id":"8145b3bf-a202-46d8-a3bc-fad366dd5d4a","year":2019},"citing_paper":{"arxiv_id":"2511.19704","last_updated":"2026-04-10T00:50:03Z","snapshot_observed_at":"2026-08-11T14:26:55.408955Z","submitted_at":"2025-11-24T21:15:01Z","title":"RADSeg: Unleashing Parameter and Compute Efficient Zero-Shot Open-Vocabulary Segmentation Using Agglomerative Models","version":2},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-05-17T05:23:11.261860Z"},"links":{"cited_paper":"/paper/1906.05797","citing_paper":"/paper/2511.19704"},"observation_digest":"sha256:7e0aaa8580ddfaa40cbb26df3f219f4f806bd60625bdba68db133d0af5df28f1","observation_id":"2982d568-7b07-46b6-89a8-bdb99c07c794","resolution":{"observed_at":"2026-05-17T05:24:04.712930Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-05-23T17:23:43.390569+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-23T17:23:43.390569+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.14786","last_updated":"2025-02-20T18:08:29Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-02-20T18:08:29Z","title":"SigLIP 2: Multilingual Vision-Language Encoders with Improved Semantic Understanding, Localization, and Dense Features","version":1},"cited_work":{"arxiv_id":"2502.14786","doi":"10.48550/arxiv.2502.14786","metadata_source":"pith","pith_arxiv_id":"2502.14786","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SigLIP 2: Multilingual Vision-Language Encoders with Improved Semantic Understanding, Localization, and Dense Features","venue":"cs.CV","work_id":"50eec732-2d41-432f-9dcf-ac7fff235ea5","year":2025},"citing_paper":{"arxiv_id":"2511.19704","last_updated":"2026-04-10T00:50:03Z","snapshot_observed_at":"2026-08-11T14:26:55.408955Z","submitted_at":"2025-11-24T21:15:01Z","title":"RADSeg: Unleashing Parameter and Compute Efficient Zero-Shot Open-Vocabulary Segmentation Using Agglomerative Models","version":2},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-05-17T05:23:11.261860Z"},"links":{"cited_paper":"/paper/2502.14786","citing_paper":"/paper/2511.19704"},"observation_digest":"sha256:a080865b2463e882b823f29ed57963f7b81ad4cc888c01e34794fdb6fdea1cb3","observation_id":"c60b05bd-3436-4c82-9eec-c4e4414562b4","resolution":{"observed_at":"2026-05-17T05:24:04.692945Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-07-10T23:49:08.777694+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-10T23:49:08.777694+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Sclip: Rethink- ing self-attention for dense vision-language inference","venue":null,"work_id":"70df32e7-08a8-4961-a121-f8f93e826fde","year":2024},"citing_paper":{"arxiv_id":"2511.19704","last_updated":"2026-04-10T00:50:03Z","snapshot_observed_at":"2026-08-11T14:26:55.408955Z","submitted_at":"2025-11-24T21:15:01Z","title":"RADSeg: Unleashing Parameter and Compute Efficient Zero-Shot Open-Vocabulary Segmentation Using Agglomerative Models","version":2},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-05-17T05:23:11.261860Z"},"links":{"citing_paper":"/paper/2511.19704"},"observation_digest":"sha256:b96641c3e1ed4bd34e6288b8f7255bc157c2750acb11652007be9ffd28dd6672","observation_id":"bf2aa05b-62c7-4258-ac5d-6db5f0ff5ea2","resolution":{"observed_at":"2026-05-17T05:24:05.302008Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Hierarchical open- vocabulary 3d scene graphs for language-grounded robot navigation","venue":null,"work_id":"9b3479df-74d5-4621-b5d7-8d09e532287c","year":2024},"citing_paper":{"arxiv_id":"2511.19704","last_updated":"2026-04-10T00:50:03Z","snapshot_observed_at":"2026-08-11T14:26:55.408955Z","submitted_at":"2025-11-24T21:15:01Z","title":"RADSeg: Unleashing Parameter and Compute Efficient Zero-Shot Open-Vocabulary Segmentation Using Agglomerative Models","version":2},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-05-17T05:23:11.261860Z"},"links":{"citing_paper":"/paper/2511.19704"},"observation_digest":"sha256:4945970ac54bbb537ae6aadded8c0758ffaba21ba4c17a1f2fe5bbd95d5e4837","observation_id":"c02a8007-0fcb-4ef0-8103-f9bfdeea6bfe","resolution":{"observed_at":"2026-05-17T05:24:05.309207Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Clip-diy: Clip dense infer- ence yields open-vocabulary semantic segmentation for-free","venue":null,"work_id":"3565324c-7d43-4848-aadf-4e1fc26908f5","year":2024},"citing_paper":{"arxiv_id":"2511.19704","last_updated":"2026-04-10T00:50:03Z","snapshot_observed_at":"2026-08-11T14:26:55.408955Z","submitted_at":"2025-11-24T21:15:01Z","title":"RADSeg: Unleashing Parameter and Compute Efficient Zero-Shot Open-Vocabulary Segmentation Using Agglomerative Models","version":2},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-05-17T05:23:11.261860Z"},"links":{"citing_paper":"/paper/2511.19704"},"observation_digest":"sha256:928df3372814e5acfe4563c66181a441dbb09837402659d63fa07c4201426ce8","observation_id":"0d53da3e-309b-48b9-bfba-872d955f547a","resolution":{"observed_at":"2026-05-17T05:24:05.313942Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2505.23769","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Textregion: Text-aligned region tokens from frozen image-text models.arXiv preprint arXiv:2505.23769","venue":null,"work_id":"6d24d983-d361-40f4-9006-7baabefb7ed5","year":null},"citing_paper":{"arxiv_id":"2511.19704","last_updated":"2026-04-10T00:50:03Z","snapshot_observed_at":"2026-08-11T14:26:55.408955Z","submitted_at":"2025-11-24T21:15:01Z","title":"RADSeg: Unleashing Parameter and Compute Efficient Zero-Shot Open-Vocabulary Segmentation Using Agglomerative Models","version":2},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-05-17T05:23:11.261860Z"},"links":{"citing_paper":"/paper/2511.19704"},"observation_digest":"sha256:7c8c37b63bc6c824691e36ebca2d8345080a29c6e73a2f282598b8d7189a0a48","observation_id":"ba0ee568-9fac-4c99-a0de-4ec9583008e9","resolution":{"observed_at":"2026-05-17T05:24:04.722394Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Side adapter network for open-vocabulary semantic segmentation","venue":null,"work_id":"82701329-cc1d-4c67-a185-1662f48446e7","year":2023},"citing_paper":{"arxiv_id":"2511.19704","last_updated":"2026-04-10T00:50:03Z","snapshot_observed_at":"2026-08-11T14:26:55.408955Z","submitted_at":"2025-11-24T21:15:01Z","title":"RADSeg: Unleashing Parameter and Compute Efficient Zero-Shot Open-Vocabulary Segmentation Using Agglomerative Models","version":2},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-05-17T05:23:11.261860Z"},"links":{"citing_paper":"/paper/2511.19704"},"observation_digest":"sha256:819dbd0f5dc707036ebd56eef7f5b1267a952c6e4c6640acc7af4e3dcba1dc94","observation_id":"9ebdc7a1-2d13-4255-9513-e20bb4ab1b0f","resolution":{"observed_at":"2026-05-17T05:24:05.271249Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Resclip: Residual attention for training-free dense vision- language inference","venue":null,"work_id":"953fbc9f-7685-4c0c-980b-e00d4950a518","year":null},"citing_paper":{"arxiv_id":"2511.19704","last_updated":"2026-04-10T00:50:03Z","snapshot_observed_at":"2026-08-11T14:26:55.408955Z","submitted_at":"2025-11-24T21:15:01Z","title":"RADSeg: Unleashing Parameter and Compute Efficient Zero-Shot Open-Vocabulary Segmentation Using Agglomerative Models","version":2},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-05-17T05:23:11.261860Z"},"links":{"citing_paper":"/paper/2511.19704"},"observation_digest":"sha256:5b9452327fdbc62acb4523870a44f218dd0174304e92f76f99c1dcf5e1164a45","observation_id":"62f57de1-b752-4a6c-8723-ead103737c96","resolution":{"observed_at":"2026-05-17T05:24:05.291279Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Scannet++: A high-fidelity dataset of 3d in- door scenes","venue":null,"work_id":"83bba9e5-acc3-4027-8c22-c3807c240d90","year":2023},"citing_paper":{"arxiv_id":"2511.19704","last_updated":"2026-04-10T00:50:03Z","snapshot_observed_at":"2026-08-11T14:26:55.408955Z","submitted_at":"2025-11-24T21:15:01Z","title":"RADSeg: Unleashing Parameter and Compute Efficient Zero-Shot Open-Vocabulary Segmentation Using Agglomerative Models","version":2},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-05-17T05:23:11.261860Z"},"links":{"citing_paper":"/paper/2511.19704"},"observation_digest":"sha256:9fdcdeee3ec4128646e3ed87ee62bdc2a51a3e7afa0309a247c60f8ebf8edc4a","observation_id":"ddacc47a-119b-4796-b331-cd3512c6605b","resolution":{"observed_at":"2026-05-17T05:24:05.299423Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Vlfm: Vision-language frontier maps for zero-shot semantic navigation","venue":null,"work_id":"5a50905e-d2ba-45b3-9a92-b97d6060fa81","year":2024},"citing_paper":{"arxiv_id":"2511.19704","last_updated":"2026-04-10T00:50:03Z","snapshot_observed_at":"2026-08-11T14:26:55.408955Z","submitted_at":"2025-11-24T21:15:01Z","title":"RADSeg: Unleashing Parameter and Compute Efficient Zero-Shot Open-Vocabulary Segmentation Using Agglomerative Models","version":2},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-05-17T05:23:11.261860Z"},"links":{"citing_paper":"/paper/2511.19704"},"observation_digest":"sha256:0d0e6d552f55e0b492eb78f9c876b31cc0353b6727e54cdc408c7af0c3f30770","observation_id":"2941d9f3-a2e7-4db3-8cc9-5211d8693172","resolution":{"observed_at":"2026-05-17T05:24:05.281069Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Sigmoid loss for language image pre-training","venue":null,"work_id":"c6e7d01c-735a-474e-a036-a3e8e9764e34","year":2023},"citing_paper":{"arxiv_id":"2511.19704","last_updated":"2026-04-10T00:50:03Z","snapshot_observed_at":"2026-08-11T14:26:55.408955Z","submitted_at":"2025-11-24T21:15:01Z","title":"RADSeg: Unleashing Parameter and Compute Efficient Zero-Shot Open-Vocabulary Segmentation Using Agglomerative Models","version":2},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-05-17T05:23:11.261860Z"},"links":{"citing_paper":"/paper/2511.19704"},"observation_digest":"sha256:72796ead1cf02201c705f2f2d1e5df23d1f55ac7f7c5bb73961decfd3dd290f5","observation_id":"4c53d405-67d5-42aa-93fb-1e3d7d95f371","resolution":{"observed_at":"2026-05-17T05:24:05.288416Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-06T00:31:42.395652Z","title":"Scene parsing through ade20k dataset","venue":null,"work_id":"ff369fcc-fa0b-4113-8574-0761cffc851c","year":null},"citing_paper":{"arxiv_id":"2511.19704","last_updated":"2026-04-10T00:50:03Z","snapshot_observed_at":"2026-08-11T14:26:55.408955Z","submitted_at":"2025-11-24T21:15:01Z","title":"RADSeg: Unleashing Parameter and Compute Efficient Zero-Shot Open-Vocabulary Segmentation Using Agglomerative Models","version":2},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-05-17T05:23:11.261860Z"},"links":{"citing_paper":"/paper/2511.19704"},"observation_digest":"sha256:83885a3ee0da67d217879ffb6d3c4778d2995a953e42451b675f9cdc10efa8e3","observation_id":"8ad2737d-e72b-4869-999f-a27514fd933d","resolution":{"observed_at":"2026-05-17T05:24:05.266306Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"RGB” and “GT","venue":null,"work_id":"1eeb44f0-4959-4980-985b-443072fd4ef3","year":2022},"citing_paper":{"arxiv_id":"2511.19704","last_updated":"2026-04-10T00:50:03Z","snapshot_observed_at":"2026-08-11T14:26:55.408955Z","submitted_at":"2025-11-24T21:15:01Z","title":"RADSeg: Unleashing Parameter and Compute Efficient Zero-Shot Open-Vocabulary Segmentation Using Agglomerative Models","version":2},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-05-17T05:23:11.261860Z"},"links":{"citing_paper":"/paper/2511.19704"},"observation_digest":"sha256:211238123394d2dfe1abc3a139c467a27380dd26de3d957b152f365ebb542db2","observation_id":"9b302991-e3ab-402f-8eb0-265700f1a5e7","resolution":{"observed_at":"2026-05-17T05:24:05.276434Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2511.19704","last_updated":"2026-04-10T00:50:03Z","latest_version":2,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-11T14:26:55.408955Z","submitted_at":"2025-11-24T21:15:01Z","title":"RADSeg: Unleashing Parameter and Compute Efficient Zero-Shot Open-Vocabulary Segmentation Using Agglomerative Models"},"reference_resolution":{"displayed":49,"state_counts":{"malformed_identifier":0,"metadata_mismatch":1,"parse_uncertain":0,"unresolved":2,"verified_exact":16,"verified_fuzzy":30},"total_outbound_references":49},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"thesis":"As of 13 August 2026, this Paper Citation Record lists 49 of 49 outbound references and 4 inbound Pith citation observations for arXiv:2511.19704."}