{"as_of":"2026-08-05T08:09:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:77c069f87fb2ee5fd06646311d785fa0995aace5935aa9df9bc83caa4921b47d","coverage":[{"denominator":168,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":100,"source":"paper_references, paper_reference_links","source_observed_at":"2026-05-17T20:22:46.220021Z","state":"measured"},{"denominator":200,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":200,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-05T06:32:48.257954+00:00","state":"measured"},{"denominator":480,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":100,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-03T21:06:03.793460Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":1,"source":"pith","source_observed_at":"2026-08-05T02:28:24.338817Z","state":"measured"}],"external_citation_measurements":[{"count":7,"observed_at":"2026-08-05T02:28:24.338817Z","source":"pith"}],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-03T21:06:03.793460Z","title":"Sam 3: Segment anything with concepts,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2511.16887","last_updated":"2026-07-22T07:50:27Z","snapshot_observed_at":"2026-08-03T21:05:54.812786Z","submitted_at":"2025-11-21T02:00:17Z","title":"Glass Surface Detection: Leveraging Reflection Dynamics in Flash/No-flash Imagery","version":5},"reference_index":75,"source":"pdf_text","source_observed_at":"2026-08-03T21:06:03.793460Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2511.16887"},"observation_digest":"sha256:01e4c9488851b168418970a632a16198dec19e2c26af78b8671c1a0415b0897f","observation_id":"c7d6612a-9866-4d67-aad0-d2e308e7b740","resolution":{"observed_at":"2026-08-03T21:06:03.793460Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2511.21926","last_updated":"2026-04-06T04:06:21Z","snapshot_observed_at":"2026-07-06T22:37:08.086095Z","submitted_at":"2025-11-26T21:36:58Z","title":"Comparing SAM 2 and SAM 3 for Zero-Shot Segmentation of 3D Medical Data","version":2},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-05-17T04:38:15.388826Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2511.21926"},"observation_digest":"sha256:085e820e265c12e85faf9cbf6a3b75f027b45b8fd07da6a9b1d60329eb119067","observation_id":"a2a1316d-48e4-4cb7-a584-3c9f1cef3bb1","resolution":{"observed_at":"2026-05-17T04:39:02.640472Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2512.04585","last_updated":"2026-04-16T07:12:40Z","snapshot_observed_at":"2026-07-06T22:37:45.947333Z","submitted_at":"2025-12-04T09:00:25Z","title":"SAM3-I: Segment Anything with Instructions","version":4},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-05-17T02:06:35.864921Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2512.04585"},"observation_digest":"sha256:02c7434269fce19fe2ddabdc3085f242a79b0118070e7250e0965d2ac509491b","observation_id":"78513a55-c14a-4302-95ec-dbfebe3777ca","resolution":{"observed_at":"2026-05-17T02:08:51.402283Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2512.06171","last_updated":"2026-04-23T10:42:33Z","snapshot_observed_at":"2026-07-06T22:37:54.987339Z","submitted_at":"2025-12-05T21:49:58Z","title":"Automated Annotation of Shearographic Measurements Enabling Weakly Supervised Defect Detection","version":2},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-05-17T00:19:31.834944Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2512.06171"},"observation_digest":"sha256:aad2adf39dfbf8139b3177012d465b9f1d00c756ec897c409402428064b27fb4","observation_id":"8915b35e-d61f-45d5-8a7d-37a833b1c4ee","resolution":{"observed_at":"2026-05-17T00:21:23.687542Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2512.08730","last_updated":"2026-04-22T15:53:23Z","snapshot_observed_at":"2026-07-06T22:38:13.359449Z","submitted_at":"2025-12-09T15:42:28Z","title":"SegEarth-OV3: Exploring SAM 3 for Open-Vocabulary Semantic Segmentation in Remote Sensing Images","version":2},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-05-16T23:53:27.350335Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2512.08730"},"observation_digest":"sha256:1fdad10eed174b1d4934912a282d3968016ab722bee57f4172d388df05ef2281","observation_id":"a0adf699-8f55-4589-85d0-623d4cbbecbf","resolution":{"observed_at":"2026-05-16T23:53:42.289506Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-03T14:08:20.175019Z","title":"Carion, L","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2512.21545","last_updated":"2026-06-27T07:08:38Z","snapshot_observed_at":"2026-08-05T00:52:22.372509Z","submitted_at":"2025-12-25T07:34:38Z","title":"EraseLoRA: MLLM-Driven Foreground Exclusion and Background Subtype Aggregation for Dataset-Free Object Removal","version":2},"reference_index":2,"source":"arxiv_source","source_observed_at":"2026-08-03T14:08:20.175019Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2512.21545"},"observation_digest":"sha256:684b002030ab67bc60467675d6ee24150940b6f494a7f05f2289676506290a26","observation_id":"d77353ef-1ffc-4652-80d9-0ec8b0696e77","resolution":{"observed_at":"2026-08-03T14:08:20.175019Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2601.07447","last_updated":"2026-04-24T07:40:21Z","snapshot_observed_at":"2026-07-06T22:41:29.385811Z","submitted_at":"2026-01-12T11:39:36Z","title":"PanoSAMic: Panoramic Image Segmentation from SAM Feature Encoding and Dual View Fusion","version":3},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-05-16T14:43:18.682997Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2601.07447"},"observation_digest":"sha256:302c4a3f4fdb941921462e3ae08fbc447354fbf0c856bdd98b2c059a53616fe9","observation_id":"72b453bb-ccd7-4704-a95c-eee46a382203","resolution":{"observed_at":"2026-05-16T14:48:00.841351Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2601.10611","last_updated":"2026-04-02T16:01:02Z","snapshot_observed_at":"2026-08-02T06:45:30.387180Z","submitted_at":"2026-01-15T17:27:44Z","title":"Molmo2: Open Weights and Data for Vision-Language Models with Video Understanding and Grounding","version":4},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-05-16T04:21:29.526008Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2601.10611"},"observation_digest":"sha256:46b163824a0ce4d0059a67c109b8be21215c4534eeaf246a1058cb403c5945f3","observation_id":"4965ad96-e67f-4501-bbf1-aceb171e428a","resolution":{"observed_at":"2026-05-16T04:21:29.693226Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2601.13895","last_updated":"2026-04-24T08:12:11Z","snapshot_observed_at":"2026-07-06T22:42:12.887517Z","submitted_at":"2026-01-20T12:25:41Z","title":"OmniOVCD: Streamlining Open-Vocabulary Change Detection with SAM 3","version":2},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-05-16T12:51:19.006683Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2601.13895"},"observation_digest":"sha256:c72b6c614d00b8783b2b125d7baddafaf99a02fe91ee85a2235c20033f2df839","observation_id":"5eb98aaf-4d8f-456b-aa38-e587bcff05f0","resolution":{"observed_at":"2026-05-16T12:52:53.290843Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-03T03:41:53.322999Z","title":"Sam 3: Segment anything with concepts","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2602.07413","last_updated":"2026-06-09T17:59:12Z","snapshot_observed_at":"2026-08-03T03:41:52.703042Z","submitted_at":"2026-02-07T07:18:00Z","title":"Going with the Flow: Koopman Behavioral Models as Pseudo Planners for Visuo-Motor Dexterity","version":3},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-08-03T03:41:53.322999Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2602.07413"},"observation_digest":"sha256:c366ecfe65c9e7e49a9b0e16011de5123d6442838d57be133465b01d173af464","observation_id":"ae6b334e-9778-49bb-9120-fdc5985a915f","resolution":{"observed_at":"2026-08-03T03:41:53.322999Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-02T23:23:33.793070Z","title":"V ., Khedr, H., Huang, A., et al","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2602.14134","last_updated":"2026-06-01T07:44:02Z","snapshot_observed_at":"2026-08-02T23:23:32.424074Z","submitted_at":"2026-02-15T13:12:28Z","title":"DenseMLLM: Standard Multimodal LLMs for Dense Prediction","version":2},"reference_index":2025,"source":"pdf_text","source_observed_at":"2026-08-02T23:23:33.793070Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2602.14134"},"observation_digest":"sha256:65cc916f2848d1c22daee1a634259925f145c4e24847772157a42779b7db60a2","observation_id":"081f0df1-d124-4cc9-abe0-b3032db15574","resolution":{"observed_at":"2026-08-02T23:23:33.793070Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-02T21:20:16.412077Z","title":"Sam 3: Segment anything with concepts,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2602.20551","last_updated":"2026-06-07T03:20:42Z","snapshot_observed_at":"2026-08-02T21:20:15.945931Z","submitted_at":"2026-02-24T05:10:22Z","title":"CAD-Prompted SAM3: Geometry-Conditioned Instance Segmentation for Industrial Objects","version":3},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-02T21:20:16.412077Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2602.20551"},"observation_digest":"sha256:16b689e7dae6f5dca0335525306e982bdc3636fe8d82ae026f5319deed49d2d7","observation_id":"09ff3779-20df-43e0-b28e-45d8bef4a2f0","resolution":{"observed_at":"2026-08-02T21:20:16.412077Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2603.03577","last_updated":"2026-05-13T20:51:32Z","snapshot_observed_at":"2026-07-06T22:47:46.409064Z","submitted_at":"2026-03-03T23:11:17Z","title":"From Local Matches to Global Masks: Template-Guided Instance Detection and Segmentation in Open-World Scenes","version":2},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-05-15T16:19:15.019645Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2603.03577"},"observation_digest":"sha256:11c7f9a5a6c5b3e1281decc12c6f751d32cfa8fae240a6126fb7057eabf09864","observation_id":"9feaafd4-f012-4749-8bf5-90300f24e668","resolution":{"observed_at":"2026-05-15T16:20:09.762625Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2603.05377","last_updated":"2026-06-28T10:51:16Z","snapshot_observed_at":"2026-08-02T22:18:25.610175Z","submitted_at":"2026-03-05T17:02:22Z","title":"OpenFrontier: General Navigation with Visual-Language Grounded Frontiers","version":2},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-05-21T11:36:24.334860Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2603.05377"},"observation_digest":"sha256:8321d3bfcf2211e26200422fdcda1d66a43fc149ded69ee3871f6e630b1b1a66","observation_id":"db76e5a5-b39d-41d9-a1d3-6d8f21d6cbf8","resolution":{"observed_at":"2026-05-21T11:40:03.479992Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-07-15T13:59:22.441714Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2603.06168","last_updated":"2026-06-27T17:44:17Z","snapshot_observed_at":"2026-08-02T23:39:17.774806Z","submitted_at":"2026-03-06T11:22:14Z","title":"JOPP-3D: Joint Open Vocabulary Semantic Segmentation on Point Clouds and Panoramas","version":3},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-07-15T13:59:22.441714Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2603.06168"},"observation_digest":"sha256:20e3e4b686e0ed8319d3b5289f1384dd7d4fa78c7a667684afcfe9cb9edfbea4","observation_id":"5c6c7158-acad-4860-b5b0-45c668acc6f1","resolution":{"observed_at":"2026-07-15T13:59:22.441714Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2603.06885","last_updated":"2026-04-15T13:11:59Z","snapshot_observed_at":"2026-07-06T22:48:13.053749Z","submitted_at":"2026-03-06T21:07:08Z","title":"OPTED: Open Preprocessed Trachoma Eye Dataset Using Zero-Shot SAM 3 Segmentation","version":2},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-05-15T14:33:41.517193Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2603.06885"},"observation_digest":"sha256:c09d25a509aee164a6e767881890593698327d933c0f39059dccf7c46362f50a","observation_id":"126b5957-dd78-473d-b437-8dcdd122f9b2","resolution":{"observed_at":"2026-05-15T14:35:55.821273Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-07-15T13:24:52.801459Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2603.07221","last_updated":"2026-06-24T11:18:49Z","snapshot_observed_at":"2026-07-31T01:15:29.359580Z","submitted_at":"2026-03-07T14:02:39Z","title":"Margin in Abstract Spaces","version":2},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-07-15T13:24:52.801459Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2603.07221"},"observation_digest":"sha256:fed772ea9a56615a7d418bcccce7051123181d88865d667a423bb45b92d31845","observation_id":"94da9b96-4aa6-4a33-8b4f-6a230249df7d","resolution":{"observed_at":"2026-07-15T13:24:52.801459Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2603.08096","last_updated":"2026-04-20T17:13:59Z","snapshot_observed_at":"2026-07-06T22:48:21.803560Z","submitted_at":"2026-03-09T08:37:05Z","title":"TrianguLang: Geometry-Aware Semantic Consensus for Pose-Free 3D Localization","version":3},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-05-15T15:12:23.459579Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2603.08096"},"observation_digest":"sha256:4620e9f6d3c5597f5f95d550127ae7ac36fe04417b03406b562a415a8f97f2a1","observation_id":"0b8d655f-7c74-4a8f-81f2-ffe1340474fc","resolution":{"observed_at":"2026-05-15T15:16:09.686640Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2603.09283","last_updated":"2026-04-22T02:23:50Z","snapshot_observed_at":"2026-07-06T22:48:26.306802Z","submitted_at":"2026-03-10T07:07:27Z","title":"From Ideal to Real: Stable Video Object Removal under Imperfect Conditions","version":2},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-05-15T13:34:45.782304Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2603.09283"},"observation_digest":"sha256:10579589a50fd1a7ec9ade32b0201234b0ac63378cb2b1fd3e4edce7e8ba1018","observation_id":"e49c95a7-8b16-4f41-bb02-ecc23d38fa61","resolution":{"observed_at":"2026-05-15T13:35:51.535952Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-07-14T22:32:36.129552Z","title":"Sam 3: Segment anything with concepts,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2603.11917","last_updated":"2026-06-19T06:31:29Z","snapshot_observed_at":"2026-07-14T22:32:35.882828Z","submitted_at":"2026-03-12T13:31:43Z","title":"PicoSAM3: Real-Time In-Sensor Region-of-Interest Segmentation","version":4},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-07-14T22:32:36.129552Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2603.11917"},"observation_digest":"sha256:3d94e2ccb781d040a36ea85fc9c70ed202e3ceb211d231b7fe112a7699475afb","observation_id":"8eaae0ae-bbad-4dd4-b1fa-0a2fb56006ca","resolution":{"observed_at":"2026-07-14T22:32:36.129552Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2603.12639","last_updated":"2026-04-13T08:38:18Z","snapshot_observed_at":"2026-07-06T22:48:57.474272Z","submitted_at":"2026-03-13T04:16:19Z","title":"RoboStereo: Dual-Tower 4D Embodied World Models for Unified Policy Optimization","version":2},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-05-15T12:13:50.985113Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2603.12639"},"observation_digest":"sha256:f8edbd286e8b9abc2383717e32fed884b8c2f73fdfccd7aa91d1f3d26c3ff76a","observation_id":"9bcbec49-edef-4cfb-b0a0-b68d235da30d","resolution":{"observed_at":"2026-05-15T12:15:34.455286Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-07-14T21:59:30.885343Z","title":"arXiv preprint arXiv:2511.16719 (2025)","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2603.12939","last_updated":"2026-07-10T06:29:36Z","snapshot_observed_at":"2026-08-02T06:49:44.101445Z","submitted_at":"2026-03-13T12:34:26Z","title":"RoboStream: Weaving Spatio-Temporal Reasoning with Memory in Vision-Language Models for Robotics","version":2},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-07-14T21:59:30.885343Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2603.12939"},"observation_digest":"sha256:f868d9d1ec0f3bdfd13f836c3c830c106a2805aa5e71701ce8d184371798b837","observation_id":"092e52dc-08d3-44e7-9df4-d23ba99ce77e","resolution":{"observed_at":"2026-07-14T21:59:30.885343Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-07-14T21:36:54.932072Z","title":"arXiv preprint arXiv:2511.16719 (2025)","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2603.13854","last_updated":"2026-06-09T19:28:03Z","snapshot_observed_at":"2026-07-14T21:36:53.944875Z","submitted_at":"2026-03-14T09:22:52Z","title":"Power Term Polynomial Algebra for Boolean Logic","version":2},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-07-14T21:36:54.932072Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2603.13854"},"observation_digest":"sha256:222a3c2f55c6a0b16551121dd865ac464d3b403fcd6194488ffcf7f22d87229f","observation_id":"c32198cc-c6e5-4956-b528-4d10afc70f62","resolution":{"observed_at":"2026-07-14T21:36:54.932072Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2603.16869","last_updated":"2026-05-10T17:28:14Z","snapshot_observed_at":"2026-07-06T22:49:28.866886Z","submitted_at":"2026-03-17T17:59:51Z","title":"SegviGen: Repurposing 3D Generative Model for Part Segmentation","version":3},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-05-15T09:29:08.513396Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2603.16869"},"observation_digest":"sha256:eb262c08a38ce572217a0fb2b54a8ff9e08cd96494220b5658c6d9a5c4b02c28","observation_id":"51eaa2b5-760c-45ac-bd63-64c39a805556","resolution":{"observed_at":"2026-05-15T09:29:53.456812Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-07-13T22:49:03.259461Z","title":"arXiv preprint arXiv:2511.16719 (2025)","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2603.17975","last_updated":"2026-06-28T11:09:46Z","snapshot_observed_at":"2026-07-13T22:49:02.789541Z","submitted_at":"2026-03-18T17:39:05Z","title":"AHOY! Animatable Humans under Occlusion from YouTube Videos with Gaussian Splatting and Video Diffusion Priors","version":2},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-07-13T22:49:03.259461Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2603.17975"},"observation_digest":"sha256:38a9eb9ca2747ef612c6b1220e1dfb64cef8dbfdeaf5fdfc2bd70af95af4766f","observation_id":"42ab5947-cc6c-4164-90b3-5b8acf0b23a1","resolution":{"observed_at":"2026-07-13T22:49:03.259461Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-07-13T22:46:37.452699Z","title":"Sam 3: Segment anything with concepts.arXiv preprint arXiv:2511.16719, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2603.18208","last_updated":"2026-03-18T18:57:45Z","snapshot_observed_at":"2026-07-13T22:46:35.768521Z","submitted_at":"2026-03-18T18:57:45Z","title":"Quantum orientation entanglement analysis of the interpolating helicity states between the instant form dynamics and the light-front dynamics","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-07-13T22:46:37.452699Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2603.18208"},"observation_digest":"sha256:4095a62cb21ea8b3519ce745069fcb07c236abcbaacedc3dbc1cfb285de01336","observation_id":"e5f99ae3-0b2f-492f-9fe1-4a862a441a80","resolution":{"observed_at":"2026-07-13T22:46:37.452699Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2603.19684","last_updated":"2026-06-23T07:31:57Z","snapshot_observed_at":"2026-08-02T15:49:19.766021Z","submitted_at":"2026-03-20T06:32:16Z","title":"TSegAgent: Zero-Shot Tooth Segmentation via Geometry-Aware Vision-Language Agents","version":2},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-05-15T09:02:41.561003Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2603.19684"},"observation_digest":"sha256:4c9b2ede3968b738cd2f55d15877ca344c4a6b1fd5ba6a914628cd9a2395965b","observation_id":"7752e164-6d7c-4fe9-8b44-10694d87b156","resolution":{"observed_at":"2026-05-15T09:05:20.301673Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-02T17:52:42.442465Z","title":"arXiv preprint arXiv:2511.16719 (2025)","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2603.19802","last_updated":"2026-07-15T11:20:38Z","snapshot_observed_at":"2026-08-02T17:52:37.547780Z","submitted_at":"2026-03-20T09:40:41Z","title":"Evaluating Vision Foundation Models for Pixel and Object Classification in Microscopy","version":2},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-02T17:52:42.442465Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2603.19802"},"observation_digest":"sha256:967ee9cd2ba762179788ab5eeb8d4832bab6136a157933ed382b68de9ff29c24","observation_id":"adf46d93-1c07-46fe-b065-32748c062e05","resolution":{"observed_at":"2026-08-02T17:52:42.442465Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2603.20530","last_updated":"2026-04-20T18:41:00Z","snapshot_observed_at":"2026-07-06T22:49:57.103822Z","submitted_at":"2026-03-20T21:57:51Z","title":"Memory Over Maps: 3D Object Localization Without Reconstruction","version":2},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-05-15T07:46:50.155348Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2603.20530"},"observation_digest":"sha256:109d91a0f40769e84ba305e2c06a2f10c090ce0020a4b4572d718fa29faf80d5","observation_id":"2aebbc26-3777-408f-8d1e-4c709306d49f","resolution":{"observed_at":"2026-05-15T07:49:50.970769Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2603.22003","last_updated":"2026-05-09T06:59:22Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2026-03-23T14:08:58Z","title":"VP-VLA: Visual Prompting as an Interface for Vision-Language-Action Models","version":3},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-05-15T00:54:36.125258Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2603.22003"},"observation_digest":"sha256:034e7ae6a7ad46341501e1f2a9d2b1704f006c33f941f3192331cccca41b47b2","observation_id":"ba339898-8eff-4884-abb5-bddf055d02c1","resolution":{"observed_at":"2026-05-15T00:59:36.760683Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-07-13T19:38:38.452093Z","title":"arXiv preprint arXiv:2511.16719 (2025) 7, 8, 31","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2603.23455","last_updated":"2026-06-30T04:40:05Z","snapshot_observed_at":"2026-07-13T19:37:58.775361Z","submitted_at":"2026-03-24T17:26:55Z","title":"DetPO: In-Context Learning with Multi-Modal LLMs for Few-Shot Object Detection","version":2},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-07-13T19:38:38.452093Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2603.23455"},"observation_digest":"sha256:4ce54356e41a67684f1f5237514cccb0adbbf3998113fc749864cdc2098a543c","observation_id":"c0d0156a-a90a-48c1-ace4-14a837ab822c","resolution":{"observed_at":"2026-07-13T19:38:38.452093Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-07-13T18:24:58.428701Z","title":"arXiv preprint arXiv:2511.16719 (2025) 4","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2603.25168","last_updated":"2026-06-26T05:47:15Z","snapshot_observed_at":"2026-07-13T18:24:45.051987Z","submitted_at":"2026-03-26T08:37:32Z","title":"ET-SAM: Efficient Point Prompt Prediction in SAM for Unified Scene Text Detection and Layout Analysis","version":2},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-07-13T18:24:58.428701Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2603.25168"},"observation_digest":"sha256:fcfb1d06190c58d94f96e92ccc8148f1ee59ef35ac65028c1b67137acb19d1be","observation_id":"5472be30-f914-4499-a6cf-ec200f207c56","resolution":{"observed_at":"2026-07-13T18:24:58.428701Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-07-13T18:06:18.381075Z","title":"arXiv preprint arXiv:2511.16719 (2025) 6, 7","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2603.25645","last_updated":"2026-06-28T15:36:28Z","snapshot_observed_at":"2026-07-13T18:06:15.323946Z","submitted_at":"2026-03-26T16:58:43Z","title":"Colon-Bench: An Agentic Workflow for Scalable Dense Lesion Annotation in Full-Procedure Colonoscopy Videos","version":3},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-07-13T18:06:18.381075Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2603.25645"},"observation_digest":"sha256:dc1004930b99b4c6326b0bbbce42d59269ddd816c926cba84b84e611db2690da","observation_id":"9638240a-360f-4a82-b14b-0369bfdeb041","resolution":{"observed_at":"2026-07-13T18:06:18.381075Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-07-15T11:46:27.666803Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2603.27706","last_updated":"2026-07-14T09:25:01Z","snapshot_observed_at":"2026-07-17T23:18:00.863114Z","submitted_at":"2026-03-29T14:09:50Z","title":"MAR3: Multi-Agent Recognition, Reasoning, and Reflection for Reference Audio-Visual Segmentation","version":2},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-07-15T11:46:27.666803Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2603.27706"},"observation_digest":"sha256:c3dc13e7c68c540b4965727880146649710dc3e9053dc14167fb8a6b657a9ab7","observation_id":"b2a0b6d2-0404-4899-820f-7c5e40e342b0","resolution":{"observed_at":"2026-07-15T11:46:27.666803Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-07-13T16:13:29.986339Z","title":"SAM 3: Segment anything with concepts,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2603.28691","last_updated":"2026-06-27T17:09:33Z","snapshot_observed_at":"2026-07-13T16:13:29.034875Z","submitted_at":"2026-03-30T17:12:17Z","title":"DRIVE-Nav: Directional Reasoning, Inspection, and Verification for Efficient Open-Vocabulary Navigation","version":2},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-07-13T16:13:29.986339Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2603.28691"},"observation_digest":"sha256:bba59c3b54007e03c3be63850c8d0f35d48d929a2f44f7434d5aaaa43bf4f30c","observation_id":"4c51da11-f076-4283-b950-bd59cd1cfa9a","resolution":{"observed_at":"2026-07-13T16:13:29.986339Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-07-13T14:40:24.818842Z","title":"arXiv preprint arXiv:2511.16719 (2025)","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2604.01001","last_updated":"2026-07-01T11:45:08Z","snapshot_observed_at":"2026-07-13T14:40:18.914932Z","submitted_at":"2026-04-01T15:00:46Z","title":"EgoSim: Egocentric World Simulator for Embodied Interaction Generation","version":2},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-07-13T14:40:24.818842Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2604.01001"},"observation_digest":"sha256:c339aeaa621a8491bf4496e695d570c6d18d5a871355ffa854d36453a5d219e6","observation_id":"9db1ee9b-c2d5-465d-b2ec-d45e6f9eb50c","resolution":{"observed_at":"2026-07-13T14:40:24.818842Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2604.02509","last_updated":"2026-04-02T21:07:17Z","snapshot_observed_at":"2026-08-02T19:41:06.034769Z","submitted_at":"2026-04-02T21:07:17Z","title":"Rapidly deploying on-device eye tracking by distilling visual foundation models","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-05-13T21:38:55.454229Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2604.02509"},"observation_digest":"sha256:d53a6fda23f000af8d53279661c3f98008c412fe4fd653ddb328731c73dbb843","observation_id":"d54443b7-90a1-4b65-8230-c44d5aed02b5","resolution":{"observed_at":"2026-05-13T21:43:19.083820Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2604.02593","last_updated":"2026-04-03T00:09:14Z","snapshot_observed_at":"2026-08-02T13:20:20.763889Z","submitted_at":"2026-04-03T00:09:14Z","title":"Moondream Segmentation: From Words to Masks","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-05-13T20:14:45.629804Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2604.02593"},"observation_digest":"sha256:63b10aa30bec42ac0260475e8e9493cf990f59548730a238f02d5b5a45fc7ac5","observation_id":"ebfc333b-c4b7-4910-b9fb-54e1ec95ff0b","resolution":{"observed_at":"2026-05-13T20:18:13.720508Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2604.02773","last_updated":"2026-04-03T06:32:18Z","snapshot_observed_at":"2026-07-06T22:52:06.460027Z","submitted_at":"2026-04-03T06:32:18Z","title":"Generalized Small Object Detection:A Point-Prompted Paradigm and Benchmark","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-05-13T20:11:01.013199Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2604.02773"},"observation_digest":"sha256:fc0eb977e8fed118e11b7432791aa0ea2c227cdcc77943e527d4f5b3ac91af7f","observation_id":"7a048f2d-b3bd-4b28-b6c7-4248c5a52b06","resolution":{"observed_at":"2026-05-13T20:13:13.562022Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-07-13T13:13:22.720572Z","title":"Sam 3: Segment anything with concepts,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2604.03547","last_updated":"2026-04-04T02:24:19Z","snapshot_observed_at":"2026-08-02T07:30:16.328049Z","submitted_at":"2026-04-04T02:24:19Z","title":"KappaFormer: Physics-aware Transformer for lattice thermal conductivity via cross-domain transfer learning","version":1},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-07-13T13:13:22.720572Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2604.03547"},"observation_digest":"sha256:486f02e8c483f8bc7219ae7495142de532c3e22c417baff7163cf95f0aa4c171","observation_id":"da13772d-4786-412b-b9df-2e7d51ff0359","resolution":{"observed_at":"2026-07-13T13:13:22.720572Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2604.04473","last_updated":"2026-04-06T06:48:32Z","snapshot_observed_at":"2026-08-02T18:27:59.544984Z","submitted_at":"2026-04-06T06:48:32Z","title":"Beyond Standard Benchmarks: A Systematic Audit of Vision-Language Model's Robustness to Natural Semantic Variation Across Diverse Tasks","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-05-10T20:26:42.128350Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2604.04473"},"observation_digest":"sha256:2857cada8db7d30a016662b4d85ade6e3c85e516c97cc288d543d88d24500b22","observation_id":"e914de77-f439-4790-90f7-f5d68b1d84f0","resolution":{"observed_at":"2026-05-10T21:55:52.598548Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2604.04500","last_updated":"2026-04-06T07:51:59Z","snapshot_observed_at":"2026-07-30T06:32:02.248567Z","submitted_at":"2026-04-06T07:51:59Z","title":"Saliency-R1: Enforcing Interpretable and Faithful Vision-language Reasoning via Saliency-map Alignment Reward","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-05-10T19:59:19.379119Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2604.04500"},"observation_digest":"sha256:810717b7d970f4b557b54de038a1529f4745e92b716d34fbac9312ed9df30f41","observation_id":"8ca6106a-f7f0-49e3-b8c6-73c28b58fa9a","resolution":{"observed_at":"2026-05-10T22:20:47.635927Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2604.04911","last_updated":"2026-04-08T04:54:06Z","snapshot_observed_at":"2026-07-06T22:53:46.999911Z","submitted_at":"2026-04-06T17:54:42Z","title":"SpatialEdit: Benchmarking Fine-Grained Image Spatial Editing","version":2},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-05-10T19:23:50.614589Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2604.04911"},"observation_digest":"sha256:6e2ae0f1f273748aef45c10b4f9084f372460940f686cd9356e0e292d0608f46","observation_id":"f447b548-1169-41c7-ba3d-97cd4f91449e","resolution":{"observed_at":"2026-05-10T23:00:50.445165Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2604.04969","last_updated":"2026-07-12T03:10:04Z","snapshot_observed_at":"2026-08-02T05:19:18.740150Z","submitted_at":"2026-04-04T07:14:01Z","title":"MG$^2$-RAG: Multi-Granularity Graph for Multimodal Retrieval-Augmented Generation","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-05-13T17:12:38.618233Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2604.04969"},"observation_digest":"sha256:fd5ae39314c46036e578ff9bf8c140d6d9fee2ead410d8f98b082c2059350836","observation_id":"4a4c7d2c-9aad-415f-af00-b2ce866881a4","resolution":{"observed_at":"2026-05-13T17:13:00.935260Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-07-14T19:54:25.294067Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2604.04969","last_updated":"2026-07-12T03:10:04Z","snapshot_observed_at":"2026-08-02T05:19:18.740150Z","submitted_at":"2026-04-04T07:14:01Z","title":"MG$^2$-RAG: Multi-Granularity Graph for Multimodal Retrieval-Augmented Generation","version":2},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-07-14T19:54:25.294067Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2604.04969"},"observation_digest":"sha256:b448853b0438cd6e15024a044381ff2aa98fee068e0c7c9ee8721977b7d978f9","observation_id":"42201890-872f-42b3-9d46-69034011ef9a","resolution":{"observed_at":"2026-07-14T19:54:25.294067Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2604.06332","last_updated":"2026-04-07T18:13:55Z","snapshot_observed_at":"2026-08-04T02:17:12.505784Z","submitted_at":"2026-04-07T18:13:55Z","title":"Telescope: Learnable Hyperbolic Foveation for Ultra-Long-Range Object Detection","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-05-10T18:46:48.550865Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2604.06332"},"observation_digest":"sha256:86450be980e19d0b6db949de1783be1ceb3e4a282e2b63fe3a1990c9cbe12d27","observation_id":"baa270bb-7c9a-46c3-bc6a-45b790afcd3f","resolution":{"observed_at":"2026-05-10T23:55:52.689492Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2604.06671","last_updated":"2026-04-08T04:45:17Z","snapshot_observed_at":"2026-07-06T22:55:06.424752Z","submitted_at":"2026-04-08T04:45:17Z","title":"4D Vessel Reconstruction for Benchtop Thrombectomy Analysis","version":1},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-05-10T18:16:38.725224Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2604.06671"},"observation_digest":"sha256:d8f1eafc137d2bbc9e98700f74b6a3d7345efc54935efb342cb0c75adb7e5494","observation_id":"4ac903d8-f713-4e18-9c55-8725f3500c0e","resolution":{"observed_at":"2026-05-10T20:25:46.968273Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2604.06725","last_updated":"2026-04-08T06:47:55Z","snapshot_observed_at":"2026-08-02T13:53:46.692069Z","submitted_at":"2026-04-08T06:47:55Z","title":"Enhancing MLLM Spatial Understanding via Active 3D Scene Exploration for Multi-Perspective Reasoning","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-05-10T19:16:46.753641Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2604.06725"},"observation_digest":"sha256:000b1ca7b6102ec06f0d4b2c1695224d192193a6a5bcb64e0b3a5e0bdf23bcb6","observation_id":"d9efa4a3-eb88-4be4-b8d0-cf4b27715efc","resolution":{"observed_at":"2026-05-10T23:15:47.800798Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2604.06870","last_updated":"2026-04-08T09:32:15Z","snapshot_observed_at":"2026-07-06T22:55:15.791334Z","submitted_at":"2026-04-08T09:32:15Z","title":"RefineAnything: Multimodal Region-Specific Refinement for Perfect Local Details","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-05-10T19:11:43.172296Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2604.06870"},"observation_digest":"sha256:b0c10c9b50e43ee67374ac2bb2408319da3a86b233e18006e781fbd34e57b9db","observation_id":"03c6a5e6-ec34-460a-b925-1e56076ae1cb","resolution":{"observed_at":"2026-05-10T23:20:54.386356Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2604.07230","last_updated":"2026-04-09T11:01:59Z","snapshot_observed_at":"2026-07-06T22:55:33.502756Z","submitted_at":"2026-04-08T15:53:57Z","title":"PhyEdit: Towards Real-World Object Manipulation via Physically-Grounded Image Editing","version":2},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-05-10T18:54:47.451103Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2604.07230"},"observation_digest":"sha256:fc7065391d6a7dc469d9c0d122aac74718a49c3b35f36109acccabab5bd5b194","observation_id":"5d5b4d84-0efc-4650-a5bb-0b9d7b8c2737","resolution":{"observed_at":"2026-05-10T23:45:50.485034Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2604.07818","last_updated":"2026-04-24T03:58:27Z","snapshot_observed_at":"2026-07-06T22:57:05.146373Z","submitted_at":"2026-04-09T05:20:07Z","title":"Open-Ended Video Game Glitch Detection with Agentic Reasoning and Temporal Grounding","version":2},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-05-10T18:24:23.950768Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2604.07818"},"observation_digest":"sha256:77c65dfb958a859220c51c24d81ffbb50fa7f2850757602834dddb81b4ed9ba3","observation_id":"60c2e017-28a6-4ccb-8bea-055e2ab0cee2","resolution":{"observed_at":"2026-05-11T00:41:03.640206Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2604.07916","last_updated":"2026-08-04T13:16:24Z","snapshot_observed_at":"2026-08-05T07:53:06.638032Z","submitted_at":"2026-04-09T07:37:09Z","title":"Tarot-SAM3: Training-free SAM3 for Any Referring Expression Segmentation","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-05-10T18:11:13.376684Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2604.07916"},"observation_digest":"sha256:785986cb70237bcf493fb916ba3e1e08f00e1623d10911406ee678b33726c1b2","observation_id":"d044e828-4fe9-421c-a63c-eaace89e72c8","resolution":{"observed_at":"2026-05-11T05:21:00.963573Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2604.08626","last_updated":"2026-04-17T22:49:55Z","snapshot_observed_at":"2026-07-06T22:57:40.773778Z","submitted_at":"2026-04-09T16:00:10Z","title":"WildDet3D: Scaling Promptable 3D Detection in the Wild","version":2},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-05-10T17:38:13.336003Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2604.08626"},"observation_digest":"sha256:d065b5a4bfa7f1d5c2d7c3f9d19817cfca082290c6b65dbe330d1b8e7593d102","observation_id":"6d65d5da-0eeb-4d8b-baf6-0fe445301d1e","resolution":{"observed_at":"2026-05-11T06:26:00.661055Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2604.08664","last_updated":"2026-04-09T18:00:05Z","snapshot_observed_at":"2026-07-06T22:57:45.084987Z","submitted_at":"2026-04-09T18:00:05Z","title":"Generative Simulation for Policy Learning in Physical Human-Robot Interaction","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-05-10T17:06:58.099881Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2604.08664"},"observation_digest":"sha256:a7b34fee90c5dec171b81e19299faddde17734c634c713428abc017ea2952a5e","observation_id":"fa03c5b0-d5af-403f-8b03-6252d42c85d4","resolution":{"observed_at":"2026-05-11T07:35:57.540789Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2604.09045","last_updated":"2026-04-10T07:07:10Z","snapshot_observed_at":"2026-07-06T22:57:59.555972Z","submitted_at":"2026-04-10T07:07:10Z","title":"Scene-Agnostic Object-Centric Representation Learning for 3D Gaussian Splatting","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-05-10T17:29:54.777890Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2604.09045"},"observation_digest":"sha256:4f088ba753c410b3c22899a530f4ca0b854580ba89dd26cc2836ae8fa3c9781a","observation_id":"111f5887-71b7-455a-a96f-a6bf345e301b","resolution":{"observed_at":"2026-05-11T06:41:32.285461Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2604.09167","last_updated":"2026-04-10T09:51:42Z","snapshot_observed_at":"2026-07-06T22:58:08.579543Z","submitted_at":"2026-04-10T09:51:42Z","title":"MAG-3D: Multi-Agent Grounded Reasoning for 3D Understanding","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-05-10T17:25:31.097385Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2604.09167"},"observation_digest":"sha256:7cf8d0bc38390d98bfa4dc7f2d97c981adc182830ea2f907ab856508c68f34f6","observation_id":"7ab1e9fa-41d6-4c0a-af89-462ecc305336","resolution":{"observed_at":"2026-05-11T06:51:23.259350Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2604.09690","last_updated":"2026-04-06T13:17:14Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2026-04-06T13:17:14Z","title":"Are We Recognizing the Jaguar or Its Background? A Diagnostic Framework for Jaguar Re-Identification","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-05-10T19:55:56.261335Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2604.09690"},"observation_digest":"sha256:a097dc884ce45af599111032a5d97d491b060ffa3e27fda0dd5d83696ee6a6eb","observation_id":"a13b7f91-19a5-4b3a-ac7f-43c4c9217348","resolution":{"observed_at":"2026-05-10T22:20:48.897656Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2604.09781","last_updated":"2026-06-29T09:07:17Z","snapshot_observed_at":"2026-07-12T23:07:19.301563Z","submitted_at":"2026-04-10T18:06:02Z","title":"Text-Guided 6D Object Pose Rearrangement via Closed-Loop VLM Agents","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-05-10T16:55:48.506049Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2604.09781"},"observation_digest":"sha256:84aadb2d3d5649b7d61d7a6cee29429a1faf4beda4ab771834f96f9643acbe9b","observation_id":"168d0bfb-403b-44a3-834c-c1486b418785","resolution":{"observed_at":"2026-05-11T07:51:01.853072Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-07-12T23:07:40.699400Z","title":"arXiv:2511.16719 (2025)","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2604.09781","last_updated":"2026-06-29T09:07:17Z","snapshot_observed_at":"2026-07-12T23:07:19.301563Z","submitted_at":"2026-04-10T18:06:02Z","title":"Text-Guided 6D Object Pose Rearrangement via Closed-Loop VLM Agents","version":2},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-07-12T23:07:40.699400Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2604.09781"},"observation_digest":"sha256:cd60c7d6645a4a01605935fdbceba5a4d7073c8781a3f1bc9ecec7f2f41757ca","observation_id":"503401ba-2aca-47fe-8edc-bbc80f3a4436","resolution":{"observed_at":"2026-07-12T23:07:40.699400Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2604.09920","last_updated":"2026-04-10T21:30:45Z","snapshot_observed_at":"2026-07-06T22:58:38.807460Z","submitted_at":"2026-04-10T21:30:45Z","title":"Does Your VFM Speak Plant? The Botanical Grammar of Vision Foundation Models for Object Detection","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-05-10T17:07:17.174815Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2604.09920"},"observation_digest":"sha256:66d95231e48306942800c0041d60f0a0ebcb3452aa1c2e33179af209ab69c1e9","observation_id":"36de14df-a458-47f3-8c76-e5127d64d09c","resolution":{"observed_at":"2026-05-11T07:35:57.207091Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2604.10108","last_updated":"2026-05-17T04:14:11Z","snapshot_observed_at":"2026-07-06T22:58:47.599383Z","submitted_at":"2026-04-11T09:00:54Z","title":"JARVIS: A Just-in-Time Augmented Reality VLM-Powered Instruction System for Cross-Reality Task Guidance","version":2},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-05-10T16:21:13.585442Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2604.10108"},"observation_digest":"sha256:0de0ca5e6da5fe55e9a15f543a25c1f5a65a76ac62d1a903bb7ae7a845b32cb6","observation_id":"155740ca-7606-4443-9319-dddf1e39e451","resolution":{"observed_at":"2026-05-11T09:00:59.364126Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2604.10108","last_updated":"2026-05-17T04:14:11Z","snapshot_observed_at":"2026-07-06T22:58:47.599383Z","submitted_at":"2026-04-11T09:00:54Z","title":"JARVIS: A Just-in-Time Augmented Reality VLM-Powered Instruction System for Cross-Reality Task Guidance","version":3},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-05-21T01:28:04.022431Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2604.10108"},"observation_digest":"sha256:18a0199b76797e3fd0b5f184940e20ff0cb769b65e3d603968dd7cc53f1e82f0","observation_id":"e6aafc74-8893-4eea-8ee0-df58b60350e3","resolution":{"observed_at":"2026-05-21T01:29:22.420008Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2604.10132","last_updated":"2026-04-11T09:53:09Z","snapshot_observed_at":"2026-07-06T22:58:51.931478Z","submitted_at":"2026-04-11T09:53:09Z","title":"Semantic Manipulation Localization","version":1},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-05-10T15:32:30.345594Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2604.10132"},"observation_digest":"sha256:8d382d62b384c39df1781e993fd6e2d4830937f13414b1bebdf69c0c8938437c","observation_id":"f24a5d44-da82-4fd1-9018-6315e28ade09","resolution":{"observed_at":"2026-05-11T10:21:01.491684Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-07-12T22:20:17.569353Z","title":"arXiv preprint arXiv:2511.16719 (2025)","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2604.10788","last_updated":"2026-05-30T11:15:20Z","snapshot_observed_at":"2026-08-02T23:24:19.874403Z","submitted_at":"2026-04-12T19:38:19Z","title":"TInR: Exploring Tool-Internalized Reasoning in Large Language Models","version":2},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-07-12T22:20:17.569353Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2604.10788"},"observation_digest":"sha256:fab6172138684a4cb67230c695ffb2c725048deaaaa851011ce17c280e883f7f","observation_id":"8e1bdb60-defa-46c5-beb0-53e74c76113a","resolution":{"observed_at":"2026-07-12T22:20:17.569353Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2604.10789","last_updated":"2026-04-12T19:42:12Z","snapshot_observed_at":"2026-07-06T22:59:18.756089Z","submitted_at":"2026-04-12T19:42:12Z","title":"ReplicateAnyScene: Zero-Shot Video-to-3D Composition via Textual-Visual-Spatial Alignment","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-05-10T15:38:03.377613Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2604.10789"},"observation_digest":"sha256:7c3a01096c70998d58b4ee8b16a73bd6177cb40d03d480e7a97b469be061873d","observation_id":"7114e3ae-75e2-4503-8731-ef478c689045","resolution":{"observed_at":"2026-05-11T10:06:05.406339Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2604.11170","last_updated":"2026-04-13T08:29:49Z","snapshot_observed_at":"2026-08-02T16:48:26.324351Z","submitted_at":"2026-04-13T08:29:49Z","title":"Do Instance Priors Help Weakly Supervised Semantic Segmentation?","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-07-11T11:50:26.030339Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2604.11170"},"observation_digest":"sha256:64a538efa2ceff1b540df8eace031cdbda7081e46caa7679a02ddfeecae5f0e1","observation_id":"0ae249b4-71a2-4afb-9457-fa6e5d8886b6","resolution":{"observed_at":"2026-05-10T16:30:36.253577Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2604.11231","last_updated":"2026-04-13T09:35:14Z","snapshot_observed_at":"2026-07-06T22:59:40.900521Z","submitted_at":"2026-04-13T09:35:14Z","title":"Seg2Change: Adapting Open-Vocabulary Semantic Segmentation Model for Remote Sensing Change Detection","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-05-10T15:25:38.084033Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2604.11231"},"observation_digest":"sha256:ff0d5dc1559710b1ff17856a7e25e4cbd9807dafacb90639dfe6f73356150bc5","observation_id":"466f2ded-49f6-4c17-af7c-d20ee55e9d6c","resolution":{"observed_at":"2026-05-11T10:36:02.861601Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2604.11411","last_updated":"2026-04-13T12:55:56Z","snapshot_observed_at":"2026-08-03T03:42:10.089825Z","submitted_at":"2026-04-13T12:55:56Z","title":"Online Reasoning Video Object Segmentation","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-05-10T15:29:03.843441Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2604.11411"},"observation_digest":"sha256:fdcde78d5518a5f74aa0af5246ae6f224a68837bfd71e596f91445ce79ee4101","observation_id":"3c196fa8-8911-4584-b03b-baf8eaa14d18","resolution":{"observed_at":"2026-05-11T10:31:00.229266Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2604.11789","last_updated":"2026-04-20T14:38:53Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2026-04-13T17:55:02Z","title":"LMMs Meet Object-Centric Vision: Understanding, Segmentation, Editing and Generation","version":2},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-05-10T15:35:37.095627Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2604.11789"},"observation_digest":"sha256:0ddf3f9a4d82275b21068a7b06c09d4c36756f2e3571deef2250ecdf8b2930bb","observation_id":"a152bce4-0965-4234-8347-7b5e03f39abb","resolution":{"observed_at":"2026-05-11T10:11:04.723038Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2604.11998","last_updated":"2026-04-13T19:38:49Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2026-04-13T19:38:49Z","title":"The Second Challenge on Cross-Domain Few-Shot Object Detection at NTIRE 2026: Methods and Results","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-05-10T16:05:16.070144Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2604.11998"},"observation_digest":"sha256:a9ca989b24913fab017e58800aba5c428feb9ee780410790f572e867aca6ae3a","observation_id":"47ed6040-9e11-43fd-8c76-a79a40a03696","resolution":{"observed_at":"2026-05-11T09:21:00.673605Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2604.12144","last_updated":"2026-07-01T22:34:57Z","snapshot_observed_at":"2026-08-02T05:18:11.450897Z","submitted_at":"2026-04-13T23:48:35Z","title":"VERITAS: A Multi-Agent Co-Scientist for Verifiable Image-Derived Hypothesis Testing","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-05-10T14:56:06.372343Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2604.12144"},"observation_digest":"sha256:960d2025fdb8c57f9a5acacb064565faddde8fd6824d68420e75370f5f20c1a3","observation_id":"648187dc-b0eb-4d1e-8f4b-9360cb3fb12e","resolution":{"observed_at":"2026-05-11T11:26:02.154005Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2604.12929","last_updated":"2026-04-14T16:19:12Z","snapshot_observed_at":"2026-07-06T23:01:00.830207Z","submitted_at":"2026-04-14T16:19:12Z","title":"Grasp in Gaussians: Fast Monocular Reconstruction of Dynamic Hand-Object Interactions","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-05-10T15:01:52.354836Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2604.12929"},"observation_digest":"sha256:89b4967ba5d6cc730a0d8cf06d781e12505ecaa01c9075dc83ab8da95cfda3fb","observation_id":"cba463f3-cbad-47d6-ab9e-86ce111eed04","resolution":{"observed_at":"2026-05-11T11:21:00.599214Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2604.14268","last_updated":"2026-04-15T17:59:17Z","snapshot_observed_at":"2026-07-06T23:02:05.116649Z","submitted_at":"2026-04-15T17:59:17Z","title":"HY-World 2.0: A Multi-Modal World Model for Reconstructing, Generating, and Simulating 3D Worlds","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-05-10T13:45:24.961208Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2604.14268"},"observation_digest":"sha256:34e184bdeb8f8a50a913e5ac1ecaf18a737cba11b1d65913d25ea3bf1ff0dc94","observation_id":"1cbf2e62-bd01-40e5-8945-0f6bbbf8c937","resolution":{"observed_at":"2026-05-10T13:45:27.276875Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2604.14302","last_updated":"2026-07-18T16:18:55Z","snapshot_observed_at":"2026-08-02T16:18:23.645980Z","submitted_at":"2026-04-15T18:00:45Z","title":"Geometrically Consistent Multi-View Scene Generation from Freehand Sketches","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-05-10T13:23:58.169493Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2604.14302"},"observation_digest":"sha256:55e4c0c3464592b36577b8f76ed08df0bafbfc2f484dc0271281c9ac1597396e","observation_id":"0d52dcb8-965e-4f38-bb0c-c15c017c191c","resolution":{"observed_at":"2026-05-10T13:25:26.025915Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-02T16:18:25.431492Z","title":"arXiv preprint arXiv:2511.16719 (2025)","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2604.14302","last_updated":"2026-07-18T16:18:55Z","snapshot_observed_at":"2026-08-02T16:18:23.645980Z","submitted_at":"2026-04-15T18:00:45Z","title":"Geometrically Consistent Multi-View Scene Generation from Freehand Sketches","version":2},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-02T16:18:25.431492Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2604.14302"},"observation_digest":"sha256:d7e0ce44ed3fd3e57f309207d887ce3cd7e9973f85b715f867f614c009d1d110","observation_id":"9c1b7aa7-1c15-487f-9810-a14ae8d8bd88","resolution":{"observed_at":"2026-08-02T16:18:25.431492Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2604.14944","last_updated":"2026-06-19T11:34:50Z","snapshot_observed_at":"2026-07-12T19:56:30.664166Z","submitted_at":"2026-04-16T12:38:14Z","title":"HRDexDB: A Paired Human-Robot Dataset for Cross-Embodiment Dexterous Grasping","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-05-10T10:48:27.803794Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2604.14944"},"observation_digest":"sha256:e32bd7192c476a13c05d7793d512eeac6b3f17f8f4e9bf6ac5ba172bf410e578","observation_id":"45d12334-47f7-402c-8d6c-0cf00521b1ee","resolution":{"observed_at":"2026-05-10T10:49:56.066534Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-07-12T19:56:39.725379Z","title":"Carion, L","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2604.14944","last_updated":"2026-06-19T11:34:50Z","snapshot_observed_at":"2026-07-12T19:56:30.664166Z","submitted_at":"2026-04-16T12:38:14Z","title":"HRDexDB: A Paired Human-Robot Dataset for Cross-Embodiment Dexterous Grasping","version":2},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-07-12T19:56:39.725379Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2604.14944"},"observation_digest":"sha256:6bf8e51b539f340155473743b7cefbc972e2383105d9685bbd09a40ecbb8398a","observation_id":"cccf5171-c325-4e0a-b3e4-5ef86068b754","resolution":{"observed_at":"2026-07-12T19:56:39.725379Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2604.15299","last_updated":"2026-04-16T17:57:08Z","snapshot_observed_at":"2026-07-06T23:02:53.677687Z","submitted_at":"2026-04-16T17:57:08Z","title":"AnimationBench: Are Video Models Good at Character-Centric Animation?","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-05-10T11:47:27.131584Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2604.15299"},"observation_digest":"sha256:310e2f3f8f30976440d02f74e24caab899a2d1917058b1f45dbb634651c1cd16","observation_id":"fb75c2a4-e3ae-4692-a011-dbe84b769df2","resolution":{"observed_at":"2026-05-10T11:50:20.361212Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2604.17231","last_updated":"2026-04-19T03:31:02Z","snapshot_observed_at":"2026-07-06T23:04:23.730478Z","submitted_at":"2026-04-19T03:31:02Z","title":"Fringe Projection Based Vision Pipeline for Autonomous Hard Drive Disassembly","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-05-10T06:56:58.988003Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2604.17231"},"observation_digest":"sha256:a173e21c3adeb3b122a11f7b29448620562b2bad810c532bea047fe977ecef10","observation_id":"32a54687-4f8d-4393-aefe-39c01e1bff4b","resolution":{"observed_at":"2026-05-10T07:01:49.507373Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2604.17530","last_updated":"2026-04-19T16:45:19Z","snapshot_observed_at":"2026-08-03T16:41:36.178114Z","submitted_at":"2026-04-19T16:45:19Z","title":"Real-Time Cellist Postural Evaluation With On-Device Computer Vision","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-05-10T05:38:35.025500Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2604.17530"},"observation_digest":"sha256:df83f1c1b2653ce9a5882079ac0f619ddd21e98e6607e1b31d73ae54ab4c70cf","observation_id":"f6859d33-06f2-4c06-95f5-47bb7e761261","resolution":{"observed_at":"2026-05-10T05:41:02.090528Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2604.18201","last_updated":"2026-04-20T12:50:26Z","snapshot_observed_at":"2026-07-06T23:05:08.728446Z","submitted_at":"2026-04-20T12:50:26Z","title":"DiffuSAM: Diffusion Guided Zero-Shot Object Grounding for Remote Sensing Imagery","version":1},"reference_index":4,"source":"arxiv_source","source_observed_at":"2026-05-10T05:06:06.159341Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2604.18201"},"observation_digest":"sha256:6cb81ef9464061cbe302d7a0dd1728a40375fc80545d151d6210c464fce6a7b3","observation_id":"4b0be7c8-f854-4344-a1cb-ebe14ff0cc55","resolution":{"observed_at":"2026-05-10T10:09:08.371735Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2604.18225","last_updated":"2026-05-12T20:12:53Z","snapshot_observed_at":"2026-08-02T15:30:09.471390Z","submitted_at":"2026-04-20T13:10:07Z","title":"Is SAM3 ready for pathology segmentation?","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-05-10T04:40:57.284414Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2604.18225"},"observation_digest":"sha256:99af79ca787556decf270829ea7e1f05d3786f0a3cf294c2141dd87d053cfda6","observation_id":"362df0b7-cba3-4610-837c-a1504f26bbaf","resolution":{"observed_at":"2026-05-10T12:05:23.345458Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2604.18225","last_updated":"2026-05-12T20:12:53Z","snapshot_observed_at":"2026-08-02T15:30:09.471390Z","submitted_at":"2026-04-20T13:10:07Z","title":"Is SAM3 ready for pathology segmentation?","version":2},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-05-14T21:54:14.103479Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2604.18225"},"observation_digest":"sha256:acea990797bb56cd8ec50e21052e43b092564d40a5eb65f919bb83e3820f9d28","observation_id":"5cbea517-b97b-4613-bff1-5c6693c024e4","resolution":{"observed_at":"2026-05-14T21:58:03.759666Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2604.18665","last_updated":"2026-04-20T14:40:43Z","snapshot_observed_at":"2026-07-06T23:05:30.879164Z","submitted_at":"2026-04-20T14:40:43Z","title":"APRVOS: 1st Place Winner of 5th PVUW MeViS-Audio Track","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-05-10T03:29:12.783993Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2604.18665"},"observation_digest":"sha256:1a06f257c9f0316a7d913b8d1703cf7537465237258f5e8f7426604941d69d90","observation_id":"b9a83476-dbcf-4bfc-8484-6763f09f7e8e","resolution":{"observed_at":"2026-05-10T03:29:21.560465Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2604.19324","last_updated":"2026-04-21T10:46:42Z","snapshot_observed_at":"2026-07-31T05:48:40.687274Z","submitted_at":"2026-04-21T10:46:42Z","title":"PLaMo 2.1-VL Technical Report","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-05-10T03:09:05.436900Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2604.19324"},"observation_digest":"sha256:231e844f56717065e0b3f598e8f4d784ad5cada81414e8ca12a4e40130a83ab0","observation_id":"3c57f55c-6c02-44ac-91cb-7da2f25c5a26","resolution":{"observed_at":"2026-05-11T12:46:02.530784Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2604.19591","last_updated":"2026-04-22T09:29:15Z","snapshot_observed_at":"2026-08-01T17:38:58.491384Z","submitted_at":"2026-04-21T15:42:34Z","title":"Structure-Semantic Decoupled Modulation of Global Geospatial Embeddings for High-Resolution Remote Sensing Mapping","version":2},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-05-10T02:06:05.390586Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2604.19591"},"observation_digest":"sha256:bf8ce6b031255880bee2d86f621d191c86554ef30d30b5df78a1e6d3a350825d","observation_id":"7da65e17-8820-4d66-b5b0-e1e53d27b9d8","resolution":{"observed_at":"2026-05-11T13:16:06.030436Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2604.19609","last_updated":"2026-04-21T15:56:26Z","snapshot_observed_at":"2026-08-03T15:00:15.869225Z","submitted_at":"2026-04-21T15:56:26Z","title":"Volume Transformer: Revisiting Vanilla Transformers for 3D Scene Understanding","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-05-10T03:33:56.211514Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2604.19609"},"observation_digest":"sha256:8e0e6c2d39762fa11a432437f204f5402884231e772a88463beefd7c2f1be3a6","observation_id":"ec16c773-e5db-44b5-a1a6-2b7cc9db6172","resolution":{"observed_at":"2026-05-11T12:31:03.223382Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2604.19636","last_updated":"2026-04-21T16:25:43Z","snapshot_observed_at":"2026-07-06T23:06:16.972438Z","submitted_at":"2026-04-21T16:25:43Z","title":"CoInteract: Physically-Consistent Human-Object Interaction Video Synthesis via Spatially-Structured Co-Generation","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-05-10T03:14:45.834520Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2604.19636"},"observation_digest":"sha256:7911f86251af597432603fc441d8328265575da62f38a1a3755c72f3cab51bb5","observation_id":"61af82f1-2a10-4679-8046-ee4ed96b151e","resolution":{"observed_at":"2026-05-11T12:41:04.398550Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2604.19648","last_updated":"2026-04-21T16:37:18Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2026-04-21T16:37:18Z","title":"CoCo-SAM3: Harnessing Concept Conflict in Open-Vocabulary Semantic Segmentation","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-05-10T03:10:10.398335Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2604.19648"},"observation_digest":"sha256:c713bf0fea44b736358833b590e3b2229a52b784c59ca6246556d0e8d7c3a09a","observation_id":"0dbc6211-a1cc-4abe-9791-3e525f6677a2","resolution":{"observed_at":"2026-05-10T03:14:08.261351Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2604.20258","last_updated":"2026-04-22T07:08:01Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2026-04-22T07:08:01Z","title":"Rethinking Where to Edit: Task-Aware Localization for Instruction-Based Image Editing","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-05-10T00:39:21.872643Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2604.20258"},"observation_digest":"sha256:e6c093e0f503c2876e60750476a7cddb5a815ec612aec000fda9c5730e55b892","observation_id":"903d84b6-01e3-4107-b94d-d3c08964ba91","resolution":{"observed_at":"2026-05-10T00:39:48.383774Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2604.20329","last_updated":"2026-06-03T18:02:24Z","snapshot_observed_at":"2026-07-06T23:06:49.210284Z","submitted_at":"2026-04-22T08:23:48Z","title":"Image Generators are Generalist Vision Learners","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-05-10T01:14:05.034951Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2604.20329"},"observation_digest":"sha256:5936a9ea55a563245ac4df4eb9c883922e56ee4db2a8059768bb0318d71acbb7","observation_id":"4b49dd26-26a1-45f9-8555-f9c41b018145","resolution":{"observed_at":"2026-05-11T13:41:06.000924Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2604.20329","last_updated":"2026-06-03T18:02:24Z","snapshot_observed_at":"2026-07-06T23:06:49.210284Z","submitted_at":"2026-04-22T08:23:48Z","title":"Image Generators are Generalist Vision Learners","version":2},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-05-15T07:40:46.090808Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2604.20329"},"observation_digest":"sha256:49025b2b2d7e459759b506d5ec818ea8216b5d352bfb77306f4b56d3092c3398","observation_id":"af3f4cfb-5577-4c1b-8d9a-e907ffa95274","resolution":{"observed_at":"2026-05-15T07:45:14.665089Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2604.21502","last_updated":"2026-05-22T15:42:19Z","snapshot_observed_at":"2026-08-04T02:10:59.871292Z","submitted_at":"2026-04-23T10:04:36Z","title":"VFM$^{4}$SDG: Unveiling the Power of VFMs for Single-Domain Generalized Object Detection","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-05-09T21:38:15.877684Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2604.21502"},"observation_digest":"sha256:0fc5092ca14c6efc617961b1fd0a92a889a7245c2fe621cf9ac7dabed0487a83","observation_id":"a9475ea4-8305-4180-84e9-853d6caf6d10","resolution":{"observed_at":"2026-05-11T14:31:07.242137Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2604.21502","last_updated":"2026-05-22T15:42:19Z","snapshot_observed_at":"2026-08-04T02:10:59.871292Z","submitted_at":"2026-04-23T10:04:36Z","title":"VFM$^{4}$SDG: Unveiling the Power of VFMs for Single-Domain Generalized Object Detection","version":2},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-05-25T06:03:56.040612Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2604.21502"},"observation_digest":"sha256:2d58b0a8698899bfe493687eceab16d36927c950971edc0bbf365947a8cf1f73","observation_id":"7810396d-378d-4d90-a2ab-75b1d2f5e5dd","resolution":{"observed_at":"2026-05-25T06:05:26.336364Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2604.21952","last_updated":"2026-04-23T05:27:39Z","snapshot_observed_at":"2026-08-02T07:47:00.163884Z","submitted_at":"2026-04-23T05:27:39Z","title":"Focus Session: Hardware and Software Techniques for Accelerating Multimodal Foundation Models","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-05-09T23:02:07.554154Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2604.21952"},"observation_digest":"sha256:23e500f9d1c34ff53d231fae9dfa20e1daf320ca0ac3b82c2f4b8c63d8cf8b06","observation_id":"d73ef529-973a-48c6-904c-7c4e31211a91","resolution":{"observed_at":"2026-05-09T23:04:17.777657Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2604.22836","last_updated":"2026-04-20T14:36:19Z","snapshot_observed_at":"2026-07-06T23:09:10.050398Z","submitted_at":"2026-04-20T14:36:19Z","title":"AgentRVOS for MeViS-Text Track of 5th PVUW Challenge: 3rd Method","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-05-10T05:14:25.423302Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2604.22836"},"observation_digest":"sha256:36c84213948ed24f44079bed80a935552b1cbd10b94b417a9fc741850877bc87","observation_id":"10f282c6-061e-4efb-9c2a-0bcb6fcbebc3","resolution":{"observed_at":"2026-05-10T09:33:41.849359Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2604.22837","last_updated":"2026-04-20T14:43:43Z","snapshot_observed_at":"2026-07-06T23:09:10.050398Z","submitted_at":"2026-04-20T14:43:43Z","title":"OAMVOS:2nd Report for 5th PVUW MOSE Track","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-05-10T05:04:56.366788Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2604.22837"},"observation_digest":"sha256:dcc9f0471dc4969eda1a1107fc0681eba69a14b1367a481553af0039b43e267f","observation_id":"f2f1cc56-c5dd-4618-a64e-a4a3eb0c0e6c","resolution":{"observed_at":"2026-05-10T10:09:08.582081Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2604.22875","last_updated":"2026-04-28T04:48:22Z","snapshot_observed_at":"2026-08-03T02:43:56.733861Z","submitted_at":"2026-04-23T22:33:15Z","title":"SketchVLM: Vision language models can annotate images to explain thoughts and guide users","version":2},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-05-09T21:32:08.503584Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2604.22875"},"observation_digest":"sha256:b8dadba8ac62cda161d362039102b6499e4a252a9177894c68b95ede07ba7653","observation_id":"43085a33-3f39-40ad-86d0-acc05b57c688","resolution":{"observed_at":"2026-05-11T14:36:05.439172Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2604.23095","last_updated":"2026-04-25T01:17:45Z","snapshot_observed_at":"2026-07-06T23:09:23.591544Z","submitted_at":"2026-04-25T01:17:45Z","title":"INSIGHT: Indoor Scene Intelligence from Geometric-Semantic Hierarchy Transfer for Public~Safety","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-05-08T08:38:36.288087Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2604.23095"},"observation_digest":"sha256:2e3b568328914dcb837e53ac450c681752bf7b281c75bea3db3387b737336ede","observation_id":"7dabd64c-9557-42c7-934f-77ebfa623672","resolution":{"observed_at":"2026-05-11T20:31:14.589423Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2604.23249","last_updated":"2026-05-04T07:58:10Z","snapshot_observed_at":"2026-07-29T22:29:59.267331Z","submitted_at":"2026-04-25T11:01:27Z","title":"BridgeACT: Bridging Human Demonstrations to Robot Actions via Unified Tool-Target Affordances","version":2},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-05-08T07:55:35.304883Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2604.23249"},"observation_digest":"sha256:66b8f2fd14ad919f410b3253238de33356965beb403d0f1072980ee556cf3de2","observation_id":"0fb5166b-55d3-42dc-bedb-745b02c5409e","resolution":{"observed_at":"2026-05-11T20:46:15.633882Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2511.16719/citation-record","integrity":"/paper/2511.16719/integrity","json":"/paper/2511.16719/citation-record.json","paper":"/paper/2511.16719"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-09T00:05:48.311803Z","title":"write newline","venue":null,"work_id":"8e5fda61-e601-4df4-8204-015bee341570","year":null},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":1,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:015aced6bb3a62ce3e11057ec00026319c36a9fc3584479ebe3162efeb30513b","observation_id":"f7eaaba7-80d5-40fa-a4e9-01409a57e85c","resolution":{"observed_at":"2026-05-17T20:25:12.542845Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Greenhouse gas equivalencies calculator","venue":null,"work_id":"6fdd7255-496e-4272-9155-2c7d016e5c4f","year":2022},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":2,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:73238e10d51719dc222b544762346965dd924218dc549a9815a6bfc29ee865e4","observation_id":"0fd3a05d-19e2-40f6-bfaf-7bf887d25995","resolution":{"observed_at":"2026-05-17T20:25:12.539950Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Multi-label cluster discrimination for visual representation learning","venue":null,"work_id":"9f8ab084-1917-4978-b58b-43ebdf4809df","year":2024},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":3,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:c3de00a1a6df8f862530efc763be7fe69100658fd14d73b95a7fd1d9737dcf4b","observation_id":"bc05edb3-0754-4fca-be81-836b309940bf","resolution":{"observed_at":"2026-05-17T20:25:12.261769Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Burst: A benchmark for unifying object recognition, segmentation and tracking in video","venue":null,"work_id":"f683a8ec-ee9b-4032-b7a1-7176e9925b05","year":2023},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":4,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:a615eda53464eada7ea3fe9cb0bbb6ef23e7eafd0563f0f233dbfce0aaa15d5f","observation_id":"f7709340-2aa5-486c-9c89-e19951adb736","resolution":{"observed_at":"2026-05-17T20:25:12.315170Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Gmot-40: A benchmark for generic multiple object tracking","venue":null,"work_id":"1c33e85d-12bb-4045-8291-ae5a83f5fa7d","year":2021},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":5,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:8ecc70bee9c6490a2acb5294fd1d783f9c002f064f7a560aec7dd1b3d25610a3","observation_id":"4c8b77f4-a836-4cd7-b6d2-c92cf9ded9a3","resolution":{"observed_at":"2026-05-17T20:25:12.309425Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2509.03499","last_updated":"2025-09-03T17:30:53Z","snapshot_observed_at":"2026-08-05T02:49:44.363970Z","submitted_at":"2025-09-03T17:30:53Z","title":"DeepSea MOT: A benchmark dataset for multi-object tracking on deep-sea video","version":1},"cited_work":{"arxiv_id":"2509.03499","doi":"10.48550/arxiv.2509.03499","metadata_source":"arxiv_reference","pith_arxiv_id":"2509.03499","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"DeepSea MOT : A benchmark dataset for multi-object tracking on deep-sea video","venue":"ArXiv.org","work_id":"babe81fb-9c12-440e-b34b-08bebd29fc6a","year":2025},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":6,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"cited_paper":"/paper/2509.03499","citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:b971fd2fb84889cf87611be8e2dae64b0bba86b7e9dc9b687e8b4c2268101eb3","observation_id":"2a4cb442-38fc-4d48-ae34-39900a057a93","resolution":{"observed_at":"2026-05-17T20:25:11.218899Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Tracking without bells and whistles","venue":null,"work_id":"037c49e1-5060-419c-bba4-f6977d8282fb","year":2019},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":7,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:51d5b7f96ece55bb27733cdf9171e0a2678306727eb2c7ff86ff1989b00905a2","observation_id":"dd2d058e-26a5-4deb-ac77-25956a0821df","resolution":{"observed_at":"2026-05-17T20:25:12.320140Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Simple online and realtime tracking","venue":null,"work_id":"06c4aad4-0f08-4345-ad7c-64310be173b0","year":2016},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:23ab6cd9bd74c4e90b1b96b15b8bb40a034b46bf7150793c498814df4e105a97","observation_id":"a5229b02-6507-4bc9-92e9-a3734517d11d","resolution":{"observed_at":"2026-05-17T20:25:12.343909Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2407.07726","last_updated":"2024-10-10T17:28:23Z","snapshot_observed_at":"2026-08-04T10:58:38.750074Z","submitted_at":"2024-07-10T14:57:46Z","title":"PaliGemma: A versatile 3B VLM for transfer","version":2},"cited_work":{"arxiv_id":"2407.07726","doi":"10.48550/arxiv.2407.07726","metadata_source":"pith","pith_arxiv_id":"2407.07726","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"PaliGemma: A versatile 3B VLM for transfer","venue":"cs.CV","work_id":"df6f48b3-5792-47c7-9614-cb856ea31ad9","year":2024},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":9,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"cited_paper":"/paper/2407.07726","citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:6a12219cdd77ec806fa2258aa9c69625f47ad8ee735ab3a6c7ad4af0091b7278","observation_id":"024d1d5e-f874-45fa-8e01-a2f7a1bf7662","resolution":{"observed_at":"2026-05-17T20:25:11.400127Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-05-22T13:22:31.538292+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-22T13:22:31.538292+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2004.10934","last_updated":"2020-04-23T02:10:02Z","snapshot_observed_at":"2026-07-06T09:14:32.318388Z","submitted_at":"2020-04-23T02:10:02Z","title":"YOLOv4: Optimal Speed and Accuracy of Object Detection","version":1},"cited_work":{"arxiv_id":"2004.10934","doi":"10.3390/s22052341","metadata_source":"pith","pith_arxiv_id":"2004.10934","snapshot_observed_at":"2026-07-11T11:50:26.030339Z","title":"YOLOv4: Optimal Speed and Accuracy of Object Detection","venue":"cs.CV","work_id":"7057aaee-27f6-4209-a83c-f59727f937a8","year":2020},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":10,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"cited_paper":"/paper/2004.10934","citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:517d8ba67472b5e014b0113d396f4a120ba1d707dead8f429fb2d3b9a9f62fd1","observation_id":"84dfcf42-aef0-41a5-8fe4-7c0887d10576","resolution":{"observed_at":"2026-05-17T20:25:11.417758Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Window attention is bugged: How not to interpolate position embeddings","venue":null,"work_id":"b44fb7a2-258b-4da3-9078-5bacea345f24","year":2024},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:d5507e2d580d2e8bb1096cf5daa3dc3e48aaa47279fa22bdad5f41140678225c","observation_id":"050aa7e8-efa4-4c69-9730-0ee1b4ba80d8","resolution":{"observed_at":"2026-05-17T20:25:12.346435Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.13181","last_updated":"2025-04-28T18:01:39Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-04-17T17:59:57Z","title":"Perception Encoder: The best visual embeddings are not at the output of the network","version":2},"cited_work":{"arxiv_id":"2504.13181","doi":"10.48550/arxiv.2504.13181","metadata_source":"pith","pith_arxiv_id":"2504.13181","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Perception Encoder: The best visual embeddings are not at the output of the network","venue":"cs.CV","work_id":"409be941-4d4a-4ceb-a28a-eaa2d7709a1c","year":2025},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":12,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"cited_paper":"/paper/2504.13181","citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:428b1955dbc82c1155153f2d488734f07147b73032b10a7494993486a121ea60","observation_id":"865c7c00-f5dc-43fd-ad58-056324471fea","resolution":{"observed_at":"2026-05-17T20:25:11.582160Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Align-detr: Enhancing end-to-end object detection with aligned loss","venue":null,"work_id":"d381dc15-426b-4d76-887a-85c87e17253f","year":2024},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":13,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:cc669714e47c03df23b7da8d76c163081bde8b88fa92f2f3b48ef206d745b781","observation_id":"3984b8c9-bfd6-4b9b-aba5-751c86b575a0","resolution":{"observed_at":"2026-05-17T20:25:12.379980Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Observation-centric sort: Rethinking sort for robust multi-object tracking","venue":null,"work_id":"e7cb507e-e16a-4425-bf82-c108ce5d379d","year":2023},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":14,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:c88c7f8bdc177436ca5e851337fe7b9bc47367d384de0942cdd3a5b7c38b1300","observation_id":"93e86de1-a942-4a4d-a285-10faee621bd0","resolution":{"observed_at":"2026-05-17T20:25:12.267527Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"End-to-end object detection with transformers","venue":null,"work_id":"a84866a3-c2bc-4855-97ad-883d58e4abe4","year":2020},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":15,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:304b391052be71a7fce9fd7a5347a519ee1e225c6e65eab065e58e0231a0d849","observation_id":"1f7bd674-e994-43e3-afb6-59ed2f422ba4","resolution":{"observed_at":"2026-05-17T20:25:12.368528Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.03459","last_updated":"2024-06-05T17:07:24Z","snapshot_observed_at":"2026-07-06T18:26:01.717861Z","submitted_at":"2024-06-05T17:07:24Z","title":"LW-DETR: A Transformer Replacement to YOLO for Real-Time Detection","version":1},"cited_work":{"arxiv_id":"2406.03459","doi":"10.48550/arxiv.2406.03459","metadata_source":"arxiv_reference","pith_arxiv_id":"2406.03459","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Lw-detr: A transformer replacement to yolo for real-time detection","venue":"arXiv (Cornell University)","work_id":"6bdf4f58-6cf5-4c7a-9bf8-3329106fde70","year":2024},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":16,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"cited_paper":"/paper/2406.03459","citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:ed343820df059b368043444b2bfe92e7a5e774c112bf581e15bbcfd788246857","observation_id":"896aa9b6-6c1d-4d56-a6e2-844dd3e2287c","resolution":{"observed_at":"2026-05-17T20:25:11.523127Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Sam4mllm: Enhance multi-modal large language model for referring expression segmentation","venue":null,"work_id":"a53769ab-9bbd-4644-8ca6-f4d42acdde9e","year":2024},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":17,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:bb7db1308781627f9039f3bb04c6dae69c5b9110d35a988042fd72da17f9e460","observation_id":"1ce86dc3-e963-4488-8dd7-817603627b44","resolution":{"observed_at":"2026-05-17T20:25:12.264670Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Re-aligning language to visual objects with an agentic workflow","venue":null,"work_id":"6bff4888-b635-4b6b-acab-44111923d810","year":2025},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":18,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:87d3222d79f333f0d792e809f70a492b9da3d242bd08b3866b66f8ddbfea1faf","observation_id":"5b6836b1-e40e-45b9-8df0-66b4806832a0","resolution":{"observed_at":"2026-05-17T20:25:12.366172Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Schwing, and Alexander Kirillov","venue":null,"work_id":"183c58ab-8b72-4e51-b453-02f3aa09813c","year":2021},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":19,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:6911b00a55ef3dc14bf09ecc471d63e4b4634be109561a2559c7db0f49d2cd07","observation_id":"ee75b244-db1b-4a98-b3eb-9e737e84aac9","resolution":{"observed_at":"2026-05-17T20:25:12.370781Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.13180","last_updated":"2025-07-23T19:22:35Z","snapshot_observed_at":"2026-08-01T20:25:06.490113Z","submitted_at":"2025-04-17T17:59:56Z","title":"PerceptionLM: Open-Access Data and Models for Detailed Visual Understanding","version":3},"cited_work":{"arxiv_id":"2504.13180","doi":"10.48550/arxiv.2504.13180","metadata_source":"arxiv_reference","pith_arxiv_id":"2504.13180","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Perceptionlm: Open-access data and models for detailed visual understanding","venue":"ArXiv.org","work_id":"e81d25ef-bd13-473f-8fc5-0cf3dc3af528","year":2025},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":20,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"cited_paper":"/paper/2504.13180","citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:9f9dc5d0d7a66268a8c0c0d994668042804b136b7a0fac00cedbf97ce157cda7","observation_id":"9b872af4-c4ba-4404-9fd9-9132e9fe6cae","resolution":{"observed_at":"2026-05-17T20:25:11.567661Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"ELECTRA : Pre-training text encoders as discriminators rather than generators","venue":null,"work_id":"d5902951-e3aa-4f92-8d05-ff7fcb4b6d53","year":2020},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":21,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:78a9932c68888e511a504cff2bfa3cacdcf6b0b082209527ca6d87e2ea90b84d","observation_id":"9f83b8a9-4cfb-4b17-8a0c-76d1acb456cc","resolution":{"observed_at":"2026-05-17T20:25:12.394649Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2507.06261","last_updated":"2025-12-19T14:25:46Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-07-07T17:36:04Z","title":"Gemini 2.5: Pushing the Frontier with Advanced Reasoning, Multimodality, Long Context, and Next Generation Agentic Capabilities","version":6},"cited_work":{"arxiv_id":"2507.06261","doi":"10.48550/arxiv.2503.19","metadata_source":"pith","pith_arxiv_id":"2507.06261","snapshot_observed_at":"2026-07-11T03:17:51.364436Z","title":"Gemini 2.5: Pushing the Frontier with Advanced Reasoning, Multimodality, Long Context, and Next Generation Agentic Capabilities","venue":"cs.CL","work_id":"008df105-2fdd-45d8-857a-8e35868aecb6","year":2025},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":22,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"cited_paper":"/paper/2507.06261","citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:7686f00ba4509e49bcf036333cccf7e54990393232c714fe6e1755098abdc8a2","observation_id":"af39eeb1-05bd-4021-a9ba-e91c9578a9d6","resolution":{"observed_at":"2026-05-17T20:25:11.445429Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-08T02:14:26.692730Z","title":"The cityscapes dataset for semantic urban scene understanding","venue":null,"work_id":"f41683b7-7f15-49c2-ab76-6c949f21dfad","year":2016},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":23,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:b24ea3fdc310a50daefc5808f42ddef7a0e200495ac800b2feef6e6ed5a4de6c","observation_id":"362b6e69-149a-4c61-a79e-b87270247b25","resolution":{"observed_at":"2026-05-17T20:25:12.504006Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2102.01066","last_updated":"2022-03-15T06:04:05Z","snapshot_observed_at":"2026-08-05T01:25:27.457563Z","submitted_at":"2021-02-01T18:56:02Z","title":"Evaluating Large-Vocabulary Object Detectors: The Devil is in the Details","version":2},"cited_work":{"arxiv_id":"2102.01066","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2102.01066","snapshot_observed_at":"2026-07-03T08:07:45.472888Z","title":"Evaluating large-vocabulary object detectors: The devil is in the details","venue":null,"work_id":"6e954585-f1a1-4e94-94e6-33bc03e53417","year":2021},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":24,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"cited_paper":"/paper/2102.01066","citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:dcf05ca7d1097fe5c2b0d1ac95d2fc16a59ff1e26d36fe5089a358eddb646c2c","observation_id":"1c0563c3-0eef-463a-bbab-44b536af04da","resolution":{"observed_at":"2026-05-17T20:25:11.431192Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Molmo and pixmo: Open weights and open data for state-of-the-art vision-language models","venue":null,"work_id":"a3de73ca-2059-46f0-94a4-14090442be2a","year":2025},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":25,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:c1e6afe061ef0cb38eab1c4eade7edc68cfdfdd9826e3074e5d39440c3853087","observation_id":"627f38f0-73bb-4bd9-aa7a-a851a3934ff2","resolution":{"observed_at":"2026-05-17T20:25:12.353805Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2508.05630","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-10T03:06:43.588255Z","title":"MOSEv2: A more challenging dataset for video object segmentation in complex scenes","venue":null,"work_id":"5e313244-b6d4-40a5-ba70-6f060ec438db","year":2025},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":26,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:a4a69a8127ea1480f4a8253a802407d87a55b8e9e96f9f1b4dea905ea07af0b1","observation_id":"f7264151-16c5-4eac-bce8-4a8c3b6eab47","resolution":{"observed_at":"2026-05-17T20:25:11.492824Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"A large-scale synthetic pathological dataset for deep learning-enabled segmentation of breast cancer","venue":null,"work_id":"5032417a-9e16-42a0-9ff8-42a77abf9c01","year":2023},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":27,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:c9810cde1c958af8394041a025e4e2620df806d46d5ed4a614b1c5f40a7d3f11","observation_id":"583d7e1f-081e-4836-8b46-80cfef7a4eb4","resolution":{"observed_at":"2026-05-17T20:25:12.336699Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.16268","last_updated":"2025-07-29T00:08:01Z","snapshot_observed_at":"2026-07-06T19:37:16.203681Z","submitted_at":"2024-10-21T17:59:19Z","title":"SAM2Long: Enhancing SAM 2 for Long Video Segmentation with a Training-Free Memory Tree","version":3},"cited_work":{"arxiv_id":"2410.16268","doi":"10.48550/arxiv.2410.16268","metadata_source":"arxiv_reference","pith_arxiv_id":"2410.16268","snapshot_observed_at":"2026-07-10T06:15:00.866473Z","title":"Sam2long: Enhancing sam 2 for long video seg- mentation with a training-free memory tree","venue":null,"work_id":"981cafae-2190-4d02-9705-052f04a0a251","year":2024},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":28,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"cited_paper":"/paper/2410.16268","citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:c3293ac7008118a33da65450de4c95d31de3ea3e7f3c0d5710491bd4b0fef4d0","observation_id":"554a66a4-47b5-4ab3-972c-077b837a3b5e","resolution":{"observed_at":"2026-05-17T20:25:11.455972Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2208.08984","last_updated":"2023-06-08T06:35:33Z","snapshot_observed_at":"2026-07-06T13:43:13.386359Z","submitted_at":"2022-08-18T17:55:37Z","title":"Open-Vocabulary Universal Image Segmentation with MaskCLIP","version":2},"cited_work":{"arxiv_id":"2208.08984","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2208.08984","snapshot_observed_at":"2026-07-03T08:57:47.583278Z","title":"Open- vocabulary universal image segmentation with MaskCLIP","venue":null,"work_id":"afa10fbf-e01c-48af-9b6b-e884ef223fc6","year":2022},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":29,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"cited_paper":"/paper/2208.08984","citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:a43889872a85abf0b874e72f18df7c64530fcde0e2276f11d31f9c90bf01a3ea","observation_id":"35a756b1-9635-468d-9d7b-f68ac69fd194","resolution":{"observed_at":"2026-05-17T20:25:11.372628Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2206.07643","last_updated":"2022-11-18T18:23:08Z","snapshot_observed_at":"2026-07-06T13:21:17.244944Z","submitted_at":"2022-06-15T16:41:29Z","title":"Coarse-to-Fine Vision-Language Pre-training with Fusion in the Backbone","version":2},"cited_work":{"arxiv_id":"2206.07643","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2206.07643","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Coarse-to-fine vision-language pre-training with fusion in the backbone","venue":null,"work_id":"4fe3323a-87ba-4386-8376-deee76f67d5a","year":2022},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":30,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"cited_paper":"/paper/2206.07643","citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:fdf1be671453414b894db8df89e99c6c3898c78788ea6e0a415d1e32b929377d","observation_id":"6b3b6b65-7403-46f3-bc16-68e45c254392","resolution":{"observed_at":"2026-05-17T20:25:11.547619Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"The llama 3 herd of models","venue":null,"work_id":"f87f4155-69a0-4eb8-a0e8-8dbb5dadd63c","year":2024},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":31,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:a4ef29069533d528f95f277fe21a66fa5ab3724f2c059cd3eda24d8fcbcee4ba","observation_id":"0520de45-545c-4b27-8c35-5c99a2ccbba0","resolution":{"observed_at":"2026-05-17T20:25:12.331965Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Livecell—a large-scale dataset for label-free live cell segmentation","venue":null,"work_id":"4b1ce876-c14c-4959-9d33-c8102bee20c5","year":2021},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":32,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:2e70165917990b889936217f55481eaec30aea563b4d245d627e0761b617cff7","observation_id":"ed25477c-7854-4e8c-9377-8f6fe279cdd2","resolution":{"observed_at":"2026-05-17T20:25:12.361465Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Detect to track and track to detect","venue":null,"work_id":"1f6a33fa-1c34-42ae-8c4f-1d915f38f211","year":2017},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":33,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:8c654d276c7f5a6b9dad4aeda72f4b5bc76c15d41208a36fd2d2063b109f1dbf","observation_id":"5a7738df-9ac7-4812-91ad-c42bc6cc4caf","resolution":{"observed_at":"2026-05-17T20:25:12.327049Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"0e19cccc-e7a9-449b-8846-e932dc933da3","year":null},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":34,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:c3fe2eecee59b644c8c8bfc5fa3c347f5d815521f70f69baa367a46b27929423","observation_id":"587a2805-146d-42db-a560-871b140a842b","resolution":{"observed_at":"2026-05-17T20:25:12.334280Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Llmdet: Learning strong open-vocabulary object detectors under the supervision of large language models","venue":null,"work_id":"b21c1a6d-e2a9-4151-b36b-66f18928f21d","year":2025},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":35,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:72594e636eb1a0bde17adafc480255049136acf4fb3d884a86f4e4b9cef16da7","observation_id":"f3ab3d75-285a-4ec9-806c-df92ec7ad3ec","resolution":{"observed_at":"2026-05-17T20:25:12.341566Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Pannuke: an open pan-cancer histology dataset for nuclei instance segmentation and classification","venue":null,"work_id":"a0e392e9-a3d7-4137-a30c-6becc46ac0bb","year":2019},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":36,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:772bdcf4e9d902ff8d7da9e1034d9a7da195433bc9a926dba77cad7275cc79da","observation_id":"4688bd17-a57f-4e42-a9d6-d3342f4343cb","resolution":{"observed_at":"2026-05-17T20:25:12.359197Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2003.10778","last_updated":"2020-04-22T08:52:04Z","snapshot_observed_at":"2026-07-06T09:06:56.021993Z","submitted_at":"2020-03-24T11:25:12Z","title":"PanNuke Dataset Extension, Insights and Baselines","version":7},"cited_work":{"arxiv_id":"2003.10778","doi":"10.48550/arxiv.2003.10778","metadata_source":"arxiv_reference","pith_arxiv_id":"2003.10778","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Gamper, N","venue":"arXiv (Cornell University)","work_id":"6339ffd6-9882-4acb-a5d2-4b0160adf8b3","year":2003},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":37,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"cited_paper":"/paper/2003.10778","citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:8d236529f1ead1bc7fb7eecaecf8e8184ecd60ec05f202e4161020d17b5ef470","observation_id":"7f2a7a4e-512f-4469-ba0d-a9bb15fb0cf4","resolution":{"observed_at":"2026-05-17T20:25:11.411152Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"a32242a1-1289-4a3f-9838-4be3e5e83b98","year":2022},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":38,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:618fdd271a9864530af492499ee5f2f3e58689383f4065bbba52abe6c2552391","observation_id":"e01f6412-018d-4dda-8069-0b69a99a2a77","resolution":{"observed_at":"2026-05-17T20:25:12.317721Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2104.13921","last_updated":"2022-05-12T01:27:40Z","snapshot_observed_at":"2026-07-06T11:04:30.929441Z","submitted_at":"2021-04-28T17:58:57Z","title":"Open-vocabulary Object Detection via Vision and Language Knowledge Distillation","version":3},"cited_work":{"arxiv_id":"2104.13921","doi":null,"metadata_source":"pith","pith_arxiv_id":"2104.13921","snapshot_observed_at":"2026-07-04T16:29:57.549240Z","title":"Open-vocabulary Object Detection via Vision and Language Knowledge Distillation","venue":"cs.CV","work_id":"59541f32-18cf-4328-ad60-c9018e1401cf","year":2021},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":39,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"cited_paper":"/paper/2104.13921","citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:279406021375e4b7a067da23888b734f240066894cbe011708a78f7d0427ad42","observation_id":"48f4fb16-2244-44e7-9240-93aa5bf2ddb4","resolution":{"observed_at":"2026-05-17T20:25:11.414343Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Lvis: A dataset for large vocabulary instance segmentation","venue":null,"work_id":"2bf44652-e2eb-466e-9581-53e6d2a54069","year":2019},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":40,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:32b89b23e79ca94fb7363117bccf99337e876c0161d99d0d5481150f10028230","observation_id":"98d8d0b7-3353-4487-8414-b87483717b75","resolution":{"observed_at":"2026-05-17T20:25:12.322438Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Masked autoencoders are scalable vision learners","venue":null,"work_id":"7ae4bc08-37d5-4c44-a102-d5a1144c7b18","year":2022},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":41,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:e643ba053f5cf2a245c05d5cda3cf88b554fb80cde749f0227683ae5288de605","observation_id":"b3086335-c28b-4a49-ab4c-806d3097b412","resolution":{"observed_at":"2026-05-17T20:25:12.305113Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2403.13298","last_updated":"2024-07-16T04:54:56Z","snapshot_observed_at":"2026-07-06T17:47:30.518562Z","submitted_at":"2024-03-20T04:47:13Z","title":"Rotary Position Embedding for Vision Transformer","version":2},"cited_work":{"arxiv_id":"2403.13298","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2403.13298","snapshot_observed_at":"2026-07-01T13:45:45.264681Z","title":"Rotary position embedding for vision transformer","venue":null,"work_id":"87a5840d-d8e6-4dfc-b4b5-adcba3bde377","year":2024},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":42,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"cited_paper":"/paper/2403.13298","citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:50e47cffeb9a8e4857dc372d5e305d35584bb9825c710cd94668bf8c81890254","observation_id":"7e5377c8-8e90-4d66-99c8-52d0b7d89111","resolution":{"observed_at":"2026-05-17T20:25:11.403942Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2404.19326","last_updated":"2024-05-01T01:30:58Z","snapshot_observed_at":"2026-07-06T18:07:30.890899Z","submitted_at":"2024-04-30T07:50:29Z","title":"LVOS: A Benchmark for Large-scale Long-term Video Object Segmentation","version":2},"cited_work":{"arxiv_id":"2404.19326","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2404.19326","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Lvos: A benchmark for large- scale long-term video object segmentation","venue":null,"work_id":"4b9c54a4-eb00-4397-8960-5c8ba7aa28c6","year":2024},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":43,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"cited_paper":"/paper/2404.19326","citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:a05892465ca4bde08c402de1044765133af8d6d3229f3d550872e2f65d058f4f","observation_id":"d3d4a90f-82dc-4e05-b183-0fa5a395fbdc","resolution":{"observed_at":"2026-05-17T20:25:11.554163Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.5281/zenodo.1212303","metadata_source":"doi_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-02T13:16:57.966770Z","title":"author Montani, I","venue":"Zenodo (CERN European Organization for Nuclear Research)","work_id":"21f271e7-c123-4e13-856a-053f2da387a8","year":2020},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":44,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:c96fcfc9bb33df50a81729d728a44f601ec1ca29adf01dc15c78c7058a780bf5","observation_id":"6eb8893b-5295-49d8-b31d-0131b2258318","resolution":{"observed_at":"2026-05-17T20:25:11.211822Z","resolver_source":"doi","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-03T15:08:47.996614+00:00","source":"crossref_status_cache"},{"observed_at":"2026-08-03T15:08:47.996614+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1707.06642","last_updated":"2018-04-10T20:22:13Z","snapshot_observed_at":"2026-07-06T05:52:02.469903Z","submitted_at":"2017-07-20T17:59:55Z","title":"The iNaturalist Species Classification and Detection Dataset","version":2},"cited_work":{"arxiv_id":"1707.06642","doi":null,"metadata_source":"pith","pith_arxiv_id":"1707.06642","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"The iNaturalist Species Classification and Detection Dataset","venue":"cs.CV","work_id":"4dd193ee-7628-47ee-bc94-b9175ca934c6","year":2017},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":45,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"cited_paper":"/paper/1707.06642","citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:f8e9b03b068d848ef48fb8a2f26fdeb6c5d2cb77b9cd7a1994ef0487fa725554","observation_id":"f8f33147-a929-4cd6-bba5-99402320d17f","resolution":{"observed_at":"2026-05-17T20:25:11.387280Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"DAC-DETR : Divide the attention layers and conquer","venue":null,"work_id":"383a274a-fcd8-4b97-88e7-392af65bdabe","year":2023},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":46,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:81cfe4cec201b085d56145cb39c7c7ccdf1d642f36e8d98f9da44fcacdd21546","observation_id":"873da0d1-28ce-4939-88b1-d602b1d22d33","resolution":{"observed_at":"2026-05-17T20:25:12.275119Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Densely connected parameter-efficient tuning for referring image segmentation","venue":null,"work_id":"cc8312ed-c61a-4e6b-8783-d6b19fcc164a","year":2025},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":47,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:d740d9902949ac14ac09d0cbc111ed62358f753a82c5cc4f7b189d37941131a6","observation_id":"788ce252-4f03-460e-b466-5e2311b1a7d3","resolution":{"observed_at":"2026-05-17T20:25:12.277772Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2207.13080","last_updated":"2023-05-16T16:01:50Z","snapshot_observed_at":"2026-07-06T13:35:39.787667Z","submitted_at":"2022-07-26T17:52:14Z","title":"DETRs with Hybrid Matching","version":3},"cited_work":{"arxiv_id":"2207.13080","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2207.13080","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Detrs with hybrid matching","venue":null,"work_id":"d0c64444-c991-4baf-9623-002319300391","year":2022},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":48,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"cited_paper":"/paper/2207.13080","citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:dcd5dd6875d68f78239f6359d44e131b9f2ecd70abd9056bd3a1dccaf913e771","observation_id":"83f811d4-57be-402f-ae80-99729362d600","resolution":{"observed_at":"2026-05-17T20:25:11.393789Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2004.12276","last_updated":"2020-07-18T21:02:49Z","snapshot_observed_at":"2026-08-02T07:13:53.469881Z","submitted_at":"2020-04-26T02:38:26Z","title":"Fashionpedia: Ontology, Segmentation, and an Attribute Localization Dataset","version":2},"cited_work":{"arxiv_id":"2004.12276","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2004.12276","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Belongie","venue":null,"work_id":"072a21e5-f899-46ea-bf06-36ea0868c605","year":2004},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":49,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"cited_paper":"/paper/2004.12276","citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:d5aa97cfaac54818190dbd338742cb7e14a56e4d754a0460f0371061a5fab9d5","observation_id":"9e67fa51-2b6c-4f95-8a2b-531112388380","resolution":{"observed_at":"2026-05-17T20:25:11.564043Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2504.04519","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-02T03:06:30.030991Z","title":"Sam2mot: A novel paradigm of multi-object tracking by segmentation","venue":null,"work_id":"6628c27b-4f5b-4b4b-897b-c69907f56e08","year":2025},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":50,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:ae8be8beed14ec25ac8bf4de1c3a8399957b14984e747dfc18a2c2991af66a83","observation_id":"e218e163-f96b-4d89-a6c9-78a09ef003bd","resolution":{"observed_at":"2026-05-17T20:25:11.427667Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"T-rex2: Towards generic object detection via text-visual prompt synergy","venue":null,"work_id":"81b67627-f67d-4e81-8679-9e25ca58d9e1","year":2024},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":51,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:5c5f927cf43585b2eb047b9fed0a84ead2a0f7176e444645789b64f580dc8468","observation_id":"3152d228-2a7f-4164-8087-beb56eeb75bf","resolution":{"observed_at":"2026-05-17T20:25:12.282360Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Trackeval","venue":null,"work_id":"d647e808-06ed-4813-849a-5f8e059474c9","year":2020},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":52,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:81250153fd6c7ad1e54a6411a91029716a3a0b74056165d81f1a3f3b1d526852","observation_id":"a8cac499-bd4c-413d-91a6-a35043d95ce1","resolution":{"observed_at":"2026-05-17T20:25:12.272382Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Mdetr-modulated detection for end-to-end multi-modal understanding","venue":null,"work_id":"c72d57d5-2efe-463a-a1c7-50cfbaf4c396","year":2021},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":53,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:77358e6652a8f253c2fd8fa7c2e299b6a4345466b2bde5c056dc304c8988c01b","observation_id":"09ea6adb-25e3-4a50-aad8-cd84ca1c8415","resolution":{"observed_at":"2026-05-17T20:25:12.501610Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Your large vision-language model only needs a few attention heads for visual grounding","venue":null,"work_id":"52e45c95-58c6-4066-a7bf-68034336064e","year":2025},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":54,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:844994412b43de37c110a4850f6a534d9e6f380a91833a8464a97deaee53af68","observation_id":"4de276d2-9682-45f1-9322-b2fd161654e0","resolution":{"observed_at":"2026-05-17T20:25:12.509076Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2109.14646","last_updated":"2022-09-07T19:46:22Z","snapshot_observed_at":"2026-07-06T11:52:46.440188Z","submitted_at":"2021-09-29T18:08:42Z","title":"FathomNet: A global image database for enabling artificial intelligence in the ocean","version":4},"cited_work":{"arxiv_id":"2109.14646","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2109.14646","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Orenstein, Brian Schlining, Lonny Lundsten, Kevin Barnard, Giovanna Sainz, Oceane Boulais, Benjamin G","venue":null,"work_id":"9f288fb6-c753-451d-ba07-779d1cf98300","year":2021},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":55,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"cited_paper":"/paper/2109.14646","citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:3c8d53239487a08a43e8edbb3f34b8e9639faa8386d792d348f1bc6f56ebb1d6","observation_id":"ce1242b9-1664-4dd6-96f0-66606f860f1e","resolution":{"observed_at":"2026-05-17T20:25:11.452255Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Referitgame: Referring to objects in photographs of natural scenes","venue":null,"work_id":"e6916b00-bf15-4819-b2dd-dbf453ab2d53","year":2014},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":56,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:606eb3a42e9ec04761705f12ae1bc914102696b1b5f317607624043e5b1ab296","observation_id":"afe9653d-60bb-4f74-85e7-f6c8c3823c7c","resolution":{"observed_at":"2026-05-17T20:25:12.522938Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Video mask transfiner for high-quality video instance segmentation","venue":null,"work_id":"2cae86f5-bb2c-499b-9da8-c6c028bce050","year":2022},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":57,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:33598b1f8aae1c100deff04457852e656b5491f9efe7415bc705ae2f46996d07","observation_id":"d2aced7b-0cc2-4614-a2da-b97f731320e4","resolution":{"observed_at":"2026-05-17T20:25:12.525138Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"442c5fa3-4e5c-421b-89f0-4a482e427b58","year":2024},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":58,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:b7c2b3e9b9c574c67296acaebdb1053b1b266784551587fceba4132206a1e72e","observation_id":"25413d31-9899-4168-8023-3c5d5f5d607f","resolution":{"observed_at":"2026-05-17T20:25:12.284584Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2408.12569","last_updated":"2024-08-27T02:31:42Z","snapshot_observed_at":"2026-07-06T19:04:43.716629Z","submitted_at":"2024-08-22T17:37:27Z","title":"Sapiens: Foundation for Human Vision Models","version":3},"cited_work":{"arxiv_id":"2408.12569","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2408.12569","snapshot_observed_at":"2026-07-03T14:48:32.980001Z","title":"arXiv preprint arXiv:2408.12569 , year=","venue":null,"work_id":"80c75e15-0fc5-43f3-836c-b196beae0c71","year":2024},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":59,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"cited_paper":"/paper/2408.12569","citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:bc1ac781e328953552e072b1004cecdb47388566ce03395bf73d0ff04be6feb0","observation_id":"099940c9-e8a8-43ea-b9ca-db73581b73a9","resolution":{"observed_at":"2026-05-17T20:25:11.397091Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Segment anything","venue":null,"work_id":"35bf84c4-84b4-4295-a6eb-b3949c3a4d93","year":2023},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":60,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:5a1dbf479c1a98cf441104c98dacb6afa20939f287be029e09546b8966c60d93","observation_id":"7842e0dd-ce7f-42c8-b965-14f56feb827e","resolution":{"observed_at":"2026-05-17T20:25:12.520673Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Visual genome: Connecting language and vision using crowdsourced dense image annotations","venue":null,"work_id":"dc2c2250-c626-4027-a67b-60e8f5f9c6d8","year":2017},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":61,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:f2a0ea159d4dd51a4a1b5679b5bd87e0e72f2d27ac93f42b6327062b57442d72","observation_id":"cd4630c6-c9e6-47db-95c6-8108338b864c","resolution":{"observed_at":"2026-05-17T20:25:12.488089Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"The open images dataset v4: Unified image classification, object detection, and visual relationship detection at scale","venue":null,"work_id":"6c068f3c-8c2e-4fae-b898-678c2e00f99a","year":1956},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":62,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:37985dab5f6b111c21e545c9ba49cf33981c9077c45d58f911ec5914b05f53d0","observation_id":"dd70c8d3-d410-47c9-9600-3f5d16fdfb0c","resolution":{"observed_at":"2026-05-17T20:25:12.485512Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1910.09700","last_updated":"2019-11-04T20:37:33Z","snapshot_observed_at":"2026-07-06T08:31:12.309726Z","submitted_at":"2019-10-21T23:57:32Z","title":"Quantifying the Carbon Emissions of Machine Learning","version":2},"cited_work":{"arxiv_id":"1910.09700","doi":"10.48550/arxiv.1910.09700","metadata_source":"pith","pith_arxiv_id":"1910.09700","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Quantifying the Carbon Emissions of Machine Learning","venue":"cs.CY","work_id":"7bc98d11-b344-40f2-b27f-2e08a08d1b95","year":2019},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":63,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"cited_paper":"/paper/1910.09700","citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:8b87278423d883f438733a30344125800620f08eb83910d30fe0b1f512f87f27","observation_id":"bec7a059-59f9-46ec-9ffd-f790b2dad35d","resolution":{"observed_at":"2026-05-17T20:25:11.533141Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Lisa: Reasoning segmentation via large language model","venue":null,"work_id":"5bcb5f91-e81c-4024-a945-4e080b9637b7","year":2024},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":64,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:3f4e0f48c1afdee4add32409d2ce3ccd826847684cc6c5ba6dda08d6e9394a68","observation_id":"fd775d3b-89cc-48b0-afa8-04208920e2cd","resolution":{"observed_at":"2026-05-17T20:25:12.480111Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"EDEN: Multimodal Synthetic Dataset of Enclosed garDEN Scenes","venue":null,"work_id":"70aea9a9-377d-4c14-bb82-464ad3465593","year":2021},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":65,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:59a9b886e263d2f17f5e8b7043276adaecc6519423f7ad450b6adb2315938c64","observation_id":"a196dd1d-ca4a-410a-847e-bb692f7d71ce","resolution":{"observed_at":"2026-05-17T20:25:12.483043Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Elevater: A benchmark and toolkit for evaluating language-augmented visual models","venue":null,"work_id":"838cb97b-312b-492f-a4c6-3caac76e4fa6","year":2022},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":66,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:4a00540a2f91e0eb3b8518160e5dd6856fdb6b706a64f52aeb1fc32ba61ea47f","observation_id":"88a82047-d6d1-4d2d-89d2-a54958474c81","resolution":{"observed_at":"2026-05-17T20:25:12.490448Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Visual in-context prompting","venue":null,"work_id":"9b4f218f-2fca-4581-8a99-9cde85fa24a3","year":2024},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":67,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:d539ee0c3a6e169a965cbd764bc1bf172ad89f983aa81246b720e13feedfb131","observation_id":"1b54f151-db12-45d8-8211-1d4fd855b7fa","resolution":{"observed_at":"2026-05-17T20:25:12.456912Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.14467","last_updated":"2025-05-01T14:14:05Z","snapshot_observed_at":"2026-07-06T21:12:04.001107Z","submitted_at":"2025-04-20T02:51:11Z","title":"LGD: Leveraging Generative Descriptions for Zero-Shot Referring Image Segmentation","version":2},"cited_work":{"arxiv_id":"2504.14467","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2504.14467","snapshot_observed_at":"2026-06-29T12:53:26.682105Z","title":"Lgd: Leveraging generative descriptions for zero-shot referring image segmentation","venue":null,"work_id":"af5e1988-49c2-4857-afd2-0fef7a97f399","year":2025},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":68,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"cited_paper":"/paper/2504.14467","citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:23417321f555e39f0047b5958b732b36be599f49a005e3893dea873aa17153fc","observation_id":"ad5b1c4f-3d5a-458f-9e1b-771cd5d2e131","resolution":{"observed_at":"2026-05-17T20:25:11.512700Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2301.12597","last_updated":"2023-06-15T07:57:29Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-01-30T00:56:51Z","title":"BLIP-2: Bootstrapping Language-Image Pre-training with Frozen Image Encoders and Large Language Models","version":3},"cited_work":{"arxiv_id":"2301.12597","doi":"10.48550/arxiv.2301.12597","metadata_source":"pith","pith_arxiv_id":"2301.12597","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"BLIP-2: Bootstrapping Language-Image Pre-training with Frozen Image Encoders and Large Language Models","venue":"cs.CV","work_id":"63d03f4d-15f4-4583-8286-913c19f02294","year":2023},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":69,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"cited_paper":"/paper/2301.12597","citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:9a3fbc71a49c340377014f9c4f8d303d8d69bd4ee04d6363912a26cf96f4031b","observation_id":"43f652ba-fd9c-42eb-b0b1-6c1f6b2836fb","resolution":{"observed_at":"2026-05-17T20:25:11.526629Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-12T21:49:47.354893+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T21:49:47.354893+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Desco: Learning object recognition with rich language descriptions","venue":null,"work_id":"e7973800-97f6-470d-ba2a-d0d6b814b95d","year":2023},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":70,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:9541587e1c7e0cf9e1c6af875433bf9dd2fdf47b7e7088bcd38139c84ecd96a4","observation_id":"be2e9b4c-5ab2-4862-83c4-5982e0401363","resolution":{"observed_at":"2026-05-17T20:25:12.444256Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Grounded language-image pre-training","venue":null,"work_id":"7b037c9a-5645-4489-a39a-8f3b6f7a9e20","year":2022},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":71,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:d2482e4f711ac20dd34723f0fafad25bbaaf6b7f8a0d252b954645d0f80d591d","observation_id":"c4aa71e5-931f-4dce-8cc4-07aa9d1e97c4","resolution":{"observed_at":"2026-05-17T20:25:12.448988Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Tracking every thing in the wild","venue":null,"work_id":"a6dfbb18-895c-47e0-8d32-cf9b2f4628d1","year":2022},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":72,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:7a122e6baa3abe5cc3e53e128e00393e3cc743cdd6714de7dedd91dd14e81c81","observation_id":"fb425089-a84b-43ad-93e7-26770477cf52","resolution":{"observed_at":"2026-05-17T20:25:12.451329Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Exploring plain vision transformer backbones for object detection","venue":null,"work_id":"7484a1cc-1a77-44a5-8d77-b26a17e973f8","year":2022},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":73,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:558d1da109a49214015afe181ed3f69e55dd98d3569818624427e4d396d13e45","observation_id":"0145bc80-2102-4375-b791-279d6e4253de","resolution":{"observed_at":"2026-05-17T20:25:12.462655Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Open-vocabulary semantic segmentation with mask-adapted clip","venue":null,"work_id":"035d3fea-7b96-489c-ac5c-044cf24117c0","year":2023},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":74,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:6b882c8647bebc8d0d68eb6558d01dc12981714dfe912980b0b74130e35e1dd3","observation_id":"54b4156a-7465-4598-bb8c-91ea648ed839","resolution":{"observed_at":"2026-05-17T20:25:12.498861Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"WCS camera traps","venue":null,"work_id":"59351513-59f4-4c88-8d63-8011e2b3a148","year":null},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":75,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:297794dfa094610d3d6e877719ab8afe0b5e02c6b15019c0e7e94d96f0943143","observation_id":"fd2cbe59-9dbc-4e9b-aa1b-ad523c3a66b4","resolution":{"observed_at":"2026-05-17T20:25:12.515829Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-07T20:44:09.209541Z","title":"Microsoft coco: Common objects in context","venue":null,"work_id":"e38750ad-2d86-48f0-ba45-b8dcc36dbd20","year":2014},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":76,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:3ca29f134145332134636c1da2d009c21ce0597b04b67f317e95c90ce3dab54c","observation_id":"45f56976-187f-4dcf-8e61-18c0d1d2bafe","resolution":{"observed_at":"2026-05-17T20:25:12.434784Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Detr doesn't need multi-scale or locality design","venue":null,"work_id":"acf1cf07-def2-4785-a18a-2c9dde341ee3","year":2023},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":77,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:f1d7c29fae1822594563464bbb8073ff71d7e65e92d242454e141b97f1c6eb41","observation_id":"106a1d16-3f00-4810-a4a6-a2d6ff0c8cc9","resolution":{"observed_at":"2026-05-17T20:25:12.425347Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Grounding dino: Marrying dino with grounded pre-training for open-set object detection","venue":null,"work_id":"e1283563-079b-4417-a0cb-46874e82dce9","year":2023},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":78,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:ca382795503cdbff4a123c9b9de9692503ac360d6ef862cda39854b3e4873ab2","observation_id":"2c219341-cf14-4b04-8196-3c456d3d4c62","resolution":{"observed_at":"2026-05-17T20:25:12.427753Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Grounding dino: Marrying dino with grounded pre-training for open-set object detection","venue":null,"work_id":"2b142fe0-a6ab-471e-b932-09fe7d857812","year":2024},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":79,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:c84ae6dfc9736370afc96879e234a5a26c235db6f1d61d6f1d27d6b3f4470a93","observation_id":"0ff77ce1-269d-44a0-8e0a-388ed9268477","resolution":{"observed_at":"2026-05-17T20:25:12.430218Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Hybrid global-local representation with augmented spatial guidance for zero-shot referring image segmentation","venue":null,"work_id":"d3128924-c8a2-46e5-9f1e-f158b12028f2","year":2025},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":80,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:1a7abfe0be338803e4f885c8bbbdaf7d400268ab28f502ea17559081877d1777","observation_id":"bed37865-5700-40ed-b380-841760ae1b3a","resolution":{"observed_at":"2026-05-17T20:25:12.418584Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Universal segmentation at arbitrary granularity with language instruction","venue":null,"work_id":"ba9a7eb6-a8fb-43b4-b455-3e3627986b76","year":2024},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":81,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:0272f40339b1a2bd82aaadace049515bf73aa587ad9ea0f7658cd4e60f010ee9","observation_id":"53dd0bad-7b08-4b4a-8131-9458ff9d9389","resolution":{"observed_at":"2026-05-17T20:25:12.420702Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2503.06520","last_updated":"2026-05-31T09:29:40Z","snapshot_observed_at":"2026-08-05T01:42:30.194131Z","submitted_at":"2025-03-09T08:48:51Z","title":"Seg-Zero: Reasoning-Chain Guided Segmentation via Cognitive Reinforcement","version":3},"cited_work":{"arxiv_id":"2503.06520","doi":null,"metadata_source":"pith","pith_arxiv_id":"2503.06520","snapshot_observed_at":"2026-07-04T19:30:08.050016Z","title":"Seg-Zero: Reasoning-Chain Guided Segmentation via Cognitive Reinforcement","venue":"cs.CV","work_id":"7a33b2a4-8409-4a00-ad7b-84e810df1ba7","year":2025},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":82,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"cited_paper":"/paper/2503.06520","citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:4ed7a761487ed36e15bfcf68f20d46fb84778001a4119bbba5f522e48610270d","observation_id":"59c3125b-5bdd-4e2f-8258-eadd370c475c","resolution":{"observed_at":"2026-05-17T20:25:11.470281Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Towards end-to-end unified scene text detection and layout analysis","venue":null,"work_id":"9f800bdc-8981-4818-b449-23aa06535104","year":2022},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":83,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:19ec07579e8e6566cf9bd2bc7df4030959c36f779929fc23e84bf7fe83f04d8f","observation_id":"8b3f0a00-242b-4a55-a321-d35808c6a216","resolution":{"observed_at":"2026-05-17T20:25:12.423179Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.09750","last_updated":"2023-05-16T18:56:12Z","snapshot_observed_at":"2026-07-06T15:28:20.821441Z","submitted_at":"2023-05-16T18:56:12Z","title":"ICDAR 2023 Competition on Hierarchical Text Detection and Recognition","version":1},"cited_work":{"arxiv_id":"2305.09750","doi":null,"metadata_source":"pith","pith_arxiv_id":"2305.09750","snapshot_observed_at":"2026-07-08T02:04:26.442940Z","title":"Icdar 2023 competition on hierarchical text detection and recognition","venue":"cs.CV","work_id":"186f9df4-a008-4848-81d5-3c9d261c0cf0","year":2023},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":84,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"cited_paper":"/paper/2305.09750","citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:79946c7ae4f37bd7e5616190d739448bdbe87fd712f3d8a78e91da7d381e2d87","observation_id":"2ac156fb-0aa7-456a-af71-47e6cbe6edaf","resolution":{"observed_at":"2026-05-17T20:25:11.489087Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Decoupled weight decay regularization","venue":null,"work_id":"0e14f3c3-0bfe-48ab-9a20-6e513d32976c","year":2019},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":85,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:2842476269c4fd5ee7ffdb270824c791c28771ba34434c0809bd0ce290e6ef01","observation_id":"6a6cd9ba-f971-4c22-8148-dac8d729c35f","resolution":{"observed_at":"2026-05-17T20:25:12.439532Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2506.04277","last_updated":"2025-06-04T02:07:40Z","snapshot_observed_at":"2026-07-06T21:36:45.605379Z","submitted_at":"2025-06-04T02:07:40Z","title":"RSVP: Reasoning Segmentation via Visual Prompting and Multi-modal Chain-of-Thought","version":1},"cited_work":{"arxiv_id":"2506.04277","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2506.04277","snapshot_observed_at":"2026-07-03T00:47:30.631926Z","title":"Rsvp: Reasoning segmentation via visual prompting and multi-modal chain-of-thought","venue":null,"work_id":"026d3dfa-6287-4f2b-833c-eeb4fd2944df","year":2025},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":86,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"cited_paper":"/paper/2506.04277","citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:5e815ff9312b09770eacbc94ddec475d77cfa283d195ec0eb5da9aa9df1d9227","observation_id":"0ffd6ff2-f54f-4d8a-9d37-5055ffd9f570","resolution":{"observed_at":"2026-05-17T20:25:11.519580Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Hota: A higher order metric for evaluating multi-object tracking","venue":null,"work_id":"a444ae81-ab58-40bf-a612-73a1cd70e038","year":2021},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":87,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:ab7cfa2ecb4afcf84b4991597fb9189416917f0b4ab36dfa4488843852cf9138","observation_id":"fa41eb51-36bd-4fc8-b3c8-35912ed9ffb5","resolution":{"observed_at":"2026-05-17T20:25:12.409327Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.03600","last_updated":"2025-04-04T17:13:37Z","snapshot_observed_at":"2026-08-04T20:04:58.153270Z","submitted_at":"2025-04-04T17:13:37Z","title":"MedSAM2: Segment Anything in 3D Medical Images and Videos","version":1},"cited_work":{"arxiv_id":"2504.03600","doi":null,"metadata_source":"pith","pith_arxiv_id":"2504.03600","snapshot_observed_at":"2026-07-08T20:15:34.383554Z","title":"Medsam2: Segment anything in 3d medical images and videos","venue":"eess.IV","work_id":"4be9151e-38db-47e1-a183-b6e11b538112","year":2025},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":88,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"cited_paper":"/paper/2504.03600","citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:abe99439d1f93e945b542ba842fe8a076d45d2723422856fb962cdaf7dfba7fe","observation_id":"aee0cbd5-fbcb-41ad-8e80-5b04981a4de0","resolution":{"observed_at":"2026-05-17T20:25:11.441824Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Generation and comprehension of unambiguous object descriptions","venue":null,"work_id":"c7a4889a-285e-44b1-ad95-d2079ff7812f","year":2016},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":89,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:18c052aa39a7b4aec44cf0c1c20d743f2b21a7ecf1eba100a6d8dfafb81ce48e","observation_id":"500faeb0-3c71-46c2-ba6b-7102584321af","resolution":{"observed_at":"2026-05-17T20:25:12.411682Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2307.12730","last_updated":"2023-08-02T12:10:55Z","snapshot_observed_at":"2026-08-02T10:36:55.414042Z","submitted_at":"2023-07-24T12:22:19Z","title":"COCO-O: A Benchmark for Object Detectors under Natural Distribution Shifts","version":2},"cited_work":{"arxiv_id":"2307.12730","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2307.12730","snapshot_observed_at":"2026-07-04T12:59:53.245515Z","title":"Coco-o: A benchmark for object detectors under natural distribution shifts","venue":null,"work_id":"4a3172d2-aadf-4847-a8ca-73b35448fc08","year":2023},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":90,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"cited_paper":"/paper/2307.12730","citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:d36bb14a0d011bee8630e57ef35531078bf1df43e60b53f837b698061e3d8d4a","observation_id":"0a8a2c3c-2ab9-4f66-b15d-08f1e1a34b5e","resolution":{"observed_at":"2026-05-17T20:25:11.515971Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Trackformer: Multi-object tracking with transformers","venue":null,"work_id":"5050d38a-2e7d-4f7b-85f9-d2018ba8fc3a","year":2022},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":91,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:09f4906c637d903959f9d2650b508d01c93039833d57d4923129b5fbde64be08","observation_id":"6fb45219-baf8-4d9a-96d0-acf12d4efc07","resolution":{"observed_at":"2026-05-17T20:25:12.406571Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Simple open-vocabulary object detection","venue":null,"work_id":"f310a008-a51c-402c-b44d-88b8aaa29669","year":2022},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":92,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:dfeffaaf49b90df37f6a3363ef51316b7c0ddfa6dacbe463234ffae23e866d9f","observation_id":"d29dc5ac-1351-4538-bded-c3a70212d3d0","resolution":{"observed_at":"2026-05-17T20:25:12.413872Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.09683","last_updated":"2024-05-22T13:00:02Z","snapshot_observed_at":"2026-08-04T22:40:33.196944Z","submitted_at":"2023-06-16T08:27:46Z","title":"Scaling Open-Vocabulary Object Detection","version":3},"cited_work":{"arxiv_id":"2306.09683","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2306.09683","snapshot_observed_at":"2026-07-03T13:58:21.272558Z","title":"Scaling open-vocabulary object detection","venue":null,"work_id":"b52ffc8f-1185-4def-828a-f9e6b31eeda2","year":2024},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":93,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"cited_paper":"/paper/2306.09683","citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:125e4e3a245c17830a5fb8e92cf314aea65d6fd2e692e5aa1449c8fecf99f869","observation_id":"75978a72-45bf-4ecc-bbee-a5079baaf29a","resolution":{"observed_at":"2026-05-17T20:25:11.477247Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2303.16382","last_updated":"2023-03-29T01:42:54Z","snapshot_observed_at":"2026-07-06T15:09:21.780147Z","submitted_at":"2023-03-29T01:42:54Z","title":"ARMBench: An Object-centric Benchmark Dataset for Robotic Manipulation","version":1},"cited_work":{"arxiv_id":"2303.16382","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2303.16382","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Armbench: An object-centric benchmark dataset for robotic manipulation","venue":null,"work_id":"a95c81e1-f107-4ccb-9097-1a53daa66265","year":2023},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":94,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"cited_paper":"/paper/2303.16382","citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:cb7e24845a094afe1f61ddaea0d5fb4324bf0020cac80f58c0e7d0dc01b5a27f","observation_id":"7ce729b6-c3a1-4872-a66e-9dbb725684ed","resolution":{"observed_at":"2026-05-17T20:25:11.508856Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Model cards for model reporting","venue":null,"work_id":"bfa40a63-d40a-4ee9-9d82-5f8c0eefa8b8","year":2019},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":95,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:e913e61f694f5e59e1511ba206ed7ef6a730e9c83f2610c2c9a85e16569e9731","observation_id":"b09c948e-7c6a-4697-b790-99918bc65632","resolution":{"observed_at":"2026-05-17T20:25:12.518406Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2106.14977","last_updated":"2021-06-30T10:05:21Z","snapshot_observed_at":"2026-07-06T11:23:56.887244Z","submitted_at":"2021-06-28T20:51:26Z","title":"The Food Recognition Benchmark: Using DeepLearning to Recognize Food on Images","version":2},"cited_work":{"arxiv_id":"2106.14977","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2106.14977","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"The food recognition benchmark: Using deeplearning to recognize food on images","venue":null,"work_id":"2c2d6d6d-766f-4652-b761-7099f2552efd","year":2021},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":96,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"cited_paper":"/paper/2106.14977","citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:62a1befc1ca0f24e9edacb31c88c9ed51a42af46196271f9d4357d91c43d32a4","observation_id":"d4417618-3315-4d4f-8d99-70b93171f163","resolution":{"observed_at":"2026-05-17T20:25:11.544092Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"The role of context for object detection and semantic segmentation in the wild","venue":null,"work_id":"b840e53b-3c28-41ba-bfa2-5ca17f4f59f8","year":2014},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":97,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:88cc95c682ebab7bb5194eef52369122336f2efb5330a9f4710718ea1b467c0f","observation_id":"c0817322-7fdd-458c-b6ed-2bf12fabc242","resolution":{"observed_at":"2026-05-17T20:25:12.534751Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Public domain collection dataset","venue":null,"work_id":"be11da4d-d135-48b7-bc73-cf04a8c04a43","year":null},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":98,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:11cbeee5b8feba2748afc23b09d6075f7759dbd6165a3cafe1b1eeffc9fde524","observation_id":"889241bd-e00f-4ac2-8199-9202dba4a043","resolution":{"observed_at":"2026-05-17T20:25:12.432593Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2308.16777","last_updated":"2023-09-01T05:57:47Z","snapshot_observed_at":"2026-07-06T16:12:51.303563Z","submitted_at":"2023-08-31T14:55:30Z","title":"Ref-Diff: Zero-shot Referring Image Segmentation with Generative Models","version":2},"cited_work":{"arxiv_id":"2308.16777","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2308.16777","snapshot_observed_at":"2026-07-04T08:59:42.789854Z","title":"Ref-diff: Zero-shot referring image segmentation with generative models","venue":null,"work_id":"afd76357-8469-4e1b-b865-eb6b8b66646a","year":2023},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":99,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"cited_paper":"/paper/2308.16777","citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:e4c68cfa9271a6a5d52c2f9fcfaeaa7d52eb5c19d8cab8d5672d2f8384212d80","observation_id":"17ab4453-da85-4878-8c7a-1465cec38302","resolution":{"observed_at":"2026-05-17T20:25:11.540542Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Dinov2: Learning robust visual features without supervision","venue":null,"work_id":"711e2338-a0bc-4fd9-ba5a-35a1958ce9c5","year":2024},"citing_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"reference_index":100,"source":"arxiv_source","source_observed_at":"2026-05-17T20:22:46.220021Z"},"links":{"citing_paper":"/paper/2511.16719"},"observation_digest":"sha256:1126ebb51278783de84009d68f76de7210bf902f7da66095996adebe37110237","observation_id":"ee35e022-09f1-4f64-9675-c87e707ec7ff","resolution":{"observed_at":"2026-05-17T20:25:12.495940Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","latest_version":2,"primary_category":"cs.CV","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts"},"reference_resolution":{"displayed":100,"state_counts":{"malformed_identifier":0,"metadata_mismatch":3,"parse_uncertain":0,"unresolved":3,"verified_exact":32,"verified_fuzzy":62},"total_outbound_references":168},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"thesis":"As of 5 August 2026, this Paper Citation Record lists 100 of 168 outbound references and 100 inbound Pith citation observations for arXiv:2511.16719."}