{"as_of":"2026-08-23T00:04:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:649edbc23cbcc7702b4dfb6ac04b4ef531ab832402b078f8a1056221d82a5124","coverage":[{"denominator":16,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":16,"source":"paper_references, paper_reference_links","source_observed_at":"2026-07-10T05:41:42.555295Z","state":"measured"},{"denominator":16,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":16,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-22T06:32:14.747728+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2607.08541/citation-record","integrity":"/paper/2607.08541/integrity","json":"/paper/2607.08541/citation-record.json","paper":"/paper/2607.08541"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"1804.02767","last_updated":"2018-04-08T22:27:57Z","snapshot_observed_at":"2026-08-15T16:54:48.034047Z","submitted_at":"2018-04-08T22:27:57Z","title":"YOLOv3: An Incremental Improvement","version":1},"cited_work":{"arxiv_id":"1804.02767","doi":"10.48550/arxiv.1804.02767","metadata_source":"pith","pith_arxiv_id":"1804.02767","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"YOLOv3: An Incremental Improvement","venue":"cs.CV","work_id":"d737b3cc-9bd1-43d6-8310-c91e64b510f7","year":2018},"citing_paper":{"arxiv_id":"2607.08541","last_updated":"2026-07-09T14:35:37Z","snapshot_observed_at":"2026-08-19T01:14:44.888182Z","submitted_at":"2026-07-09T14:35:37Z","title":"VocaDet: Sample-Driven Open-Vocabulary Object Detection and Segmentation via Visual Tokenization and Vector Database Retrieval","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-07-10T05:41:42.555295Z"},"links":{"cited_paper":"/paper/1804.02767","citing_paper":"/paper/2607.08541"},"observation_digest":"sha256:c46a2e2f5f4009f1162faa5a60f683558cbe71c1292061306e083b974bfae133","observation_id":"70556889-5243-4bd0-b9a7-6ea04e5e479f","resolution":{"observed_at":"2026-07-10T05:46:50.372021Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-07-12T23:51:10.097273+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T23:51:10.097273+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-10T06:26:52.342606Z","title":"Yolov10: Real-time end-to-end object detection.Advances in neural information processing systems, 37:107984–108011, 2024","venue":null,"work_id":"14c8f9cc-fe3e-435d-8d3d-526e769033db","year":2024},"citing_paper":{"arxiv_id":"2607.08541","last_updated":"2026-07-09T14:35:37Z","snapshot_observed_at":"2026-08-19T01:14:44.888182Z","submitted_at":"2026-07-09T14:35:37Z","title":"VocaDet: Sample-Driven Open-Vocabulary Object Detection and Segmentation via Visual Tokenization and Vector Database Retrieval","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-07-10T05:41:42.555295Z"},"links":{"citing_paper":"/paper/2607.08541"},"observation_digest":"sha256:4268c84d0dd8e8e3cb723a5ecc312ec502b4b3b117ca6f8a19fba6c5dab82352","observation_id":"c2f38646-a600-48c4-8071-40d22107e002","resolution":{"observed_at":"2026-07-10T06:26:52.344134Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.16719","last_updated":"2026-03-28T16:54:56Z","snapshot_observed_at":"2026-08-07T05:59:39.049027Z","submitted_at":"2025-11-20T18:59:56Z","title":"SAM 3: Segment Anything with Concepts","version":2},"cited_work":{"arxiv_id":"2511.16719","doi":"10.48550/arxiv.2511.16719","metadata_source":"pith","pith_arxiv_id":"2511.16719","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SAM 3: Segment Anything with Concepts","venue":"cs.CV","work_id":"4a72a006-2592-4554-aad0-a9c41a9f952d","year":2025},"citing_paper":{"arxiv_id":"2607.08541","last_updated":"2026-07-09T14:35:37Z","snapshot_observed_at":"2026-08-19T01:14:44.888182Z","submitted_at":"2026-07-09T14:35:37Z","title":"VocaDet: Sample-Driven Open-Vocabulary Object Detection and Segmentation via Visual Tokenization and Vector Database Retrieval","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-07-10T05:41:42.555295Z"},"links":{"cited_paper":"/paper/2511.16719","citing_paper":"/paper/2607.08541"},"observation_digest":"sha256:50eb7f13cde4e51cf14e5a86f0e75245f3d0747d35165a749c24aa8a94e81ebd","observation_id":"2583a07e-e757-4c97-af57-add0354006c5","resolution":{"observed_at":"2026-07-10T05:46:50.380835Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T15:53:24.687473+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-10T06:26:52.337971Z","title":"Grounding dino: Marrying dino with grounded pre-training for open-set object detection","venue":null,"work_id":"98d774c8-2295-4dc2-82fb-624467a40d12","year":2024},"citing_paper":{"arxiv_id":"2607.08541","last_updated":"2026-07-09T14:35:37Z","snapshot_observed_at":"2026-08-19T01:14:44.888182Z","submitted_at":"2026-07-09T14:35:37Z","title":"VocaDet: Sample-Driven Open-Vocabulary Object Detection and Segmentation via Visual Tokenization and Vector Database Retrieval","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-07-10T05:41:42.555295Z"},"links":{"citing_paper":"/paper/2607.08541"},"observation_digest":"sha256:e32b3e270e614dbc759ca11c8241492756daa2096c1cf10aa19fa5674786e7d8","observation_id":"af2f0723-60e8-44d3-854d-0c8de1008935","resolution":{"observed_at":"2026-07-10T06:26:52.339534Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2403.14610","last_updated":"2024-03-21T17:57:03Z","snapshot_observed_at":"2026-08-16T14:07:35.127518Z","submitted_at":"2024-03-21T17:57:03Z","title":"T-Rex2: Towards Generic Object Detection via Text-Visual Prompt Synergy","version":1},"cited_work":{"arxiv_id":"2403.14610","doi":null,"metadata_source":"pith","pith_arxiv_id":"2403.14610","snapshot_observed_at":"2026-07-10T05:46:50.376442Z","title":"3D Avg.” refers to the average success rate over all 3D-printed objects, “DO Avg","venue":"cs.CV","work_id":"a7d5cc79-51ba-4f78-9917-6114664a5d36","year":2024},"citing_paper":{"arxiv_id":"2607.08541","last_updated":"2026-07-09T14:35:37Z","snapshot_observed_at":"2026-08-19T01:14:44.888182Z","submitted_at":"2026-07-09T14:35:37Z","title":"VocaDet: Sample-Driven Open-Vocabulary Object Detection and Segmentation via Visual Tokenization and Vector Database Retrieval","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-07-10T05:41:42.555295Z"},"links":{"cited_paper":"/paper/2403.14610","citing_paper":"/paper/2607.08541"},"observation_digest":"sha256:8b7f5f7f773e8e6cf7d8d23d589b1195e5b0a75ee87c90a485253f4351eb7e42","observation_id":"6b4d0b50-11d3-4e14-bb79-ecf2686a2c74","resolution":{"observed_at":"2026-07-10T05:46:50.378120Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2411.14347","last_updated":"2025-05-15T15:52:39Z","snapshot_observed_at":"2026-08-19T22:36:34.839113Z","submitted_at":"2024-11-21T17:42:20Z","title":"DINO-X: A Unified Vision Model for Open-World Object Detection and Understanding","version":3},"cited_work":{"arxiv_id":"2411.14347","doi":null,"metadata_source":"pith","pith_arxiv_id":"2411.14347","snapshot_observed_at":"2026-07-10T05:46:50.368257Z","title":"org/abs/2411.14347","venue":"cs.CV","work_id":"73ad752c-78bb-428f-b88d-6281501d8f82","year":2024},"citing_paper":{"arxiv_id":"2607.08541","last_updated":"2026-07-09T14:35:37Z","snapshot_observed_at":"2026-08-19T01:14:44.888182Z","submitted_at":"2026-07-09T14:35:37Z","title":"VocaDet: Sample-Driven Open-Vocabulary Object Detection and Segmentation via Visual Tokenization and Vector Database Retrieval","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-07-10T05:41:42.555295Z"},"links":{"cited_paper":"/paper/2411.14347","citing_paper":"/paper/2607.08541"},"observation_digest":"sha256:73e8c8d8fa3726ce0301b651956503aa27b8a95bdab11db59496e73a8bc2c881","observation_id":"8a14c43d-d37d-4003-b40f-780a694c0f02","resolution":{"observed_at":"2026-07-10T05:46:50.369655Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-10T06:26:52.340075Z","title":"Insid3: Training-free in-context segmentation with dinov3","venue":null,"work_id":"be2609ff-3558-4c7c-a180-bce494620a81","year":2026},"citing_paper":{"arxiv_id":"2607.08541","last_updated":"2026-07-09T14:35:37Z","snapshot_observed_at":"2026-08-19T01:14:44.888182Z","submitted_at":"2026-07-09T14:35:37Z","title":"VocaDet: Sample-Driven Open-Vocabulary Object Detection and Segmentation via Visual Tokenization and Vector Database Retrieval","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-07-10T05:41:42.555295Z"},"links":{"citing_paper":"/paper/2607.08541"},"observation_digest":"sha256:fe798161b8b9b61afd2959354e2262693e2d389b3381203f0b16752a161878a5","observation_id":"c32ae220-8ec5-4796-928a-87cf4650dda0","resolution":{"observed_at":"2026-07-10T06:26:52.341566Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2506.04034","last_updated":"2025-06-04T14:56:57Z","snapshot_observed_at":"2026-08-16T17:02:27.821392Z","submitted_at":"2025-06-04T14:56:57Z","title":"Rex-Thinker: Grounded Object Referring via Chain-of-Thought Reasoning","version":1},"cited_work":{"arxiv_id":"2506.04034","doi":null,"metadata_source":"pith","pith_arxiv_id":"2506.04034","snapshot_observed_at":"2026-07-10T05:46:50.382039Z","title":"Rex-thinker: Grounded object re- ferring via chain-of-thought reasoning","venue":"cs.CV","work_id":"8f7606c1-62bd-42f7-8c1d-f545ea62a1d6","year":2025},"citing_paper":{"arxiv_id":"2607.08541","last_updated":"2026-07-09T14:35:37Z","snapshot_observed_at":"2026-08-19T01:14:44.888182Z","submitted_at":"2026-07-09T14:35:37Z","title":"VocaDet: Sample-Driven Open-Vocabulary Object Detection and Segmentation via Visual Tokenization and Vector Database Retrieval","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-07-10T05:41:42.555295Z"},"links":{"cited_paper":"/paper/2506.04034","citing_paper":"/paper/2607.08541"},"observation_digest":"sha256:8aef2985390216f914a0495c70cebf0000b59b019f1e025e3b9153a65811990c","observation_id":"ea19ca4f-672a-454e-9bb0-572256233ba7","resolution":{"observed_at":"2026-07-10T05:46:50.383223Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-10T06:26:52.344770Z","title":"Detect anything via next point prediction","venue":null,"work_id":"4034011b-59b2-47cd-a39a-27afea1a2a4d","year":2026},"citing_paper":{"arxiv_id":"2607.08541","last_updated":"2026-07-09T14:35:37Z","snapshot_observed_at":"2026-08-19T01:14:44.888182Z","submitted_at":"2026-07-09T14:35:37Z","title":"VocaDet: Sample-Driven Open-Vocabulary Object Detection and Segmentation via Visual Tokenization and Vector Database Retrieval","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-07-10T05:41:42.555295Z"},"links":{"citing_paper":"/paper/2607.08541"},"observation_digest":"sha256:6b704bca11b11a93e0566afe0281e33d393743b3f30203b320c4017d3b347fc4","observation_id":"3e1a2fad-c6d4-4523-bce3-f748641cebac","resolution":{"observed_at":"2026-07-10T06:26:52.346201Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2507.02798","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-10T05:46:50.365484Z","title":"No time to train! training-free reference-based instance segmentation.arXiv preprint arXiv:2507.02798, 2025","venue":null,"work_id":"cb1866ed-b7f5-4edb-883c-211f52d6c2e8","year":2025},"citing_paper":{"arxiv_id":"2607.08541","last_updated":"2026-07-09T14:35:37Z","snapshot_observed_at":"2026-08-19T01:14:44.888182Z","submitted_at":"2026-07-09T14:35:37Z","title":"VocaDet: Sample-Driven Open-Vocabulary Object Detection and Segmentation via Visual Tokenization and Vector Database Retrieval","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-07-10T05:41:42.555295Z"},"links":{"citing_paper":"/paper/2607.08541"},"observation_digest":"sha256:0820223f9b8df456af458c91d752098d6caff9d17f066d73cc67be469c7733e0","observation_id":"3c950d6e-caa4-4295-b9a0-cc4aa3a00e1c","resolution":{"observed_at":"2026-07-10T05:46:50.367048Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2304.07193","last_updated":"2024-02-02T10:24:09Z","snapshot_observed_at":"2026-08-17T13:03:40.359628Z","submitted_at":"2023-04-14T15:12:19Z","title":"DINOv2: Learning Robust Visual Features without Supervision","version":2},"cited_work":{"arxiv_id":"2304.07193","doi":"10.48550/arxiv.2304.07193","metadata_source":"pith","pith_arxiv_id":"2304.07193","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"DINOv2: Learning Robust Visual Features without Supervision","venue":"cs.CV","work_id":"26b304e5-b54a-4f26-be7e-83299eca52e4","year":2023},"citing_paper":{"arxiv_id":"2607.08541","last_updated":"2026-07-09T14:35:37Z","snapshot_observed_at":"2026-08-19T01:14:44.888182Z","submitted_at":"2026-07-09T14:35:37Z","title":"VocaDet: Sample-Driven Open-Vocabulary Object Detection and Segmentation via Visual Tokenization and Vector Database Retrieval","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-07-10T05:41:42.555295Z"},"links":{"cited_paper":"/paper/2304.07193","citing_paper":"/paper/2607.08541"},"observation_digest":"sha256:e2db3f3cb4bb3c370670a451d1b97affdaf5bfbc53af002ea7a7dcaa8455af3a","observation_id":"e70a137f-65c9-4817-a58a-946517171727","resolution":{"observed_at":"2026-07-10T05:46:50.364243Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-10T06:26:52.333552Z","title":"Sam 2: Segment anything in images and videos","venue":null,"work_id":"ead3a57d-dc37-491e-bdc2-fce6b2fc743a","year":2025},"citing_paper":{"arxiv_id":"2607.08541","last_updated":"2026-07-09T14:35:37Z","snapshot_observed_at":"2026-08-19T01:14:44.888182Z","submitted_at":"2026-07-09T14:35:37Z","title":"VocaDet: Sample-Driven Open-Vocabulary Object Detection and Segmentation via Visual Tokenization and Vector Database Retrieval","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-07-10T05:41:42.555295Z"},"links":{"citing_paper":"/paper/2607.08541"},"observation_digest":"sha256:67e2b1640a963aff6730f058378b26fbd084bc055b194eab482ee054061244fc","observation_id":"9f901118-bf2b-4062-8848-ce82f7b4bb78","resolution":{"observed_at":"2026-07-10T06:26:52.335228Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2508.10104","last_updated":"2025-08-13T18:00:55Z","snapshot_observed_at":"2026-08-18T01:07:23.737664Z","submitted_at":"2025-08-13T18:00:55Z","title":"DINOv3","version":1},"cited_work":{"arxiv_id":"2508.10104","doi":"10.1055/a-2487-1252","metadata_source":"pith","pith_arxiv_id":"2508.10104","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"DINOv3","venue":"cs.CV","work_id":"c8b07deb-8fe7-4e18-9620-f3569d3529ce","year":2025},"citing_paper":{"arxiv_id":"2607.08541","last_updated":"2026-07-09T14:35:37Z","snapshot_observed_at":"2026-08-19T01:14:44.888182Z","submitted_at":"2026-07-09T14:35:37Z","title":"VocaDet: Sample-Driven Open-Vocabulary Object Detection and Segmentation via Visual Tokenization and Vector Database Retrieval","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-07-10T05:41:42.555295Z"},"links":{"cited_paper":"/paper/2508.10104","citing_paper":"/paper/2607.08541"},"observation_digest":"sha256:b4551a221f8d8938780ee1937a5e7f36e7be525eadbdae0a9b7588bc05a7b92f","observation_id":"1691cce1-3c41-4e75-9338-f357937d72ab","resolution":{"observed_at":"2026-07-10T05:46:50.361737Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1109.2378","last_updated":"2011-09-12T05:49:11Z","snapshot_observed_at":"2026-08-15T21:00:50.761660Z","submitted_at":"2011-09-12T05:49:11Z","title":"Modern hierarchical, agglomerative clustering algorithms","version":1},"cited_work":{"arxiv_id":"1109.2378","doi":"10.48550/arxiv.1109.2378","metadata_source":"pith","pith_arxiv_id":"1109.2378","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Modern hierarchical, agglomerative clustering algorithms","venue":"stat.ML","work_id":"531ad2b9-074c-4458-8025-7c263ce6f6b2","year":2011},"citing_paper":{"arxiv_id":"2607.08541","last_updated":"2026-07-09T14:35:37Z","snapshot_observed_at":"2026-08-19T01:14:44.888182Z","submitted_at":"2026-07-09T14:35:37Z","title":"VocaDet: Sample-Driven Open-Vocabulary Object Detection and Segmentation via Visual Tokenization and Vector Database Retrieval","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-07-10T05:41:42.555295Z"},"links":{"cited_paper":"/paper/1109.2378","citing_paper":"/paper/2607.08541"},"observation_digest":"sha256:e1f2c74303f2c7ac508cbb2f8a7b8ad919a91174265fd4a15dc60d7dc25d5c19","observation_id":"a030e303-da83-4a06-808c-f7f78949bff9","resolution":{"observed_at":"2026-07-10T05:46:50.374699Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-07-18T16:21:18.861996+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-18T16:21:18.861996+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-10T06:26:52.335824Z","title":"Milvus: A purpose-built vector data management system","venue":null,"work_id":"989c80f8-3afe-41b9-a73e-1f91e024a9c8","year":2021},"citing_paper":{"arxiv_id":"2607.08541","last_updated":"2026-07-09T14:35:37Z","snapshot_observed_at":"2026-08-19T01:14:44.888182Z","submitted_at":"2026-07-09T14:35:37Z","title":"VocaDet: Sample-Driven Open-Vocabulary Object Detection and Segmentation via Visual Tokenization and Vector Database Retrieval","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-07-10T05:41:42.555295Z"},"links":{"citing_paper":"/paper/2607.08541"},"observation_digest":"sha256:7b71d96758a0696399c87eeac94bf8a4ab72296f0f6e65df7fe1486be717a7c5","observation_id":"705a2b59-4cbe-41aa-9e9e-e8d1883643ad","resolution":{"observed_at":"2026-07-10T06:26:52.337427Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-10T06:26:52.330804Z","title":"Ua-detrac: A new benchmark and protocol for multi-object detection and tracking.Computer Vision and Image Understanding, 193:102907, 2020","venue":null,"work_id":"f651e148-f5cd-474d-aa12-398f99df02db","year":2020},"citing_paper":{"arxiv_id":"2607.08541","last_updated":"2026-07-09T14:35:37Z","snapshot_observed_at":"2026-08-19T01:14:44.888182Z","submitted_at":"2026-07-09T14:35:37Z","title":"VocaDet: Sample-Driven Open-Vocabulary Object Detection and Segmentation via Visual Tokenization and Vector Database Retrieval","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-07-10T05:41:42.555295Z"},"links":{"citing_paper":"/paper/2607.08541"},"observation_digest":"sha256:a0a65953a981126fa22c46f966404eb1ceb3a89d4cfb22c03829ed97f39af1e5","observation_id":"34ba29b0-a720-409d-8700-d15524b99d01","resolution":{"observed_at":"2026-07-10T06:26:52.332459Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2607.08541","last_updated":"2026-07-09T14:35:37Z","latest_version":1,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-19T01:14:44.888182Z","submitted_at":"2026-07-09T14:35:37Z","title":"VocaDet: Sample-Driven Open-Vocabulary Object Detection and Segmentation via Visual Tokenization and Vector Database Retrieval"},"reference_resolution":{"displayed":16,"state_counts":{"malformed_identifier":0,"metadata_mismatch":2,"parse_uncertain":0,"unresolved":0,"verified_exact":7,"verified_fuzzy":7},"total_outbound_references":16},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"thesis":"As of 23 August 2026, this Paper Citation Record lists 16 of 16 outbound references and 0 inbound Pith citation observations for arXiv:2607.08541."}