{"as_of":"2026-08-22T20:18:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:15b546e24deaabf6248d0371e9b83bfca5a116fa9a1d0ed587bbea8d9e99e493","coverage":[{"denominator":10,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":10,"source":"paper_references, paper_reference_links","source_observed_at":"2026-07-01T06:10:34.907872Z","state":"measured"},{"denominator":10,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":10,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-22T06:32:14.747728+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2606.31645/citation-record","integrity":"/paper/2606.31645/integrity","json":"/paper/2606.31645/citation-record.json","paper":"/paper/2606.31645"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2511.21631","last_updated":"2025-11-27T12:16:54Z","snapshot_observed_at":"2026-08-17T13:26:10.378579Z","submitted_at":"2025-11-26T17:59:08Z","title":"Qwen3-VL Technical Report","version":2},"cited_work":{"arxiv_id":"2511.21631","doi":"10.1016/j.neunet.2025.107777","metadata_source":"pith","pith_arxiv_id":"2511.21631","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Qwen3-VL Technical Report","venue":"cs.CV","work_id":"1fe243aa-e3c0-4da6-b391-4cbcfc88d5c0","year":2025},"citing_paper":{"arxiv_id":"2606.31645","last_updated":"2026-06-30T13:25:59Z","snapshot_observed_at":"2026-08-05T07:35:04.430557Z","submitted_at":"2026-06-30T13:25:59Z","title":"Technical Report of RoboSpatial Challenge at CVPR 2026: Selective Reasoning Activation and Reference-Frame Disambiguation for Embodied Spatial Reasoning","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-07-01T06:10:34.907872Z"},"links":{"cited_paper":"/paper/2511.21631","citing_paper":"/paper/2606.31645"},"observation_digest":"sha256:1be0956f07d9b6db0c2efb667c8ebba40b25a8cdd6696780979df24bf41ce4e5","observation_id":"ff755b4d-d502-4f6d-9e7a-0693e0db645f","resolution":{"observed_at":"2026-07-01T09:45:40.535076Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-06T19:12:51.073381Z","title":"Spatialvlm: Endow- ing vision-language models with spatial reasoning capabili- ties","venue":null,"work_id":"77a0ebe3-d7a7-4cf7-9cd2-ab8ff05d4e25","year":2024},"citing_paper":{"arxiv_id":"2606.31645","last_updated":"2026-06-30T13:25:59Z","snapshot_observed_at":"2026-08-05T07:35:04.430557Z","submitted_at":"2026-06-30T13:25:59Z","title":"Technical Report of RoboSpatial Challenge at CVPR 2026: Selective Reasoning Activation and Reference-Frame Disambiguation for Embodied Spatial Reasoning","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-07-01T06:10:34.907872Z"},"links":{"citing_paper":"/paper/2606.31645"},"observation_digest":"sha256:24e91b27285267380871b09196158524d671e2704923658b9b43e307a33a7741","observation_id":"b7f826fe-ea34-4737-b707-bcb9f64b6b9a","resolution":{"observed_at":"2026-07-06T19:12:51.074754Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-06T19:12:51.079716Z","title":"What’s “up” with vision-language models? investigating their strug- gle with spatial reasoning","venue":null,"work_id":"1ebc5d7b-2b45-4d5a-8626-a150ef8ac2d4","year":2023},"citing_paper":{"arxiv_id":"2606.31645","last_updated":"2026-06-30T13:25:59Z","snapshot_observed_at":"2026-08-05T07:35:04.430557Z","submitted_at":"2026-06-30T13:25:59Z","title":"Technical Report of RoboSpatial Challenge at CVPR 2026: Selective Reasoning Activation and Reference-Frame Disambiguation for Embodied Spatial Reasoning","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-07-01T06:10:34.907872Z"},"links":{"citing_paper":"/paper/2606.31645"},"observation_digest":"sha256:be8b75f1e5dff12ea58afcafded502ad833a5031b5cfeb352c1ff65fd1a4507b","observation_id":"eed8534e-d25d-4fd3-a56b-e725286e71f7","resolution":{"observed_at":"2026-07-06T19:12:51.081129Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-06T19:12:51.077647Z","title":"Viewspatial-bench: Evaluating multi-perspective spatial localization in vision-language models, 2025","venue":null,"work_id":"27704015-0569-4673-a1a1-5907b19c804f","year":2025},"citing_paper":{"arxiv_id":"2606.31645","last_updated":"2026-06-30T13:25:59Z","snapshot_observed_at":"2026-08-05T07:35:04.430557Z","submitted_at":"2026-06-30T13:25:59Z","title":"Technical Report of RoboSpatial Challenge at CVPR 2026: Selective Reasoning Activation and Reference-Frame Disambiguation for Embodied Spatial Reasoning","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-07-01T06:10:34.907872Z"},"links":{"citing_paper":"/paper/2606.31645"},"observation_digest":"sha256:172c3e75cca8b9f37a4718d1c609ea0d6ba50e84c2b5bd341f507e25192f3d14","observation_id":"c72daf34-0b73-45b9-ab22-295e90c381a6","resolution":{"observed_at":"2026-07-06T19:12:51.079183Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-06T19:12:51.081698Z","title":"An empirical study of catastrophic forgetting in large language models during continual fine-tuning.IEEE Transactions on Audio, Speech and Language Processing, 33:3776–3786, 2025","venue":null,"work_id":"6458ae53-d684-4245-a4a0-1b3ef6c008ad","year":2025},"citing_paper":{"arxiv_id":"2606.31645","last_updated":"2026-06-30T13:25:59Z","snapshot_observed_at":"2026-08-05T07:35:04.430557Z","submitted_at":"2026-06-30T13:25:59Z","title":"Technical Report of RoboSpatial Challenge at CVPR 2026: Selective Reasoning Activation and Reference-Frame Disambiguation for Embodied Spatial Reasoning","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-07-01T06:10:34.907872Z"},"links":{"citing_paper":"/paper/2606.31645"},"observation_digest":"sha256:12a5c107de71a5d1a41d8f4e39ea9b31013b9e739d5bc0666ceaaa6009b62880","observation_id":"2eb1ceef-f474-45bf-9716-e9ead734e80f","resolution":{"observed_at":"2026-07-06T19:12:51.083212Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-06T19:12:51.075600Z","title":"RoboSpatial: Teaching spatial understanding to 2D and 3D vision-language models for robotics","venue":null,"work_id":"2443adc0-cd62-4613-91d6-18b61c13378c","year":2025},"citing_paper":{"arxiv_id":"2606.31645","last_updated":"2026-06-30T13:25:59Z","snapshot_observed_at":"2026-08-05T07:35:04.430557Z","submitted_at":"2026-06-30T13:25:59Z","title":"Technical Report of RoboSpatial Challenge at CVPR 2026: Selective Reasoning Activation and Reference-Frame Disambiguation for Embodied Spatial Reasoning","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-07-01T06:10:34.907872Z"},"links":{"citing_paper":"/paper/2606.31645"},"observation_digest":"sha256:9cb917813c86f1fd3f1953e484b06558be239f08656c1a0a968c5549145938b3","observation_id":"a0ca1ec6-5235-4aeb-a4fc-a7e112a8ed2b","resolution":{"observed_at":"2026-07-06T19:12:51.076944Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2601.14352","doi":"10.48550/arxiv.2601.14352","metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Robobrain 2.5: Depth in sight, time in mind","venue":"arXiv (Cornell University)","work_id":"cb717066-349a-49ec-820d-6b2da63b7524","year":2026},"citing_paper":{"arxiv_id":"2606.31645","last_updated":"2026-06-30T13:25:59Z","snapshot_observed_at":"2026-08-05T07:35:04.430557Z","submitted_at":"2026-06-30T13:25:59Z","title":"Technical Report of RoboSpatial Challenge at CVPR 2026: Selective Reasoning Activation and Reference-Frame Disambiguation for Embodied Spatial Reasoning","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-07-01T06:10:34.907872Z"},"links":{"citing_paper":"/paper/2606.31645"},"observation_digest":"sha256:a6a42e225dd995e210cda4d9d2fcef9261d91df6a46b07e59f3788dffc80f8fc","observation_id":"8ff46ac1-49d0-4bc5-b2cd-6b390bfa720f","resolution":{"observed_at":"2026-07-01T09:45:40.534588Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-06T19:12:51.085810Z","title":"Qwen3.5: Accelerating productivity with na- tive multimodal agents, 2026","venue":null,"work_id":"911c14bb-d82a-427b-ac95-1680792bf1e2","year":2026},"citing_paper":{"arxiv_id":"2606.31645","last_updated":"2026-06-30T13:25:59Z","snapshot_observed_at":"2026-08-05T07:35:04.430557Z","submitted_at":"2026-06-30T13:25:59Z","title":"Technical Report of RoboSpatial Challenge at CVPR 2026: Selective Reasoning Activation and Reference-Frame Disambiguation for Embodied Spatial Reasoning","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-07-01T06:10:34.907872Z"},"links":{"citing_paper":"/paper/2606.31645"},"observation_digest":"sha256:d516f8923647db40a00938b9aa0a6e44554b5044745a962a2c0b844bf87404fa","observation_id":"99d4f1d8-e00b-464d-9e18-c74e6625cb6f","resolution":{"observed_at":"2026-07-06T19:12:51.087361Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-06T19:12:51.088166Z","title":"Embodiedscan: A holistic multi-modal 3d perception suite towards embodied ai","venue":null,"work_id":"c59b7930-acc1-444e-bf7d-dade41bb1c3b","year":2024},"citing_paper":{"arxiv_id":"2606.31645","last_updated":"2026-06-30T13:25:59Z","snapshot_observed_at":"2026-08-05T07:35:04.430557Z","submitted_at":"2026-06-30T13:25:59Z","title":"Technical Report of RoboSpatial Challenge at CVPR 2026: Selective Reasoning Activation and Reference-Frame Disambiguation for Embodied Spatial Reasoning","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-07-01T06:10:34.907872Z"},"links":{"citing_paper":"/paper/2606.31645"},"observation_digest":"sha256:d92b5a74acb26ff8d2476cfe035f114cf6a1edd6a1b063e936542872c37c7d56","observation_id":"897b3651-4f3f-4196-8b57-8389bbd54d62","resolution":{"observed_at":"2026-07-06T19:12:51.089654Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-06T19:12:51.083802Z","title":"Scannet++: A high-fidelity dataset of 3d in- door scenes","venue":null,"work_id":"e8bc1b73-28f2-4409-9798-64e4a39f165c","year":2023},"citing_paper":{"arxiv_id":"2606.31645","last_updated":"2026-06-30T13:25:59Z","snapshot_observed_at":"2026-08-05T07:35:04.430557Z","submitted_at":"2026-06-30T13:25:59Z","title":"Technical Report of RoboSpatial Challenge at CVPR 2026: Selective Reasoning Activation and Reference-Frame Disambiguation for Embodied Spatial Reasoning","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-07-01T06:10:34.907872Z"},"links":{"citing_paper":"/paper/2606.31645"},"observation_digest":"sha256:e25a0d95f17d29211bc960c31fc881858a201c5e516b71e21d65450d9f39c03a","observation_id":"3020143a-ec5f-464b-9b7f-70907f99448f","resolution":{"observed_at":"2026-07-06T19:12:51.085272Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2606.31645","last_updated":"2026-06-30T13:25:59Z","latest_version":1,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-05T07:35:04.430557Z","submitted_at":"2026-06-30T13:25:59Z","title":"Technical Report of RoboSpatial Challenge at CVPR 2026: Selective Reasoning Activation and Reference-Frame Disambiguation for Embodied Spatial Reasoning"},"reference_resolution":{"displayed":10,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":2,"verified_fuzzy":8},"total_outbound_references":10},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"thesis":"As of 22 August 2026, this Paper Citation Record lists 10 of 10 outbound references and 0 inbound Pith citation observations for arXiv:2606.31645."}