{"as_of":"2026-08-09T03:56:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:7c4a469d3a91a5ec2cf1d421502fe710c575bc1e3be1a2a026a2405259ab320f","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":9,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":9,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-08T06:32:00.761636+00:00","state":"measured"},{"denominator":9,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":9,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T11:52:05.217519Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-07-04T03:29:31.600136Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2504.02259","last_updated":"2025-08-25T02:57:46Z","snapshot_observed_at":"2026-08-07T18:21:41.515675Z","submitted_at":"2025-04-03T04:03:10Z","title":"T*: Re-thinking Temporal Search for Long-Form Video Understanding","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.02259","snapshot_observed_at":"2026-08-07T11:51:29.605698Z","title":"Re-thinking temporal search for long-form video understanding","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.01274","last_updated":"2026-06-11T17:06:50Z","snapshot_observed_at":"2026-08-08T07:40:04.746667Z","submitted_at":"2025-06-02T03:08:07Z","title":"ReFoCUS: Reinforcement-guided Frame Optimization for Contextual Understanding","version":2},"reference_index":53,"source":"pdf_text","source_observed_at":"2026-08-07T11:51:29.605698Z"},"links":{"cited_paper":"/paper/2504.02259","citing_paper":"/paper/2506.01274"},"observation_digest":"sha256:c11b32bbc9f72df787e6ad4642e7dfefae7ed7ccccb82f8415fec0ea6d881d1b","observation_id":"37a14644-df17-4e42-81f1-0d1295035920","resolution":{"observed_at":"2026-08-07T11:51:29.605698Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2504.02259","last_updated":"2025-08-25T02:57:46Z","snapshot_observed_at":"2026-08-07T18:21:41.515675Z","submitted_at":"2025-04-03T04:03:10Z","title":"T*: Re-thinking Temporal Search for Long-Form Video Understanding","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.02259","snapshot_observed_at":"2026-08-07T11:52:05.217519Z","title":"Re-thinking temporal search for long-form video understanding.arXiv preprint arXiv:2504.02259, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.01300","last_updated":"2025-06-02T04:23:21Z","snapshot_observed_at":"2026-08-07T11:42:36.031718Z","submitted_at":"2025-06-02T04:23:21Z","title":"ReAgent-V: A Reward-Driven Multi-Agent Framework for Video Understanding","version":1},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-07T11:52:05.217519Z"},"links":{"cited_paper":"/paper/2504.02259","citing_paper":"/paper/2506.01300"},"observation_digest":"sha256:a27b7c2b1c37fb7cbc76e001c3fbe0c4bfd87fe6d8197200995c624004aa41d0","observation_id":"4fec2d96-4ea0-4ee6-83dd-722aeaa204ae","resolution":{"observed_at":"2026-08-07T11:52:05.217519Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2504.02259","last_updated":"2025-08-25T02:57:46Z","snapshot_observed_at":"2026-08-07T18:21:41.515675Z","submitted_at":"2025-04-03T04:03:10Z","title":"T*: Re-thinking Temporal Search for Long-Form Video Understanding","version":3},"cited_work":{"arxiv_id":"2504.02259","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2504.02259","snapshot_observed_at":"2026-07-04T03:29:31.600136Z","title":"T*: Re-thinking temporal search for long-form video understanding.arXiv preprint arXiv:2504.02259, 2025a","venue":null,"work_id":"d41d353c-6571-4286-83d5-f47e88b30133","year":2025},"citing_paper":{"arxiv_id":"2605.09223","last_updated":"2026-07-30T19:00:53Z","snapshot_observed_at":"2026-08-05T23:10:41.225358Z","submitted_at":"2026-05-09T23:47:46Z","title":"CREST: Curvature-Regulated Event-Centric Sampling for Efficient Long-Video Understanding","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-05-12T03:06:09.753634Z"},"links":{"cited_paper":"/paper/2504.02259","citing_paper":"/paper/2605.09223"},"observation_digest":"sha256:f6616de3ed9ced3ff443f35e328686e8b7789a20d53eb6915f20a8e742721cf8","observation_id":"1f6eeeef-6662-4412-8f4d-b4f6b1d5d80d","resolution":{"observed_at":"2026-05-12T03:06:18.776307Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.02259","last_updated":"2025-08-25T02:57:46Z","snapshot_observed_at":"2026-08-07T18:21:41.515675Z","submitted_at":"2025-04-03T04:03:10Z","title":"T*: Re-thinking Temporal Search for Long-Form Video Understanding","version":3},"cited_work":{"arxiv_id":"2504.02259","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2504.02259","snapshot_observed_at":"2026-07-04T03:29:31.600136Z","title":"T*: Re-thinking temporal search for long-form video understanding.arXiv preprint arXiv:2504.02259, 2025a","venue":null,"work_id":"d41d353c-6571-4286-83d5-f47e88b30133","year":2025},"citing_paper":{"arxiv_id":"2605.09223","last_updated":"2026-07-30T19:00:53Z","snapshot_observed_at":"2026-08-05T23:10:41.225358Z","submitted_at":"2026-05-09T23:47:46Z","title":"CREST: Curvature-Regulated Event-Centric Sampling for Efficient Long-Video Understanding","version":2},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-06-30T22:47:19.742035Z"},"links":{"cited_paper":"/paper/2504.02259","citing_paper":"/paper/2605.09223"},"observation_digest":"sha256:66594cb8360338651be49efd06a28d638ed16f567b331679b00ab6bec80cc5ba","observation_id":"ddcee522-7dc9-4fb7-b695-3556614a32f0","resolution":{"observed_at":"2026-07-01T13:45:46.073121Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.02259","last_updated":"2025-08-25T02:57:46Z","snapshot_observed_at":"2026-08-07T18:21:41.515675Z","submitted_at":"2025-04-03T04:03:10Z","title":"T*: Re-thinking Temporal Search for Long-Form Video Understanding","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.02259","snapshot_observed_at":"2026-08-03T00:18:47.111823Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2605.09223","last_updated":"2026-07-30T19:00:53Z","snapshot_observed_at":"2026-08-05T23:10:41.225358Z","submitted_at":"2026-05-09T23:47:46Z","title":"CREST: Curvature-Regulated Event-Centric Sampling for Efficient Long-Video Understanding","version":3},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-03T00:18:47.111823Z"},"links":{"cited_paper":"/paper/2504.02259","citing_paper":"/paper/2605.09223"},"observation_digest":"sha256:db2c66fa9dd2427f4f94c8272aec5e1fb53407d0556e37c3ec3e23b4cdeedf5f","observation_id":"94878f95-4c49-42f2-aea4-477a0398d094","resolution":{"observed_at":"2026-08-03T00:18:47.111823Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2504.02259","last_updated":"2025-08-25T02:57:46Z","snapshot_observed_at":"2026-08-07T18:21:41.515675Z","submitted_at":"2025-04-03T04:03:10Z","title":"T*: Re-thinking Temporal Search for Long-Form Video Understanding","version":3},"cited_work":{"arxiv_id":"2504.02259","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2504.02259","snapshot_observed_at":"2026-07-04T03:29:31.600136Z","title":"T*: Re-thinking temporal search for long-form video understanding.arXiv preprint arXiv:2504.02259, 2025a","venue":null,"work_id":"d41d353c-6571-4286-83d5-f47e88b30133","year":2025},"citing_paper":{"arxiv_id":"2605.31598","last_updated":"2026-05-29T17:59:02Z","snapshot_observed_at":"2026-08-07T05:09:45.280725Z","submitted_at":"2026-05-29T17:59:02Z","title":"Linear Scaling Video VLMs for Long Video Understanding","version":1},"reference_index":75,"source":"pdf_text","source_observed_at":"2026-06-28T23:00:11.246232Z"},"links":{"cited_paper":"/paper/2504.02259","citing_paper":"/paper/2605.31598"},"observation_digest":"sha256:33eae65f71eab0ece12941211e1fda0015d933d392c5146f85522d606fac5c73","observation_id":"29b2f1df-51f5-48a8-8437-ffcf681b2239","resolution":{"observed_at":"2026-06-28T23:02:46.245746Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.02259","last_updated":"2025-08-25T02:57:46Z","snapshot_observed_at":"2026-08-07T18:21:41.515675Z","submitted_at":"2025-04-03T04:03:10Z","title":"T*: Re-thinking Temporal Search for Long-Form Video Understanding","version":3},"cited_work":{"arxiv_id":"2504.02259","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2504.02259","snapshot_observed_at":"2026-07-04T03:29:31.600136Z","title":"T*: Re-thinking temporal search for long-form video understanding.arXiv preprint arXiv:2504.02259, 2025a","venue":null,"work_id":"d41d353c-6571-4286-83d5-f47e88b30133","year":2025},"citing_paper":{"arxiv_id":"2606.20515","last_updated":"2026-06-28T15:54:03Z","snapshot_observed_at":"2026-08-02T07:26:05.359571Z","submitted_at":"2026-06-18T17:34:55Z","title":"S-Agent: Spatial Tool-Use Elicits Reasoning for Spatial Intelligence","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-06-26T17:56:07.864580Z"},"links":{"cited_paper":"/paper/2504.02259","citing_paper":"/paper/2606.20515"},"observation_digest":"sha256:27619173e0f957c1da1674f46543ac78082a106358879c5f9668d6983a264393","observation_id":"986a833b-5de0-41dd-9bd1-dced1c54e0c1","resolution":{"observed_at":"2026-07-04T03:29:31.602171Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.02259","last_updated":"2025-08-25T02:57:46Z","snapshot_observed_at":"2026-08-07T18:21:41.515675Z","submitted_at":"2025-04-03T04:03:10Z","title":"T*: Re-thinking Temporal Search for Long-Form Video Understanding","version":3},"cited_work":{"arxiv_id":"2504.02259","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2504.02259","snapshot_observed_at":"2026-07-04T03:29:31.600136Z","title":"T*: Re-thinking temporal search for long-form video understanding.arXiv preprint arXiv:2504.02259, 2025a","venue":null,"work_id":"d41d353c-6571-4286-83d5-f47e88b30133","year":2025},"citing_paper":{"arxiv_id":"2606.20515","last_updated":"2026-06-28T15:54:03Z","snapshot_observed_at":"2026-08-02T07:26:05.359571Z","submitted_at":"2026-06-18T17:34:55Z","title":"S-Agent: Spatial Tool-Use Elicits Reasoning for Spatial Intelligence","version":2},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-06-30T10:22:08.535453Z"},"links":{"cited_paper":"/paper/2504.02259","citing_paper":"/paper/2606.20515"},"observation_digest":"sha256:a9960c872d5191593c97dbeedf93a71720a019ab40dd5d86a790718878b7501f","observation_id":"a7cb09df-09ff-4610-8dc6-e62902b59224","resolution":{"observed_at":"2026-06-30T11:54:39.061797Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.02259","last_updated":"2025-08-25T02:57:46Z","snapshot_observed_at":"2026-08-07T18:21:41.515675Z","submitted_at":"2025-04-03T04:03:10Z","title":"T*: Re-thinking Temporal Search for Long-Form Video Understanding","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.02259","snapshot_observed_at":"2026-08-01T21:09:45.454954Z","title":"Re-thinking temporal search for long-form video understanding.arXiv preprint arXiv:2504.02259, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.16189","last_updated":"2026-07-17T17:59:27Z","snapshot_observed_at":"2026-08-08T05:22:33.805882Z","submitted_at":"2026-07-17T17:59:27Z","title":"Searching Videos as Trees: Self-Correcting Agents for Grounded Long Video QA","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-01T21:09:45.454954Z"},"links":{"cited_paper":"/paper/2504.02259","citing_paper":"/paper/2607.16189"},"observation_digest":"sha256:ebdca5cd4fa42b852babd408824643b1a52a6048fd0ad28c519733ef132f658e","observation_id":"3bb651f5-d098-4d35-8655-7cfa911c0247","resolution":{"observed_at":"2026-08-01T21:09:45.454954Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2504.02259/citation-record","integrity":"/paper/2504.02259/integrity","json":"/paper/2504.02259/citation-record.json","paper":"/paper/2504.02259"},"outbound":[],"paper":{"arxiv_id":"2504.02259","last_updated":"2025-08-25T02:57:46Z","latest_version":3,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-07T18:21:41.515675Z","submitted_at":"2025-04-03T04:03:10Z","title":"T*: Re-thinking Temporal Search for Long-Form Video Understanding"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"thesis":"As of 9 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 9 inbound Pith citation observations for arXiv:2504.02259."}