{"as_of":"2026-08-13T13:26:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:ed2e5c152b38ac3a3691a71df34c1f4a49308877cafb2955cd9d49a5fb82ce33","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":7,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":7,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-13T06:32:02.005865+00:00","state":"measured"},{"denominator":7,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":7,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-05T17:46:46.750179Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-07-03T10:48:03.220493Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2409.20018","last_updated":"2024-10-02T09:34:11Z","snapshot_observed_at":"2026-08-12T22:34:47.995817Z","submitted_at":"2024-09-30T07:25:16Z","title":"Visual Context Window Extension: A New Perspective for Long Video Understanding","version":2},"cited_work":{"arxiv_id":"2409.20018","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2409.20018","snapshot_observed_at":"2026-07-03T10:48:03.220493Z","title":"Visual context window extension: A new perspective for long video understanding","venue":null,"work_id":"0e516a53-ba97-454a-90b4-45a765cf60b4","year":2024},"citing_paper":{"arxiv_id":"2501.00574","last_updated":"2025-07-13T16:21:16Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-12-31T18:01:23Z","title":"VideoChat-Flash: Hierarchical Compression for Long-Context Video Modeling","version":4},"reference_index":57,"source":"pdf_text","source_observed_at":"2026-05-18T04:02:43.261543Z"},"links":{"cited_paper":"/paper/2409.20018","citing_paper":"/paper/2501.00574"},"observation_digest":"sha256:650613e5bdf9060ad66b971c7ac95db8747230e3bd37c30a8bf1bf2bf4e393eb","observation_id":"ba361a7a-a74e-498e-8ba7-1aab43771bf1","resolution":{"observed_at":"2026-05-18T04:02:43.552982Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2409.20018","last_updated":"2024-10-02T09:34:11Z","snapshot_observed_at":"2026-08-12T22:34:47.995817Z","submitted_at":"2024-09-30T07:25:16Z","title":"Visual Context Window Extension: A New Perspective for Long Video Understanding","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.20018","snapshot_observed_at":"2026-08-05T17:46:46.750179Z","title":"Visual context window extension: A new perspective for long video understanding","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2508.15717","last_updated":"2025-08-21T16:56:29Z","snapshot_observed_at":"2026-08-12T09:46:58.745729Z","submitted_at":"2025-08-21T16:56:29Z","title":"StreamMem: Query-Agnostic KV Cache Memory for Streaming Video Understanding","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-05T17:46:46.750179Z"},"links":{"cited_paper":"/paper/2409.20018","citing_paper":"/paper/2508.15717"},"observation_digest":"sha256:dd33f8e1c89c8b0ac64c5362a8c6a0116f38b824d236980a82d02194d5b7534c","observation_id":"0dde0fbc-b0b1-4d05-bb63-26e47ffb689c","resolution":{"observed_at":"2026-08-05T17:46:46.750179Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.20018","last_updated":"2024-10-02T09:34:11Z","snapshot_observed_at":"2026-08-12T22:34:47.995817Z","submitted_at":"2024-09-30T07:25:16Z","title":"Visual Context Window Extension: A New Perspective for Long Video Understanding","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.20018","snapshot_observed_at":"2026-08-04T19:28:34.247158Z","title":"Visual context window extension: A new perspective for long video understanding,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2509.09263","last_updated":"2025-09-11T08:49:22Z","snapshot_observed_at":"2026-08-11T00:06:42.054152Z","submitted_at":"2025-09-11T08:49:22Z","title":"DATE: Dynamic Absolute Time Enhancement for Long Video Understanding","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-04T19:28:34.247158Z"},"links":{"cited_paper":"/paper/2409.20018","citing_paper":"/paper/2509.09263"},"observation_digest":"sha256:4fc9041c360d744f23bf02fbcce555fe13db4b1f8c4d6c5cd69bba3e28ed8242","observation_id":"22e1696a-3703-4675-9506-ef2c6cae2c87","resolution":{"observed_at":"2026-08-04T19:28:34.247158Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.20018","last_updated":"2024-10-02T09:34:11Z","snapshot_observed_at":"2026-08-12T22:34:47.995817Z","submitted_at":"2024-09-30T07:25:16Z","title":"Visual Context Window Extension: A New Perspective for Long Video Understanding","version":2},"cited_work":{"arxiv_id":"2409.20018","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2409.20018","snapshot_observed_at":"2026-07-03T10:48:03.220493Z","title":"Visual context window extension: A new perspective for long video understanding","venue":null,"work_id":"0e516a53-ba97-454a-90b4-45a765cf60b4","year":2024},"citing_paper":{"arxiv_id":"2605.06185","last_updated":"2026-05-07T13:01:28Z","snapshot_observed_at":"2026-08-11T17:06:47.265433Z","submitted_at":"2026-05-07T13:01:28Z","title":"Event-Causal RAG: A Retrieval-Augmented Generation Framework for Long Video Reasoning in Complex Scenarios","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-05-08T10:15:15.129358Z"},"links":{"cited_paper":"/paper/2409.20018","citing_paper":"/paper/2605.06185"},"observation_digest":"sha256:23d1b71d36d1983b276a217f914f69e3be1a6e4feafc8b07ffeedddee55321a4","observation_id":"fa2f9987-8ba9-487d-9619-d60f58622ae3","resolution":{"observed_at":"2026-05-11T20:06:13.303063Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2409.20018","last_updated":"2024-10-02T09:34:11Z","snapshot_observed_at":"2026-08-12T22:34:47.995817Z","submitted_at":"2024-09-30T07:25:16Z","title":"Visual Context Window Extension: A New Perspective for Long Video Understanding","version":2},"cited_work":{"arxiv_id":"2409.20018","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2409.20018","snapshot_observed_at":"2026-07-03T10:48:03.220493Z","title":"Visual context window extension: A new perspective for long video understanding","venue":null,"work_id":"0e516a53-ba97-454a-90b4-45a765cf60b4","year":2024},"citing_paper":{"arxiv_id":"2606.05748","last_updated":"2026-06-04T06:20:23Z","snapshot_observed_at":"2026-08-10T04:46:55.422655Z","submitted_at":"2026-06-04T06:20:23Z","title":"UNIVID: Unified Vision-Language Model for Video Moderation","version":1},"reference_index":58,"source":"arxiv_source","source_observed_at":"2026-06-27T22:56:26.674841Z"},"links":{"cited_paper":"/paper/2409.20018","citing_paper":"/paper/2606.05748"},"observation_digest":"sha256:71c9fad01d521fbf6a663c2b9a14030f5e1655779d5fcd53f4fdf0c5ecb556fc","observation_id":"01534f90-515a-4f47-9d64-a5cc70a0920a","resolution":{"observed_at":"2026-07-02T16:07:09.291604Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2409.20018","last_updated":"2024-10-02T09:34:11Z","snapshot_observed_at":"2026-08-12T22:34:47.995817Z","submitted_at":"2024-09-30T07:25:16Z","title":"Visual Context Window Extension: A New Perspective for Long Video Understanding","version":2},"cited_work":{"arxiv_id":"2409.20018","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2409.20018","snapshot_observed_at":"2026-07-03T10:48:03.220493Z","title":"Visual context window extension: A new perspective for long video understanding","venue":null,"work_id":"0e516a53-ba97-454a-90b4-45a765cf60b4","year":2024},"citing_paper":{"arxiv_id":"2606.11913","last_updated":"2026-06-10T10:43:35Z","snapshot_observed_at":"2026-08-10T12:01:19.974620Z","submitted_at":"2026-06-10T10:43:35Z","title":"From Content to Knowledge: Lightning Fast Long-Video Understanding with Neural Knowledge Representations","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-06-27T10:10:37.233160Z"},"links":{"cited_paper":"/paper/2409.20018","citing_paper":"/paper/2606.11913"},"observation_digest":"sha256:09bf4b3a4474cffabf76c17f2a2427afcd4f2c1dd3cb33de98a891c9052414b5","observation_id":"b990bd02-1e6f-434b-bab0-9d981d3092d9","resolution":{"observed_at":"2026-07-03T10:17:57.400734Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2409.20018","last_updated":"2024-10-02T09:34:11Z","snapshot_observed_at":"2026-08-12T22:34:47.995817Z","submitted_at":"2024-09-30T07:25:16Z","title":"Visual Context Window Extension: A New Perspective for Long Video Understanding","version":2},"cited_work":{"arxiv_id":"2409.20018","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2409.20018","snapshot_observed_at":"2026-07-03T10:48:03.220493Z","title":"Visual context window extension: A new perspective for long video understanding","venue":null,"work_id":"0e516a53-ba97-454a-90b4-45a765cf60b4","year":2024},"citing_paper":{"arxiv_id":"2606.12195","last_updated":"2026-06-10T15:17:08Z","snapshot_observed_at":"2026-08-01T02:09:41.655807Z","submitted_at":"2026-06-10T15:17:08Z","title":"InternVideo3: Agentify Foundation Models with Multimodal Contextual Reasoning","version":1},"reference_index":284,"source":"arxiv_source","source_observed_at":"2026-06-27T09:48:27.652901Z"},"links":{"cited_paper":"/paper/2409.20018","citing_paper":"/paper/2606.12195"},"observation_digest":"sha256:5d09599d2fc229475184454621849bb61c160e1aa7f98ca74e81874181169524","observation_id":"0946f4d4-2194-425d-819c-85305bd23e68","resolution":{"observed_at":"2026-07-03T10:48:03.221935Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2409.20018/citation-record","integrity":"/paper/2409.20018/integrity","json":"/paper/2409.20018/citation-record.json","paper":"/paper/2409.20018"},"outbound":[],"paper":{"arxiv_id":"2409.20018","last_updated":"2024-10-02T09:34:11Z","latest_version":2,"primary_category":"cs.CV","snapshot_observed_at":"2026-08-12T22:34:47.995817Z","submitted_at":"2024-09-30T07:25:16Z","title":"Visual Context Window Extension: A New Perspective for Long Video Understanding"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"thesis":"As of 13 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 7 inbound Pith citation observations for arXiv:2409.20018."}