{"as_of":"2026-08-05T17:36:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:fe9202cbc4d0789030550e857ae1647e81eefc5bdf9112b1383ec07fba83b227","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":13,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":13,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-05T06:32:48.257954+00:00","state":"measured"},{"denominator":13,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":13,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-04T08:39:01.053662Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-07-04T01:19:20.288097Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2507.09313","last_updated":"2025-07-15T11:48:07Z","snapshot_observed_at":"2026-07-06T21:56:09.900882Z","submitted_at":"2025-07-12T15:11:50Z","title":"ProactiveVideoQA: A Comprehensive Benchmark Evaluating Proactive Interactions in Video Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2507.09313","snapshot_observed_at":"2026-08-04T08:39:01.053662Z","title":"Binfeng Xu, Zhiyuan PENG, Bowen Lei, Subhabrata Mukherjee, and Dongkuan Xu","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2510.19771","last_updated":"2026-07-07T17:34:06Z","snapshot_observed_at":"2026-08-04T08:38:56.874579Z","submitted_at":"2025-10-22T17:00:45Z","title":"Beyond Reactivity: Measuring Proactive Problem Solving in LLM Agents","version":4},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-04T08:39:01.053662Z"},"links":{"cited_paper":"/paper/2507.09313","citing_paper":"/paper/2510.19771"},"observation_digest":"sha256:c9f7b47f0d1629f55849fee32e5692fc6302b2f8ca9f0c7260503da7190d1063","observation_id":"5f6fd4f2-22b5-45a3-80e1-e540f88c8c65","resolution":{"observed_at":"2026-08-04T08:39:01.053662Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2507.09313","last_updated":"2025-07-15T11:48:07Z","snapshot_observed_at":"2026-07-06T21:56:09.900882Z","submitted_at":"2025-07-12T15:11:50Z","title":"ProactiveVideoQA: A Comprehensive Benchmark Evaluating Proactive Interactions in Video Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2507.09313","snapshot_observed_at":"2026-08-02T19:09:09.440908Z","title":"fill-with-silence","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2603.03447","last_updated":"2026-05-24T13:26:19Z","snapshot_observed_at":"2026-08-03T01:28:13.817274Z","submitted_at":"2026-03-03T19:02:46Z","title":"Proact-VL: A Proactive VideoLLM for Real-Time AI Companions","version":3},"reference_index":336,"source":"pdf_text","source_observed_at":"2026-08-02T19:09:09.440908Z"},"links":{"cited_paper":"/paper/2507.09313","citing_paper":"/paper/2603.03447"},"observation_digest":"sha256:b1ddb294ae6433e0ed33fb4de3c54eeea8e74ca0efbd31a18bc1b439c554c041","observation_id":"4027df28-ad4c-4b0f-917b-4e48460d6b75","resolution":{"observed_at":"2026-08-02T19:09:09.440908Z","resolver_source":null,"status":"malformed_identifier"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2507.09313","last_updated":"2025-07-15T11:48:07Z","snapshot_observed_at":"2026-07-06T21:56:09.900882Z","submitted_at":"2025-07-12T15:11:50Z","title":"ProactiveVideoQA: A Comprehensive Benchmark Evaluating Proactive Interactions in Video Large Language Models","version":2},"cited_work":{"arxiv_id":"2507.09313","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2507.09313","snapshot_observed_at":"2026-07-04T01:19:20.288097Z","title":"Proactivevideoqa: A comprehensive benchmark evaluating proactive interactions in video large language models","venue":null,"work_id":"6c44dc17-6083-4f26-b7c6-e27f819dd061","year":2025},"citing_paper":{"arxiv_id":"2604.15037","last_updated":"2026-05-02T12:33:41Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2026-04-16T14:06:30Z","title":"From Reactive to Proactive: Assessing the Proactivity of Voice Agents via ProVoice-Bench","version":3},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-05-10T11:44:12.373082Z"},"links":{"cited_paper":"/paper/2507.09313","citing_paper":"/paper/2604.15037"},"observation_digest":"sha256:6f9cca57ab5f977ea7d80e8989f3ac7b0b417b072e182f2f76f9a10dd55e034c","observation_id":"4572f0d0-b18f-45ee-804b-a0cc2491ab37","resolution":{"observed_at":"2026-05-10T11:45:21.033091Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2507.09313","last_updated":"2025-07-15T11:48:07Z","snapshot_observed_at":"2026-07-06T21:56:09.900882Z","submitted_at":"2025-07-12T15:11:50Z","title":"ProactiveVideoQA: A Comprehensive Benchmark Evaluating Proactive Interactions in Video Large Language Models","version":2},"cited_work":{"arxiv_id":"2507.09313","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2507.09313","snapshot_observed_at":"2026-07-04T01:19:20.288097Z","title":"Proactivevideoqa: A comprehensive benchmark evaluating proactive interactions in video large language models","venue":null,"work_id":"6c44dc17-6083-4f26-b7c6-e27f819dd061","year":2025},"citing_paper":{"arxiv_id":"2604.24317","last_updated":"2026-04-27T11:07:03Z","snapshot_observed_at":"2026-07-06T23:10:24.992210Z","submitted_at":"2026-04-27T11:07:03Z","title":"Don't Pause! Every prediction matters in a streaming video","version":1},"reference_index":61,"source":"pdf_text","source_observed_at":"2026-05-08T04:32:01.379605Z"},"links":{"cited_paper":"/paper/2507.09313","citing_paper":"/paper/2604.24317"},"observation_digest":"sha256:234d09569762b84adf20702102e88e6da78e6fd06fd1f5de55039ef155e6959d","observation_id":"6c92c69d-1533-4d11-8cc6-ae801125e38f","resolution":{"observed_at":"2026-05-11T21:41:18.097956Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2507.09313","last_updated":"2025-07-15T11:48:07Z","snapshot_observed_at":"2026-07-06T21:56:09.900882Z","submitted_at":"2025-07-12T15:11:50Z","title":"ProactiveVideoQA: A Comprehensive Benchmark Evaluating Proactive Interactions in Video Large Language Models","version":2},"cited_work":{"arxiv_id":"2507.09313","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2507.09313","snapshot_observed_at":"2026-07-04T01:19:20.288097Z","title":"Proactivevideoqa: A comprehensive benchmark evaluating proactive interactions in video large language models","venue":null,"work_id":"6c44dc17-6083-4f26-b7c6-e27f819dd061","year":2025},"citing_paper":{"arxiv_id":"2605.16381","last_updated":"2026-05-11T05:01:15Z","snapshot_observed_at":"2026-08-02T15:40:21.659895Z","submitted_at":"2026-05-11T05:01:15Z","title":"StreamPro: From Reactive Perception to Proactive Decision-Making in Streaming Video","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-05-20T23:32:38.332828Z"},"links":{"cited_paper":"/paper/2507.09313","citing_paper":"/paper/2605.16381"},"observation_digest":"sha256:32f78c3511bbcff5f4053c491722a62745801429a54d03e7733c32eabcd84262","observation_id":"068e4c34-be0c-4134-84b3-3b821d3b358c","resolution":{"observed_at":"2026-05-20T23:33:50.748812Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2507.09313","last_updated":"2025-07-15T11:48:07Z","snapshot_observed_at":"2026-07-06T21:56:09.900882Z","submitted_at":"2025-07-12T15:11:50Z","title":"ProactiveVideoQA: A Comprehensive Benchmark Evaluating Proactive Interactions in Video Large Language Models","version":2},"cited_work":{"arxiv_id":"2507.09313","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2507.09313","snapshot_observed_at":"2026-07-04T01:19:20.288097Z","title":"Proactivevideoqa: A comprehensive benchmark evaluating proactive interactions in video large language models","venue":null,"work_id":"6c44dc17-6083-4f26-b7c6-e27f819dd061","year":2025},"citing_paper":{"arxiv_id":"2605.17360","last_updated":"2026-07-02T12:39:55Z","snapshot_observed_at":"2026-07-06T23:28:21.395050Z","submitted_at":"2026-05-17T09:57:01Z","title":"Omni-DuplexEval: Evaluating Real-time Duplex Omni-modal Interaction","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-05-20T13:36:44.071188Z"},"links":{"cited_paper":"/paper/2507.09313","citing_paper":"/paper/2605.17360"},"observation_digest":"sha256:4d2cedf116b18600ab8f5a83bcfd9ac105119a71498ade68bfff59a36e1d2455","observation_id":"41aead78-12f1-4da2-9363-84e7b519d662","resolution":{"observed_at":"2026-05-20T13:38:19.157208Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2507.09313","last_updated":"2025-07-15T11:48:07Z","snapshot_observed_at":"2026-07-06T21:56:09.900882Z","submitted_at":"2025-07-12T15:11:50Z","title":"ProactiveVideoQA: A Comprehensive Benchmark Evaluating Proactive Interactions in Video Large Language Models","version":2},"cited_work":{"arxiv_id":"2507.09313","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2507.09313","snapshot_observed_at":"2026-07-04T01:19:20.288097Z","title":"Proactivevideoqa: A comprehensive benchmark evaluating proactive interactions in video large language models","venue":null,"work_id":"6c44dc17-6083-4f26-b7c6-e27f819dd061","year":2025},"citing_paper":{"arxiv_id":"2605.17360","last_updated":"2026-07-02T12:39:55Z","snapshot_observed_at":"2026-07-06T23:28:21.395050Z","submitted_at":"2026-05-17T09:57:01Z","title":"Omni-DuplexEval: Evaluating Real-time Duplex Omni-modal Interaction","version":2},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-07-04T01:11:42.073993Z"},"links":{"cited_paper":"/paper/2507.09313","citing_paper":"/paper/2605.17360"},"observation_digest":"sha256:1d1b7f6cbed04324ce803a7f4ec784eaa996e9e52eaf8568692b2e7782d54ce8","observation_id":"ca766167-9a2e-46cc-bc82-858fed26fffc","resolution":{"observed_at":"2026-07-04T01:19:20.291251Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2507.09313","last_updated":"2025-07-15T11:48:07Z","snapshot_observed_at":"2026-07-06T21:56:09.900882Z","submitted_at":"2025-07-12T15:11:50Z","title":"ProactiveVideoQA: A Comprehensive Benchmark Evaluating Proactive Interactions in Video Large Language Models","version":2},"cited_work":{"arxiv_id":"2507.09313","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2507.09313","snapshot_observed_at":"2026-07-04T01:19:20.288097Z","title":"Proactivevideoqa: A comprehensive benchmark evaluating proactive interactions in video large language models","venue":null,"work_id":"6c44dc17-6083-4f26-b7c6-e27f819dd061","year":2025},"citing_paper":{"arxiv_id":"2605.17921","last_updated":"2026-06-01T06:29:58Z","snapshot_observed_at":"2026-07-06T23:28:50.117638Z","submitted_at":"2026-05-18T06:29:44Z","title":"An Efficient Streaming Video Understanding Framework with Agentic Control","version":1},"reference_index":27,"source":"arxiv_source","source_observed_at":"2026-05-20T11:30:22.151045Z"},"links":{"cited_paper":"/paper/2507.09313","citing_paper":"/paper/2605.17921"},"observation_digest":"sha256:3ff9160cf4292a0afc5aaf99ff3634a5ccf2272024209af7aa8383398a1f1fbd","observation_id":"fb22717f-9da6-4ae1-acb9-ffaba75f432d","resolution":{"observed_at":"2026-05-20T11:33:14.470074Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2507.09313","last_updated":"2025-07-15T11:48:07Z","snapshot_observed_at":"2026-07-06T21:56:09.900882Z","submitted_at":"2025-07-12T15:11:50Z","title":"ProactiveVideoQA: A Comprehensive Benchmark Evaluating Proactive Interactions in Video Large Language Models","version":2},"cited_work":{"arxiv_id":"2507.09313","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2507.09313","snapshot_observed_at":"2026-07-04T01:19:20.288097Z","title":"Proactivevideoqa: A comprehensive benchmark evaluating proactive interactions in video large language models","venue":null,"work_id":"6c44dc17-6083-4f26-b7c6-e27f819dd061","year":2025},"citing_paper":{"arxiv_id":"2605.27074","last_updated":"2026-05-26T14:23:25Z","snapshot_observed_at":"2026-08-02T20:44:04.329111Z","submitted_at":"2026-05-26T14:23:25Z","title":"IPIBench: Evaluating Interactive Proactive Intelligence of MLLMs under Continuous Streams","version":1},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-06-29T18:24:57.881644Z"},"links":{"cited_paper":"/paper/2507.09313","citing_paper":"/paper/2605.27074"},"observation_digest":"sha256:2b22316154e672fe4405aa8bdb8395542c9a6eae481889318ad580672b22b574","observation_id":"3e967234-cadc-4cee-95cf-5fbac879f8a2","resolution":{"observed_at":"2026-06-29T18:33:51.065824Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2507.09313","last_updated":"2025-07-15T11:48:07Z","snapshot_observed_at":"2026-07-06T21:56:09.900882Z","submitted_at":"2025-07-12T15:11:50Z","title":"ProactiveVideoQA: A Comprehensive Benchmark Evaluating Proactive Interactions in Video Large Language Models","version":2},"cited_work":{"arxiv_id":"2507.09313","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2507.09313","snapshot_observed_at":"2026-07-04T01:19:20.288097Z","title":"Proactivevideoqa: A comprehensive benchmark evaluating proactive interactions in video large language models","venue":null,"work_id":"6c44dc17-6083-4f26-b7c6-e27f819dd061","year":2025},"citing_paper":{"arxiv_id":"2606.02482","last_updated":"2026-06-29T07:37:14Z","snapshot_observed_at":"2026-07-06T23:42:53.514946Z","submitted_at":"2026-06-01T16:52:11Z","title":"X-Stream: Exploring MLLMs as Multiplexers for Multi-Stream Understanding","version":2},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-06-28T15:06:22.102725Z"},"links":{"cited_paper":"/paper/2507.09313","citing_paper":"/paper/2606.02482"},"observation_digest":"sha256:40a2c995a379497aa01a483e5d8e428b98546e090d657821064dcad788f71018","observation_id":"c4af456a-a752-4d00-9ce7-3a177dc71ce5","resolution":{"observed_at":"2026-07-01T22:46:18.934183Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2507.09313","last_updated":"2025-07-15T11:48:07Z","snapshot_observed_at":"2026-07-06T21:56:09.900882Z","submitted_at":"2025-07-12T15:11:50Z","title":"ProactiveVideoQA: A Comprehensive Benchmark Evaluating Proactive Interactions in Video Large Language Models","version":2},"cited_work":{"arxiv_id":"2507.09313","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2507.09313","snapshot_observed_at":"2026-07-04T01:19:20.288097Z","title":"Proactivevideoqa: A comprehensive benchmark evaluating proactive interactions in video large language models","venue":null,"work_id":"6c44dc17-6083-4f26-b7c6-e27f819dd061","year":2025},"citing_paper":{"arxiv_id":"2606.02482","last_updated":"2026-06-29T07:37:14Z","snapshot_observed_at":"2026-07-06T23:42:53.514946Z","submitted_at":"2026-06-01T16:52:11Z","title":"X-Stream: Exploring MLLMs as Multiplexers for Multi-Stream Understanding","version":3},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-06-30T10:42:37.401221Z"},"links":{"cited_paper":"/paper/2507.09313","citing_paper":"/paper/2606.02482"},"observation_digest":"sha256:06aaf4c932b2a4b87d24a95a9a55ecfabda52467282b09553e9c7cf403dbf80e","observation_id":"0dd45640-7b1f-485c-b538-17545bed0fe1","resolution":{"observed_at":"2026-06-30T10:44:36.467381Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2507.09313","last_updated":"2025-07-15T11:48:07Z","snapshot_observed_at":"2026-07-06T21:56:09.900882Z","submitted_at":"2025-07-12T15:11:50Z","title":"ProactiveVideoQA: A Comprehensive Benchmark Evaluating Proactive Interactions in Video Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2507.09313","snapshot_observed_at":"2026-07-12T05:36:40.278384Z","title":"arXiv preprint arXiv:2507.09313 (2025) 5, 8","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.02991","last_updated":"2026-07-03T06:00:04Z","snapshot_observed_at":"2026-08-01T07:54:59.412079Z","submitted_at":"2026-07-03T06:00:04Z","title":"GuideMe: Multi-Domain Task Guidance and Intervention in Streaming Video","version":1},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-07-12T05:36:40.278384Z"},"links":{"cited_paper":"/paper/2507.09313","citing_paper":"/paper/2607.02991"},"observation_digest":"sha256:fc319dd3a1699b2cc4aec64f83b71523f234a4a81042d280188743874f171603","observation_id":"209f9f59-9164-42e5-8a37-7f88f5dc3007","resolution":{"observed_at":"2026-07-12T05:36:40.278384Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2507.09313","last_updated":"2025-07-15T11:48:07Z","snapshot_observed_at":"2026-07-06T21:56:09.900882Z","submitted_at":"2025-07-12T15:11:50Z","title":"ProactiveVideoQA: A Comprehensive Benchmark Evaluating Proactive Interactions in Video Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2507.09313","snapshot_observed_at":"2026-08-02T00:44:48.909159Z","title":"Proactivevideoqa: A comprehensive benchmark evaluating proactive interactions in video large language models","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.14935","last_updated":"2026-07-16T12:47:59Z","snapshot_observed_at":"2026-08-02T14:24:21.558174Z","submitted_at":"2026-07-16T12:47:59Z","title":"VideoChat3: Fully Open Video MLLM for Efficient and Generalist Video Understanding","version":1},"reference_index":89,"source":"pdf_text","source_observed_at":"2026-08-02T00:44:48.909159Z"},"links":{"cited_paper":"/paper/2507.09313","citing_paper":"/paper/2607.14935"},"observation_digest":"sha256:7bef5013aead7ab8ea38087965eec86a08fd6a2dae7185a387c3645e1e37c1ba","observation_id":"7523c331-160b-4a39-b74a-911205ccb356","resolution":{"observed_at":"2026-08-02T00:44:48.909159Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2507.09313/citation-record","integrity":"/paper/2507.09313/integrity","json":"/paper/2507.09313/citation-record.json","paper":"/paper/2507.09313"},"outbound":[],"paper":{"arxiv_id":"2507.09313","last_updated":"2025-07-15T11:48:07Z","latest_version":2,"primary_category":"cs.CV","snapshot_observed_at":"2026-07-06T21:56:09.900882Z","submitted_at":"2025-07-12T15:11:50Z","title":"ProactiveVideoQA: A Comprehensive Benchmark Evaluating Proactive Interactions in Video Large Language Models"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"thesis":"As of 5 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 13 inbound Pith citation observations for arXiv:2507.09313."}