{"as_of":"2026-08-11T11:03:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:158fcbfc195cc6d6dfe45be15bd74169c6337658d395cebcfd06ad380f718074","coverage":[{"denominator":73,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":73,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-10T13:53:27.583766Z","state":"measured"},{"denominator":75,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":75,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-11T06:34:44.6726+00:00","state":"measured"},{"denominator":2,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":2,"source":"paper_references, paper_reference_links","source_observed_at":"2026-06-26T05:23:06.035263Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-07-04T13:19:50.408368Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"cited_work":{"arxiv_id":"2501.15953","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2501.15953","snapshot_observed_at":"2026-07-04T13:19:50.408368Z","title":"Understandinglongvideosviallm-poweredentityrelation graphs","venue":null,"work_id":"aa1b6435-f1e1-446e-95d9-5fd71504aaca","year":2025},"citing_paper":{"arxiv_id":"2504.01990","last_updated":"2025-08-02T12:44:02Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-03-31T18:00:29Z","title":"Advances and Challenges in Foundation Agents: From Brain-Inspired Intelligence to Evolutionary, Collaborative, and Safe Systems","version":2},"reference_index":296,"source":"pdf_text","source_observed_at":"2026-05-22T21:39:49.832151Z"},"links":{"cited_paper":"/paper/2501.15953","citing_paper":"/paper/2504.01990"},"observation_digest":"sha256:c225285e7ecc116575b62ea234029ad77ab6933a5aa07035d320dd15fd5ce293","observation_id":"4f9a78ba-d8c6-4ab0-a9d3-6df2df5d1d56","resolution":{"observed_at":"2026-05-22T21:42:10.761885Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"cited_work":{"arxiv_id":"2501.15953","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2501.15953","snapshot_observed_at":"2026-07-04T13:19:50.408368Z","title":"Understandinglongvideosviallm-poweredentityrelation graphs","venue":null,"work_id":"aa1b6435-f1e1-446e-95d9-5fd71504aaca","year":2025},"citing_paper":{"arxiv_id":"2606.26904","last_updated":"2026-06-25T11:37:51Z","snapshot_observed_at":"2026-08-08T03:59:36.206144Z","submitted_at":"2026-06-25T11:37:51Z","title":"Confidence-Aware Tool Orchestration for Robust Video Understanding","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-06-26T05:23:06.035263Z"},"links":{"cited_paper":"/paper/2501.15953","citing_paper":"/paper/2606.26904"},"observation_digest":"sha256:524c452ecf57477087e6c62daa192c57e525f78253ff502f5446e273999eb732","observation_id":"bd75b484-e11e-4381-9dd5-bfde17ce8ec6","resolution":{"observed_at":"2026-07-04T13:19:50.410008Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2501.15953/citation-record","integrity":"/paper/2501.15953/integrity","json":"/paper/2501.15953/citation-record.json","paper":"/paper/2501.15953"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2312.11805","last_updated":"2025-05-09T21:04:06Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-12-19T02:39:27Z","title":"Gemini: A Family of Highly Capable Multimodal Models","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.11805","snapshot_observed_at":"2026-08-10T13:53:25.613270Z","title":"Dai, Anja Hauth, Katie Millican, David Silver, Slav Petrov, Melvin Johnson, Ioannis Antonoglou, Julian Schrittwieser, Amelia Glaese, Jilin Chen, Emily Pitler, Timothy P","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:25.613270Z"},"links":{"cited_paper":"/paper/2312.11805","citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:e4a5d40c208e0029d164e85117f8fe17224c588b93bf2db433468b0f07f5aab8","observation_id":"640f364f-c076-4546-a8cb-4d6d130cd6dc","resolution":{"observed_at":"2026-08-10T13:53:25.613270Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.4324/9780429449642","metadata_source":"doi_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T13:53:28.061052Z","title":"Eysenck, and Michael C","venue":null,"work_id":"7d3563b3-4092-4aff-a7ea-c6e50a795a7e","year":2020},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:25.645657Z"},"links":{"citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:da912fade371898ea1dd6956cfb1b570e83849f55f7934546cd5bcd14a7b9950","observation_id":"e6c49138-a8f8-4bf6-a6ff-7c9cdcbe4e1a","resolution":{"observed_at":"2026-08-10T13:53:28.179555Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T13:53:30.559764Z","title":null,"venue":null,"work_id":"52d7c476-3dd1-42b6-b223-7776a9a377e3","year":2024},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:25.651514Z"},"links":{"citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:4f9b0a0d8b03e4ba6f366155d84f03e869251cbb56298c840fe3eeb97a4ba1cc","observation_id":"3127ddd4-dbaf-4536-95d3-79bf27828768","resolution":{"observed_at":"2026-08-10T13:53:30.563697Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.05861","last_updated":"2024-05-31T15:22:58Z","snapshot_observed_at":"2026-08-06T20:09:28.314739Z","submitted_at":"2024-02-08T17:50:22Z","title":"Memory Consolidation Enables Long-Context Video Understanding","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.05861","snapshot_observed_at":"2026-08-10T13:53:25.655542Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:25.655542Z"},"links":{"cited_paper":"/paper/2402.05861","citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:8d115ce688fd7de8552e321f528daa0d815b902c15175bc88fbbcbe0bcee2db5","observation_id":"2c95d65d-2a66-4685-9661-4f60708a3d5a","resolution":{"observed_at":"2026-08-10T13:53:25.655542Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T13:53:30.545442Z","title":"Language models are few-shot learners","venue":null,"work_id":"ac368e24-1c6f-44a3-9843-7bc5c8c1072c","year":2020},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:25.660671Z"},"links":{"citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:0633371a6e98359919897f70e6d9d2100dc483cbc0f629675e04a73d8a82ba05","observation_id":"320e655e-e181-46b4-977e-eef8aadda1d5","resolution":{"observed_at":"2026-08-10T13:53:30.550387Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T13:53:30.531072Z","title":null,"venue":null,"work_id":"2fa470ac-4b31-4359-b3f3-bd2fd25c3503","year":2022},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:25.665664Z"},"links":{"citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:aa74c93be2b4b3b320cb491a2d4a5218337483f64ba0769db448fc8fe882ca62","observation_id":"9ac6f942-8026-422e-8c42-8038f91e495a","resolution":{"observed_at":"2026-08-10T13:53:30.535255Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2409.10078","last_updated":"2025-04-05T07:55:46Z","snapshot_observed_at":"2026-08-05T16:04:30.737837Z","submitted_at":"2024-09-16T08:30:59Z","title":"3D-TAFS: A Training-free Framework for 3D Affordance Segmentation","version":5},"cited_work":{"arxiv_id":"2409.10078","doi":null,"metadata_source":"pith","pith_arxiv_id":"2409.10078","snapshot_observed_at":"2026-08-10T13:53:29.026150Z","title":"3D-TAFS: A Training-free Framework for 3D Affordance Segmentation","venue":"cs.RO","work_id":"b874dffa-1c7b-495b-bde2-a8696f69efa1","year":2024},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:25.670925Z"},"links":{"cited_paper":"/paper/2409.10078","citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:87657d3b23cc137410cf017349c8168ec5acbf6b2c126f8877c0f85e258a480f","observation_id":"6669b946-f188-420e-b071-7a2886d0e9d4","resolution":{"observed_at":"2026-08-10T13:53:29.030905Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T13:53:30.515558Z","title":null,"venue":null,"work_id":"a0891732-84a2-4326-91f8-a56c7f99e92b","year":null},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:25.675446Z"},"links":{"citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:51c5e886058aae07b5084b6694444aea8d7ca03dad6932ae36354ba718322292","observation_id":"39357be8-5158-416f-aaa0-f4ee9f8594e6","resolution":{"observed_at":"2026-08-10T13:53:30.520676Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T13:53:30.488342Z","title":null,"venue":null,"work_id":"b8cc20a8-a5c3-4e38-8f86-df3f78d38efb","year":2019},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:25.684872Z"},"links":{"citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:7cf53735d2e1729e0925765854b61fa4fbb929447ca0f2c2648f0d47fac0825c","observation_id":"300844ad-d084-46ba-acef-7b8486621220","resolution":{"observed_at":"2026-08-10T13:53:30.492518Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2407.21783","last_updated":"2024-11-23T23:27:33Z","snapshot_observed_at":"2026-08-10T16:40:37.411115Z","submitted_at":"2024-07-31T17:54:27Z","title":"The Llama 3 Herd of Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.21783","snapshot_observed_at":"2026-08-10T13:53:25.689501Z","title":"The llama 3 herd of models.arXiv preprint arXiv:2407.21783(2024)","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:25.689501Z"},"links":{"cited_paper":"/paper/2407.21783","citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:b3ca3664c69205c00ef529fd521e8946a22eeadf236d768e39bd37b13c0134cb","observation_id":"466fe527-7f49-46ec-8303-56ecc7ed1684","resolution":{"observed_at":"2026-08-10T13:53:25.689501Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2306.08640","last_updated":"2023-06-28T05:00:35Z","snapshot_observed_at":"2026-07-06T15:42:37.055739Z","submitted_at":"2023-06-14T17:12:56Z","title":"AssistGPT: A General Multi-modal Assistant that can Plan, Execute, Inspect, and Learn","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2306.08640","snapshot_observed_at":"2026-08-10T13:53:25.707276Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:25.707276Z"},"links":{"cited_paper":"/paper/2306.08640","citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:3e49f878467eeabaab1c11e7a8fc99a86d6bfcab67efe6e7654f837aecc31f85","observation_id":"29324fd6-77ff-4f66-a885-f6da3f28266e","resolution":{"observed_at":"2026-08-10T13:53:25.707276Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T13:53:30.473034Z","title":null,"venue":null,"work_id":"3b2ca438-f3e1-41e3-8111-95ed002261c9","year":2023},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:25.776143Z"},"links":{"citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:fc98c870c7bf1ef98a9e523afc93c0286e76e0fb4c3f0fc53da3736b259937ac","observation_id":"128877d1-dd0c-4924-9291-a0934fa0621e","resolution":{"observed_at":"2026-08-10T13:53:30.478820Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T13:53:30.446233Z","title":null,"venue":null,"work_id":"cfa2a958-7957-4274-94b0-001f623f6042","year":2024},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:25.805356Z"},"links":{"citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:185a2b098b88038455a249c2ae279d9c2abe901c8da0cbfdde6a7ced0d2be04c","observation_id":"02f23a97-a5b5-4beb-be29-13b012a7d234","resolution":{"observed_at":"2026-08-10T13:53:30.458048Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T13:53:30.432954Z","title":null,"venue":null,"work_id":"049a3a11-f832-4d8e-87b7-442c49cb52cc","year":null},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:25.809721Z"},"links":{"citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:6e2ef6229a05c895ab5a8afca46e69580e099cde384dbdde71621514ae0af478","observation_id":"8fbc003e-c6c0-492d-8cc0-1ab06e779850","resolution":{"observed_at":"2026-08-10T13:53:30.436922Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T13:53:30.323966Z","title":null,"venue":null,"work_id":"633c9c78-f7bf-4545-b019-43b370b858cb","year":2024},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:25.820222Z"},"links":{"citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:b1fd9f25bfccf525bcc45a13ac215bf5074cb5d36766b298fd36ea8b360840eb","observation_id":"e3a231bf-3301-4d8b-8cf8-17a3d3f36a9e","resolution":{"observed_at":"2026-08-10T13:53:30.410874Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T13:53:30.227685Z","title":"2020.spaCy: Industrial-strength Natural Language Processing in Python","venue":null,"work_id":"e134b566-ea9d-4d3f-86d4-043d167cacf6","year":2020},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:25.827402Z"},"links":{"citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:a126218d4b2ef9038f2cda8c5c6c1bb7f8a83d2bfb4c37bde1150aa04aa73e6e","observation_id":"2b634d23-b458-4b82-8568-515271a50d10","resolution":{"observed_at":"2026-08-10T13:53:30.273686Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1905.05143","last_updated":"2019-10-13T09:44:11Z","snapshot_observed_at":"2026-08-08T10:14:34.912925Z","submitted_at":"2019-05-13T16:57:40Z","title":"VideoGraph: Recognizing Minutes-Long Human Activities in Videos","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1905.05143","snapshot_observed_at":"2026-08-10T13:53:25.832166Z","title":"VideoGraph:RecognizingMinutes-LongHumanActivitiesinVideos","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:25.832166Z"},"links":{"cited_paper":"/paper/1905.05143","citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:832f34d81008da5c6cc90da0be77ca5fe435c318b43a137d3796f4632d1647e6","observation_id":"9cbce404-ed24-4d88-ace7-294d6fb274ba","resolution":{"observed_at":"2026-08-10T13:53:25.832166Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.1007/978-3-031-19833-5_6","metadata_source":"doi_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T13:53:27.840067Z","title":null,"venue":null,"work_id":"be919e3d-506b-4311-a790-827c4438cf31","year":2022},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:25.851284Z"},"links":{"citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:ffa9c9a1829b8f1b2088008dd2aa75276e9c1f5ecfda8c72783f271dea5ae7dd","observation_id":"215ca760-78a8-4d62-9c2a-aa7f7637f153","resolution":{"observed_at":"2026-08-10T13:53:27.952472Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2401.04088","last_updated":"2024-01-08T18:47:34Z","snapshot_observed_at":"2026-08-08T06:16:25.839566Z","submitted_at":"2024-01-08T18:47:34Z","title":"Mixtral of Experts","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.04088","snapshot_observed_at":"2026-08-10T13:53:25.929319Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:25.929319Z"},"links":{"cited_paper":"/paper/2401.04088","citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:847516400dc784fe4bd2d598651cfbbcec085e7870a9837f5c9ac5838f43407a","observation_id":"cd6f2d5e-37f8-4e56-bec3-802474644d49","resolution":{"observed_at":"2026-08-10T13:53:25.929319Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2311.08046","last_updated":"2024-04-05T15:21:09Z","snapshot_observed_at":"2026-07-06T16:47:17.136168Z","submitted_at":"2023-11-14T10:11:36Z","title":"Chat-UniVi: Unified Visual Representation Empowers Large Language Models with Image and Video Understanding","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2311.08046","snapshot_observed_at":"2026-08-10T13:53:26.001124Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:26.001124Z"},"links":{"cited_paper":"/paper/2311.08046","citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:5173ebb6b48bcd297e14f058c173491060710794e1b147747e40527c1f0c53b2","observation_id":"ee612375-e69e-4e98-9083-1b6d6ddf90a0","resolution":{"observed_at":"2026-08-10T13:53:26.001124Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T13:53:30.127709Z","title":null,"venue":null,"work_id":"f4699360-3cc4-42a0-b134-93646b600413","year":null},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:26.015648Z"},"links":{"citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:94f807d5562cec2ed84246c00743d484c402b139a44bc0772b7b09078d74bc6f","observation_id":"6189cbb0-815e-4909-97eb-d6afa1e3a613","resolution":{"observed_at":"2026-08-10T13:53:30.169790Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2301.11507","last_updated":"2023-01-27T03:00:43Z","snapshot_observed_at":"2026-07-06T14:45:08.248260Z","submitted_at":"2023-01-27T03:00:43Z","title":"Semi-Parametric Video-Grounded Text Generation","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2301.11507","snapshot_observed_at":"2026-08-10T13:53:26.037466Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:26.037466Z"},"links":{"cited_paper":"/paper/2301.11507","citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:90df4c65c4fab8f31da97d5d090d3fbb34878d6bf6b3417748da8cbb3364ee42","observation_id":"cf22dd37-c93c-48e2-b764-1e3c51ee1797","resolution":{"observed_at":"2026-08-10T13:53:26.037466Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2401.13649","last_updated":"2024-06-06T02:01:09Z","snapshot_observed_at":"2026-08-05T16:01:54.847394Z","submitted_at":"2024-01-24T18:35:21Z","title":"VisualWebArena: Evaluating Multimodal Agents on Realistic Visual Web Tasks","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.13649","snapshot_observed_at":"2026-08-10T13:53:26.042016Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:26.042016Z"},"links":{"cited_paper":"/paper/2401.13649","citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:d71ea8ef0dafbb22aa8682dd42febb98c8c4ef9906de7bc64ff60649acf337c7","observation_id":"0009b40c-e752-4051-837c-6193c9e11912","resolution":{"observed_at":"2026-08-10T13:53:26.042016Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2312.11897","last_updated":"2024-08-19T12:47:11Z","snapshot_observed_at":"2026-08-02T02:18:16.022611Z","submitted_at":"2023-12-19T06:42:47Z","title":"Text-Conditioned Resampler For Long Form Video Understanding","version":3},"cited_work":{"arxiv_id":"2312.11897","doi":null,"metadata_source":"pith","pith_arxiv_id":"2312.11897","snapshot_observed_at":"2026-08-10T13:53:28.846946Z","title":"Text-Conditioned Resampler For Long Form Video Understanding","venue":"cs.CV","work_id":"8a9faa14-564f-41b6-bd81-9655c24beb63","year":2023},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:26.047053Z"},"links":{"cited_paper":"/paper/2312.11897","citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:e27e41a26b52e549622075d67bca67f62b1125e318ef77063a6a5da6286b0f70","observation_id":"7edf763c-3bd1-4d39-96d4-6e1f06330c06","resolution":{"observed_at":"2026-08-10T13:53:28.894062Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T13:53:30.044355Z","title":"Berg, Mohit Bansal, and Jingjing Liu","venue":null,"work_id":"a6adbc2f-9a44-48a3-a318-fc8f19abb11a","year":2021},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:26.087430Z"},"links":{"citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:a21062cf36ce2c881d5d61921861f1f3fb7835356f6530964926c914c86d9343","observation_id":"e89fe509-c781-4998-9668-7b08b56263da","resolution":{"observed_at":"2026-08-10T13:53:30.049607Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.06355","last_updated":"2024-01-04T02:06:07Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-05-10T17:59:04Z","title":"VideoChat: Chat-Centric Video Understanding","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.06355","snapshot_observed_at":"2026-08-10T13:53:26.184775Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:26.184775Z"},"links":{"cited_paper":"/paper/2305.06355","citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:d445178da745dce4e1e6498ef05c919c2c7277788edcd48795216a201abd8fd2","observation_id":"7813284f-12ab-4ccb-bca4-624e0e947a5e","resolution":{"observed_at":"2026-08-10T13:53:26.184775Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2408.06327","last_updated":"2024-08-12T17:44:17Z","snapshot_observed_at":"2026-08-07T04:07:56.293097Z","submitted_at":"2024-08-12T17:44:17Z","title":"VisualAgentBench: Towards Large Multimodal Models as Visual Foundation Agents","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2408.06327","snapshot_observed_at":"2026-08-10T13:53:26.277715Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:26.277715Z"},"links":{"cited_paper":"/paper/2408.06327","citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:97a9f7bc96ec08d6179c066e804290f93797fd5d886fe27f2de1e2048cc407f7","observation_id":"f98e9771-1811-4dc9-a3d6-016a8914e6db","resolution":{"observed_at":"2026-08-10T13:53:26.277715Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T13:53:30.029972Z","title":null,"venue":null,"work_id":"b63165a5-82d5-449b-bf64-2be8bffa7d70","year":2022},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:26.307554Z"},"links":{"citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:aaf45344270bd7f58675d318259a6fc227acc1dfbb427f2ed3297918b058885b","observation_id":"2373e4e5-7639-4aa4-acc5-bce61185cd65","resolution":{"observed_at":"2026-08-10T13:53:30.034548Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T13:53:30.015785Z","title":null,"venue":null,"work_id":"5804d8cb-550a-44ce-aae3-cb9cd757ee56","year":null},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:26.332395Z"},"links":{"citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:4a738372685afe54cb1810f5ebc5c9ad616918890dc60beab2fcc4035ba022a1","observation_id":"4cd0e71e-d7f4-463c-9cf8-367b6a4b343c","resolution":{"observed_at":"2026-08-10T13:53:30.020031Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T13:53:30.002553Z","title":null,"venue":null,"work_id":"1e38b027-6bb1-4a73-9940-ec76c92d928b","year":null},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:26.341898Z"},"links":{"citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:292f0452195f776e3dc308051d0f7c610372ab1d264a84551f14d69d8a9663bc","observation_id":"5c6263a1-edbb-45e9-b4b3-9e49ca116f28","resolution":{"observed_at":"2026-08-10T13:53:30.006980Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T13:53:29.975587Z","title":null,"venue":null,"work_id":"28668b99-6466-4fdb-b32e-931d09843111","year":2024},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:26.440410Z"},"links":{"citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:27a335520011844c8327b3e184103901badc8d98fdaff92179af9e7f0e46c57d","observation_id":"44c4d684-28ca-4b83-a6f6-74e24c1a254d","resolution":{"observed_at":"2026-08-10T13:53:29.979805Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T13:53:29.962541Z","title":null,"venue":null,"work_id":"a20f86ef-8232-4e4f-b19c-e312f44640e3","year":2023},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:26.551880Z"},"links":{"citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:eb515784c53d4226da6e9be995c6c8b6ac1a421293f9376892fa7ae058b2dbdd","observation_id":"eba3d0f1-bfc7-4c76-965b-5111059b5823","resolution":{"observed_at":"2026-08-10T13:53:29.966957Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T13:53:29.940045Z","title":null,"venue":null,"work_id":"5eedfa08-ffb4-455b-a8c2-3220127636ff","year":2022},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:26.563690Z"},"links":{"citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:7231f0ce612758038042f85fe692352297c04376209eae2dc1fa92f66395ff04","observation_id":"c3145615-526b-489d-b89f-9ad523008533","resolution":{"observed_at":"2026-08-10T13:53:29.948032Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2303.08774","last_updated":"2024-03-04T06:01:33Z","snapshot_observed_at":"2026-08-07T07:30:12.213965Z","submitted_at":"2023-03-15T17:15:04Z","title":"GPT-4 Technical Report","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2303.08774","snapshot_observed_at":"2026-08-10T13:53:26.602192Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:26.602192Z"},"links":{"cited_paper":"/paper/2303.08774","citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:9af69eb52892a6f8594002add0764ad1468cba7ec7f041571fbf38169209e205","observation_id":"13626071-c269-4ddf-bdce-b15ad131ab9b","resolution":{"observed_at":"2026-08-10T13:53:26.602192Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T13:53:29.988361Z","title":null,"venue":null,"work_id":"8175ce4a-55d9-4709-89d7-1eb1f64d4563","year":2023},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:26.345883Z"},"links":{"citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:c3fc25a09a81f8b0b167873236d08b0863a3747f72810996b66a704768e14614","observation_id":"5036c1f5-70d2-4649-bbe2-216257593506","resolution":{"observed_at":"2026-08-10T13:53:29.992500Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2312.07395","last_updated":"2024-12-30T09:51:22Z","snapshot_observed_at":"2026-08-01T19:49:29.658899Z","submitted_at":"2023-12-12T16:10:19Z","title":"A Simple Recipe for Contrastively Pre-training Video-First Encoders Beyond 16 Frames","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.07395","snapshot_observed_at":"2026-08-10T13:53:26.614221Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:26.614221Z"},"links":{"cited_paper":"/paper/2312.07395","citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:def9e392684941f2c67bfe491664016ab69bd65c51f646c69b650ff2b62b525e","observation_id":"c27f1958-de19-4b44-893b-4c3f244f8c2d","resolution":{"observed_at":"2026-08-10T13:53:26.614221Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T13:53:29.857825Z","title":null,"venue":null,"work_id":"a43e9f2d-7e66-4eec-948c-e8c5c12f1a1c","year":2024},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:26.618853Z"},"links":{"citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:517904730a8625384e018bd6ed9de539ad8069820ff37c842d4b977678bc2140","observation_id":"422a1e60-49a8-4d15-8747-b425cf8210c9","resolution":{"observed_at":"2026-08-10T13:53:29.905539Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T13:53:29.747571Z","title":null,"venue":null,"work_id":"b463f780-e322-4d87-a819-038f42e52ebb","year":null},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:26.623240Z"},"links":{"citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:e8f207c4409f9835152b3bd36da738da9289eb2d62b6ccfaa2eb68500caab5e8","observation_id":"65ff43dd-adc6-4f62-9703-dfed2c2c452f","resolution":{"observed_at":"2026-08-10T13:53:29.807433Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2303.15389","last_updated":"2023-03-27T17:02:21Z","snapshot_observed_at":"2026-07-06T15:08:34.018146Z","submitted_at":"2023-03-27T17:02:21Z","title":"EVA-CLIP: Improved Training Techniques for CLIP at Scale","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2303.15389","snapshot_observed_at":"2026-08-10T13:53:26.632827Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:26.632827Z"},"links":{"cited_paper":"/paper/2303.15389","citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:71a9dc212b3176fb837f1f4b9f2dd51fdc63cb1ceca789bf02203376db2ec896","observation_id":"224d9117-2ec9-43b6-9cb5-ef7b87aee715","resolution":{"observed_at":"2026-08-10T13:53:26.632827Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T13:53:29.915222Z","title":"2021.Neuralscenegraphsfordynamicscenes.In ProceedingsoftheIEEE/CVF Conference on Computer Vision and Pattern Recognition","venue":null,"work_id":"365b94a7-868b-490e-825b-5401dc89f4fb","year":2021},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:26.609765Z"},"links":{"citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:13813ac8b961cdc17ce69872e042b4b6c34519da31b295a779f4e7060284523c","observation_id":"c79bf13f-eb34-48c3-9cf6-332d8513a3c0","resolution":{"observed_at":"2026-08-10T13:53:29.922737Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T13:53:29.522214Z","title":"Long-formvideo-languagepre-trainingwithmultimodal temporal contrastive learning.Advances in neural information processing systems35 (2022), 38032–38045","venue":null,"work_id":"4e76401f-d136-4046-925d-a1de274fed73","year":2022},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:26.730099Z"},"links":{"citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:c050fe9c095f764f106ae40f82fcf07f3a16a148b4f98a6a1782e8f7af4c602f","observation_id":"26e2478d-6f19-4c56-b19e-7dbcec9b800f","resolution":{"observed_at":"2026-08-10T13:53:29.575963Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T13:53:29.485960Z","title":"ViperGPT:VisualInfer- ence via Python Execution for Reasoning.Proceedings of IEEE International Conference on Computer Vision (ICCV)(2023)","venue":null,"work_id":"1903d6d5-b6f0-4f3a-ad52-d9fab7308bcf","year":2023},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:26.797411Z"},"links":{"citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:b20b3d4492064d7365700dcdd05e824e45bb64fb7f840f90a778556f2694d33a","observation_id":"82e660f7-298f-4498-bae6-4363df9c7a70","resolution":{"observed_at":"2026-08-10T13:53:29.490359Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T13:53:26.823762Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:26.823762Z"},"links":{"citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:1819d4504b32dfd2269b49605a877d48fab2220569b33b366f261fcc43bf6e56","observation_id":"de0e4ded-b903-46e6-86fa-ec42b0fc7f87","resolution":{"observed_at":"2026-08-10T13:53:26.823762Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.1007/978-3-031-73254-6_9","metadata_source":"doi_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T13:53:27.728247Z","title":null,"venue":null,"work_id":"35891a59-9966-4f1d-b767-4aeb44370736","year":2024},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:26.904761Z"},"links":{"citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:b1c458dd71e60a4c7a1e2a99a114f35daf4c4bb783860e7f1012d0d55af641b3","observation_id":"d8777168-88a3-4f98-a6f9-1b09a6e975ed","resolution":{"observed_at":"2026-08-10T13:53:27.783255Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T13:53:29.472654Z","title":null,"venue":null,"work_id":"980f15e8-627e-44ae-8aba-c5c60bf011c5","year":2025},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:26.995735Z"},"links":{"citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:a5f1ccfd8e3f24d6de0f3a1ea177a19e9c96edf9522b4af86b8886714b6523dd","observation_id":"de1d043e-85d8-4354-b83f-31a41d3d6ceb","resolution":{"observed_at":"2026-08-10T13:53:29.477061Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.04252","last_updated":"2024-02-06T18:59:48Z","snapshot_observed_at":"2026-08-09T19:13:53.289650Z","submitted_at":"2024-02-06T18:59:48Z","title":"EVA-CLIP-18B: Scaling CLIP to 18 Billion Parameters","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.04252","snapshot_observed_at":"2026-08-10T13:53:26.637867Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:26.637867Z"},"links":{"cited_paper":"/paper/2402.04252","citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:c1ea63ee29132adff14c2932efc38f408f82c753b79e361593242796f80f841c","observation_id":"6c137259-cc76-4756-84a7-bcbf761cb4c4","resolution":{"observed_at":"2026-08-10T13:53:26.637867Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2212.03191","last_updated":"2022-12-07T12:20:55Z","snapshot_observed_at":"2026-07-06T14:27:34.639236Z","submitted_at":"2022-12-06T18:09:49Z","title":"InternVideo: General Video Foundation Models via Generative and Discriminative Learning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2212.03191","snapshot_observed_at":"2026-08-10T13:53:27.070979Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:27.070979Z"},"links":{"cited_paper":"/paper/2212.03191","citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:5965c8ca4fec59e20bfdcdf0b0e829a38931258235d8db55cad16bacd10d69bf","observation_id":"6ead0cae-eefb-4090-bfb7-abaf41a42baa","resolution":{"observed_at":"2026-08-10T13:53:27.070979Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2312.05269","last_updated":"2024-11-05T22:08:14Z","snapshot_observed_at":"2026-08-08T23:23:37.595238Z","submitted_at":"2023-12-07T19:19:25Z","title":"LifelongMemory: Leveraging LLMs for Answering Queries in Long-form Egocentric Videos","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.05269","snapshot_observed_at":"2026-08-10T13:53:27.082257Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:27.082257Z"},"links":{"cited_paper":"/paper/2312.05269","citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:2ec7821565d3f3efb6b4b69d433b885e84d68aeae7322781087b03a6df153838","observation_id":"10ffbb1a-b09e-443d-a5fb-fa813c57cb19","resolution":{"observed_at":"2026-08-10T13:53:27.082257Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T13:53:29.458654Z","title":null,"venue":null,"work_id":"37477aab-05e6-4c5d-afa6-60b65cdff9fc","year":null},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:27.087149Z"},"links":{"citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:a67bd374155d7799c6c0bd171bcf64635fffc0ff70cd87f3b367ee9e30d7ba0c","observation_id":"422e099c-08e2-4eb1-a256-d1397d64637c","resolution":{"observed_at":"2026-08-10T13:53:29.463115Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T13:53:29.387579Z","title":null,"venue":null,"work_id":"24a3fa4b-f3c4-42ef-838a-fe7af5e00027","year":2018},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:27.095912Z"},"links":{"citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:50fac411491b7a18f77133a97ecbdc1ca6ff7a27f65a81a750ef69f80043ebe4","observation_id":"4285f38f-c875-45f2-beee-554c10f27dc5","resolution":{"observed_at":"2026-08-10T13:53:29.435136Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T13:53:27.100295Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:27.100295Z"},"links":{"citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:d702357c008bd3027dd4daad9fb251162abc689efddc3307bd9accbed5d22655","observation_id":"ef29ed8c-759b-42f0-8df7-5f558cbdac9e","resolution":{"observed_at":"2026-08-10T13:53:27.100295Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T13:53:27.054986Z","title":"SupervoxelAttentionGraphsforLong-Range Video Modeling","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:27.054986Z"},"links":{"citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:10edb9a0117e2ca039a820feae029d6cec125f57a7c7fc118812ab1e2b50c267","observation_id":"23536dde-920d-49fc-a051-497b37fb9464","resolution":{"observed_at":"2026-08-10T13:53:27.054986Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T13:53:29.216543Z","title":"Contrastive VideoQuestionAnsweringviaVideoGraphTransformer","venue":null,"work_id":"0d305c24-88bb-43f8-b9c9-ffd66801829d","year":2023},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":53,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:27.109107Z"},"links":{"citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:bbf78b04c63a14d99b87ebaff8c8a3b0fb27d7c6286eca95c43b5b44b7d176b1","observation_id":"88979a82-fc07-4de1-8a2a-7ad425ac1623","resolution":{"observed_at":"2026-08-10T13:53:29.244489Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T13:53:27.113594Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":54,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:27.113594Z"},"links":{"citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:35d3bda7e5ce11164368b9becd0c63020ac4bacc2889c67a7af2178e3ef86481","observation_id":"230b071f-781f-4083-8216-b31807153408","resolution":{"observed_at":"2026-08-10T13:53:27.113594Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T13:53:29.174773Z","title":null,"venue":null,"work_id":"49958d95-bf99-4b09-907e-b10d93b00e6c","year":null},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:27.130688Z"},"links":{"citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:40d335222f151dfea3ea4953b307d47421afb46dc148fe34ee291314d9974594","observation_id":"4be26044-3c1c-4104-a87b-21968e099667","resolution":{"observed_at":"2026-08-10T13:53:29.179754Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T13:53:29.444575Z","title":"InEuropean Conference on Computer Vision","venue":null,"work_id":"66336aec-0fb8-443a-9090-4ecd708666d1","year":null},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":56,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:27.091503Z"},"links":{"citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:051418c1cf1082071b5ed79e1590358796fcb553afc240d14c2df40be29b7289","observation_id":"ef77d406-c10d-4f68-9b55-2f7fc6b526cb","resolution":{"observed_at":"2026-08-10T13:53:29.449412Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T13:53:29.131591Z","title":null,"venue":null,"work_id":"7f6b25fb-5721-42b9-b106-c317dfc6c655","year":2022},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":57,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:27.377351Z"},"links":{"citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:69d17548e3ab017b987f7d959b6a7eb3805582cea198662dfe689e90351810e7","observation_id":"905bd51d-feab-415f-b90e-0dceebb5d569","resolution":{"observed_at":"2026-08-10T13:53:29.136774Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2007.03626","last_updated":"2020-07-07T17:00:11Z","snapshot_observed_at":"2026-08-09T15:30:59.352139Z","submitted_at":"2020-07-07T17:00:11Z","title":"What Gives the Answer Away? Question Answering Bias Analysis on Video QA Datasets","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2007.03626","snapshot_observed_at":"2026-08-10T13:53:27.394918Z","title":null,"venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":58,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:27.394918Z"},"links":{"cited_paper":"/paper/2007.03626","citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:a560c783b24889090757cefba316f9cc689f239b030294893a7960603c4bd244","observation_id":"49c55f0b-8847-4cda-b056-7ee138c2cdc1","resolution":{"observed_at":"2026-08-10T13:53:27.394918Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T13:53:29.300553Z","title":null,"venue":null,"work_id":"6c21c72e-1036-4991-962f-52d197e956a6","year":2021},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":59,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:27.104915Z"},"links":{"citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:1f7161bbd134d5b6aa04a4f0e0227ec16325c3ecafcc66a8aaecc90cb76e96b1","observation_id":"5e3b4685-6f55-4505-a1a3-48a16c265c89","resolution":{"observed_at":"2026-08-10T13:53:29.334636Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T13:53:29.102148Z","title":null,"venue":null,"work_id":"1d552522-ac23-4c5e-8498-0863a9914dc9","year":2023},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":60,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:27.449721Z"},"links":{"citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:36ab0f1e5dcb0dcebe7470eef40e869f6d528fe57c4ee9f6a41396df8534137a","observation_id":"974a0545-60f5-4012-8f34-a73680b8b0b5","resolution":{"observed_at":"2026-08-10T13:53:29.106995Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T13:53:29.087023Z","title":null,"venue":null,"work_id":"a5d643ca-509f-409b-9fcf-8281b2629159","year":2024},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":61,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:27.457734Z"},"links":{"citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:1d7cfb9a7451b37f2896a80d480b4cd829e4b97c48cdc0b820691a9ac47099cb","observation_id":"20743c72-ae23-4d04-a4f1-2cefbdcc3e36","resolution":{"observed_at":"2026-08-10T13:53:29.091315Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T13:53:29.071664Z","title":null,"venue":null,"work_id":"9f0dba7d-ec53-4cdf-bf52-9aff9fd7e5fd","year":2018},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":62,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:27.462352Z"},"links":{"citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:38f6aff1bbc9712aa1d1f11ef23085ae05effadbce055f940b728900c07da7d6","observation_id":"290a944c-6b58-4710-8d28-f75bda58836f","resolution":{"observed_at":"2026-08-10T13:53:29.076624Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T13:53:29.161086Z","title":"InProceedings of the IEEE/CVF international conference on computer vision","venue":null,"work_id":"f9edaf5f-b3b0-4421-b595-3ea3afc5cc65","year":null},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":63,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:27.184683Z"},"links":{"citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:ce9745fe8a4209d4abb60f5e42cef82be38893387ef67df3767283732e38089b","observation_id":"48bd4aa9-67ab-4ade-b354-45c54c23eac9","resolution":{"observed_at":"2026-08-10T13:53:29.165531Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T13:53:29.147439Z","title":null,"venue":null,"work_id":"021465a0-940c-4e77-a9e4-5bdd614e13ce","year":2022},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":64,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:27.281551Z"},"links":{"citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:8fe6b537529bff1babb280ea584d48f5e229067fc3bdd2c7f1a9ddba76d93e38","observation_id":"3545dd8d-3d24-4abc-aeca-f588eefb5ba9","resolution":{"observed_at":"2026-08-10T13:53:29.151973Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.08085","last_updated":"2024-06-30T05:39:46Z","snapshot_observed_at":"2026-08-03T12:09:02.990658Z","submitted_at":"2024-06-12T11:07:55Z","title":"Flash-VStream: Memory-Based Real-Time Understanding for Long Video Streams","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.08085","snapshot_observed_at":"2026-08-10T13:53:27.476040Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":65,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:27.476040Z"},"links":{"cited_paper":"/paper/2406.08085","citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:a6532b81725d9199a16b987943349df5269d1bc8f78d53f31f4a0f98141e0e7e","observation_id":"e17d6756-a94b-4618-9daa-d1527357695a","resolution":{"observed_at":"2026-08-10T13:53:27.476040Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T13:53:29.057924Z","title":null,"venue":null,"work_id":"ebb0b386-bcf1-41eb-9588-edf570f26bdc","year":2023},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":66,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:27.583766Z"},"links":{"citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:71abfaf44eba4ceda1c6300eb25b1e32bf749ee6377408402ddcb075b09173ad","observation_id":"5b63260a-18df-4295-adee-28fbaaf9c22e","resolution":{"observed_at":"2026-08-10T13:53:29.061904Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T13:53:29.116503Z","title":null,"venue":null,"work_id":"411a78fb-2ef2-45a1-ba6b-706ce0ba57bf","year":2023},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":67,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:27.440666Z"},"links":{"citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:9029d76ac7dfb3354687b6cd221934bc90cc42a4f3d4824e0ec0aa76749df5e9","observation_id":"2a899241-5f13-4d2c-85b6-4d7cffd0dfcc","resolution":{"observed_at":"2026-08-10T13:53:29.120808Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2312.17235","last_updated":"2024-10-10T05:17:00Z","snapshot_observed_at":"2026-08-10T11:57:05.068914Z","submitted_at":"2023-12-28T18:58:01Z","title":"A Simple LLM Framework for Long-Range Video Question-Answering","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.17235","snapshot_observed_at":"2026-08-10T13:53:27.471264Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":72,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:27.471264Z"},"links":{"cited_paper":"/paper/2312.17235","citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:c13d26197d7b5e2fcc73fc9be2f36a64ec2884e19f75ff9760e2dae1a96dd749","observation_id":"bd39f4d2-65c3-45ea-a8ca-545b572d313c","resolution":{"observed_at":"2026-08-10T13:53:27.471264Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T13:53:30.059804Z","title":"InProceedings of the 28th ACM International Conference on Multimedia","venue":null,"work_id":"761c600c-5b16-4da6-8b87-fee126887691","year":null},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":2020,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:26.032353Z"},"links":{"citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:0a0417a9c7f3613b1424e51d47d7d1c34be1f6d6ddc18cc32053060a0c0ddc7a","observation_id":"285fce16-8799-4fed-88f3-e8aa0438cde6","resolution":{"observed_at":"2026-08-10T13:53:30.064857Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T13:53:30.420592Z","title":"In Proceedings of the IEEE/CVF Conference on Computer Vision and Pattern Recognition","venue":null,"work_id":"f1384f6c-d744-442b-8916-9ff82a65f52d","year":null},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":2021,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:25.814974Z"},"links":{"citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:aece7ebce44f91bc0db9c4b91d264e482b1e46fc32bfafa8112e490a59913834","observation_id":"057e6505-f3fd-4054-98a6-5acfea757b10","resolution":{"observed_at":"2026-08-10T13:53:30.424333Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2312.08870","last_updated":"2025-03-03T17:58:12Z","snapshot_observed_at":"2026-08-10T13:01:49.384599Z","submitted_at":"2023-12-12T09:47:59Z","title":"Vista-LLaMA: Reducing Hallucination in Video Language Models via Equal Distance to Visual Tokens","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.08870","snapshot_observed_at":"2026-08-10T13:53:26.337203Z","title":"arXiv preprint arXiv:2312.08870(2023)","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":2023,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:26.337203Z"},"links":{"cited_paper":"/paper/2312.08870","citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:a707d15c9c43ff01ac0dfb5e96295de4314032c3cf175c3f50d601d11b86472b","observation_id":"dc9704bb-eb89-4c0f-986c-d7a75ab2bbec","resolution":{"observed_at":"2026-08-10T13:53:26.337203Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T13:53:29.637712Z","title":"18221–18232","venue":null,"work_id":"b22687fb-a5eb-44da-8dbf-ddf5fb9e5981","year":null},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":2024,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:26.628315Z"},"links":{"citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:cd86ae87aa77614f4b8c0e77bb6e8a316ddf8b68b022afcecef9bcbb3cbc782b","observation_id":"e98c36c5-db03-4a6b-a269-3c4b5580cbc7","resolution":{"observed_at":"2026-08-10T13:53:29.689978Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-10T13:53:30.501913Z","title":"InEuropean Conference on Computer Vision","venue":null,"work_id":"a66e8f9e-9930-4bea-99b4-8d677a66833f","year":null},"citing_paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs","version":1},"reference_index":2025,"source":"pdf_text","source_observed_at":"2026-08-10T13:53:25.680180Z"},"links":{"citing_paper":"/paper/2501.15953"},"observation_digest":"sha256:bdf82730ff7bf6c8a4267258917542bc531457606439d9e233913eac8384bb4a","observation_id":"97e0d930-6f8e-4805-928e-c790a5d82061","resolution":{"observed_at":"2026-08-10T13:53:30.506053Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2501.15953","last_updated":"2025-01-27T10:57:24Z","latest_version":1,"primary_category":"cs.IR","snapshot_observed_at":"2026-08-10T13:47:43.852671Z","submitted_at":"2025-01-27T10:57:24Z","title":"Understanding Long Videos via LLM-Powered Entity Relation Graphs"},"reference_resolution":{"displayed":73,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":55,"verified_exact":5,"verified_fuzzy":13},"total_outbound_references":73},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"thesis":"As of 11 August 2026, this Paper Citation Record lists 73 of 73 outbound references and 2 inbound Pith citation observations for arXiv:2501.15953."}