{"as_of":"2026-08-01T22:17:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:0fe30d73e548f2376e0a43035477a1c302cc8015cf5a2a03f3f0ed21f92ca0b3","coverage":[{"denominator":89,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":89,"source":"paper_references, paper_reference_links","source_observed_at":"2026-05-08T12:28:35.008604Z","state":"measured"},{"denominator":92,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":92,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-01T06:32:01.292127+00:00","state":"measured"},{"denominator":3,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":3,"source":"paper_references, paper_reference_links","source_observed_at":"2026-07-31T04:55:07.002233Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"pith","source_observed_at":"2026-07-03T16:38:39.514289Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"cited_work":{"arxiv_id":"2605.06537","doi":null,"metadata_source":"pith","pith_arxiv_id":"2605.06537","snapshot_observed_at":"2026-07-03T16:38:39.514289Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","venue":"cs.CV","work_id":"ad226d8b-9b32-43cb-bee6-e0d5069860a2","year":2026},"citing_paper":{"arxiv_id":"2607.01751","last_updated":"2026-07-02T06:07:44Z","snapshot_observed_at":"2026-07-07T00:07:19.026006Z","submitted_at":"2026-07-02T06:07:44Z","title":"MedStreamBench: A Time-Aware Benchmark for Streaming and Proactive Medical Video Understanding","version":1},"reference_index":2,"source":"arxiv_source","source_observed_at":"2026-07-03T16:37:09.666491Z"},"links":{"cited_paper":"/paper/2605.06537","citing_paper":"/paper/2607.01751"},"observation_digest":"sha256:8fab0183702d0558b9ef7a105cfbb3de3ac290f9dba61c0175c9bbacb2c8624f","observation_id":"2f7d6379-5e45-4668-92e5-300186793dfa","resolution":{"observed_at":"2026-07-03T16:38:39.515751Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2605.06537","snapshot_observed_at":"2026-07-14T14:52:42.813309Z","title":null,"venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2607.09880","last_updated":"2026-07-10T18:13:54Z","snapshot_observed_at":"2026-07-30T01:49:09.037227Z","submitted_at":"2026-07-10T18:13:54Z","title":"CLIR-Bench: Benchmarking Multimodal Question Answering over Irregular Clinical Time Series","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-07-14T14:52:42.813309Z"},"links":{"cited_paper":"/paper/2605.06537","citing_paper":"/paper/2607.09880"},"observation_digest":"sha256:89e3decefdb60b943f06569387be37d6c5458bdc92a7b893a6902fe6d5729b97","observation_id":"c44e243c-2bba-46b7-8550-f3038b76a576","resolution":{"observed_at":"2026-07-14T14:52:42.813309Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2605.06537","snapshot_observed_at":"2026-07-31T04:55:07.002233Z","title":"Alessandro Favero, Luca Zancato, Matthew Trager, Siddharth Choudhary, Pramuditha Perera, Alessandro Achille, Ashwin Swaminathan, and Stefano Soatto","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28516","last_updated":"2026-07-30T16:55:00Z","snapshot_observed_at":"2026-08-01T21:50:07.488793Z","submitted_at":"2026-07-30T16:55:00Z","title":"Beyond Frame Selection: Generative Latent Evidence Aggregation for Long-Video Understanding","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-07-31T04:55:07.002233Z"},"links":{"cited_paper":"/paper/2605.06537","citing_paper":"/paper/2607.28516"},"observation_digest":"sha256:418159b7c6093928c70ca8d723c96ef591889c3a5971093f2b5077f1121b96be","observation_id":"23c65698-d66c-49e6-bdc3-1f9c9aa562da","resolution":{"observed_at":"2026-07-31T04:55:07.002233Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2605.06537/citation-record","integrity":"/paper/2605.06537/integrity","json":"/paper/2605.06537/citation-record.json","paper":"/paper/2605.06537"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Pixel-wise recognition for holistic surgical scene understanding","venue":null,"work_id":"531624f2-8157-46a6-a222-19b01bbdeb65","year":2025},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":1,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:146b9449ca2252ce561d63c2e59becfce303e2005d6cbef19b8d050b821cd3a8","observation_id":"a718ec8a-ec67-4d4a-8aac-406108aa6e9d","resolution":{"observed_at":"2026-05-26T13:17:49.931456Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Real-colon: A dataset for developing real-world ai applications in colonoscopy","venue":null,"work_id":"c26c1da7-e6f1-4499-9d2d-e7c99d0f3275","year":2024},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":4,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:2666ce72a0a8c147716df683b29816cebffe6d09fb95b86723a8965327fb078c","observation_id":"f563f633-34ba-4454-98ce-52fc6058e9bc","resolution":{"observed_at":"2026-05-26T13:17:50.026003Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Artificial intelligence for surgical scene understanding: a systematic review and reporting quality meta-analysis","venue":null,"work_id":"7d8a12cc-bf5c-43ea-9071-328d30b24a72","year":2025},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":5,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:bc4fed6c6277cbc77ad0379321616da84eddbd2675a21ffeb7e5fb3de2a57308","observation_id":"a6a44dd9-4f53-4e83-b302-9b027de06801","resolution":{"observed_at":"2026-05-26T13:17:50.033555Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Internvl: Scaling up vision foundation models and aligning for generic visual-linguistic tasks","venue":null,"work_id":"13871f3e-9e66-4ca8-97a0-a787cfdbab37","year":2024},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:9bfe73d9ae56d612ee24aed671c63ee2ef3039880bfa4cf248603c0e896ef3da","observation_id":"a5511266-0db8-4b9c-a8ed-e5a7f73f6a87","resolution":{"observed_at":"2026-05-26T13:17:50.040290Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Corley, Christopher D","venue":null,"work_id":"e68ab848-8163-4134-a5c4-96b91e91ce25","year":2014},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":10,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:3e1cdf9de56946ce45c53548d486be36b61fdf62ea995aaeaac9a523fd3ca58f","observation_id":"997f73c3-bf99-407e-ba97-005929064c7b","resolution":{"observed_at":"2026-05-26T13:17:50.037030Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Lovat, Peter Mountney, et al","venue":null,"work_id":"dd812922-b681-479a-81dd-bd0dab921c02","year":2023},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:6f6935141c76006a3f346bd976ae5d3e456da2098f2284b0839c416ca0640bb0","observation_id":"c660e2ec-7425-4a18-bab7-29eedcb6c27d","resolution":{"observed_at":"2026-05-26T13:17:50.054794Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Deep learning in surgical workflow analysis: A review of phase and step recognition","venue":null,"work_id":"3f306160-555f-44ce-8317-d1dc90313571","year":2023},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":12,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:c94db19542fbfb58f1e1d8f318541f9643367929d57a1662dd203fc5c7900b28","observation_id":"b0c069f8-c698-4bd8-ba6b-d32eb6c5f5e4","resolution":{"observed_at":"2026-05-26T13:17:50.072712Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Standardized cine-loop documentation in abdominal ultrasound facilitates offline image interpretation","venue":null,"work_id":"1d8680a3-8963-4b63-8fb2-b9afe622efbc","year":2015},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":13,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:84dd7a1431c6caf642650f272d2e71d425ea8f4a97aed891a125eb2a4f0a84fa","observation_id":"8c27e3e4-d092-4c5e-ab0c-b9855d2689d9","resolution":{"observed_at":"2026-05-26T13:17:49.998353Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"We're expanding our Gemini 2.5 family of models","venue":null,"work_id":"bc096867-75a5-491f-a596-605f3a1b05a6","year":2025},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":14,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:8e88b3d64ffbb641b55e5b36c36ec123d6d1cd6920398a14aa460a87af2a383f","observation_id":"bd9fd3c4-93bb-4819-99a7-b4910ec7fdf6","resolution":{"observed_at":"2026-05-26T13:17:50.061692Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Yadlapati, Mark Benson, Andrew J","venue":null,"work_id":"cf940d7a-6e8d-40d2-9c12-e64a0976754a","year":2019},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":15,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:5fdaa30c7e6f13e28cce7e701439b9ce7d8347d61d74eda7d7441c2d1560afd3","observation_id":"0843f89c-cb6e-47de-bfaa-409b970b8dbe","resolution":{"observed_at":"2026-05-26T13:17:50.057687Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Real-time ultrasound demonstration of uterine isthmus contractions during pregnancy","venue":null,"work_id":"f1aa16fe-bc2c-4f16-8442-f476a0c42da8","year":2024},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":16,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:431c0cc1c15784e14849f59c0f0dfa9b83f31e53a970d84e3f81a058582a2c21","observation_id":"fc665a48-946a-4192-bd85-39d52d93525c","resolution":{"observed_at":"2026-05-26T13:17:49.896244Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Video-mme: The first-ever comprehensive evaluation benchmark of multi-modal llms in video analysis","venue":null,"work_id":"803e6ad7-22f7-4d45-b1f9-274b99ad601a","year":2025},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":17,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:2e653fb0749d2e35313256445870b02af3816cc0d5b663c22c07f6ff59ba1d43","observation_id":"beed10f3-921c-49ab-a8f0-eb48e4a5469c","resolution":{"observed_at":"2026-05-26T13:17:49.944640Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Khan, Sophia Bano, Hani J","venue":null,"work_id":"29e0cdcb-9594-42a1-8edd-d1287376bae7","year":2024},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":18,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:6d01509d0e7f1242ebdda1142d445821d97481981ceca9c1edadab0403702e8f","observation_id":"808cda99-e8ae-4155-8085-8e4590fd4dd6","resolution":{"observed_at":"2026-05-26T13:17:49.889557Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Ophnet: A large-scale video benchmark for ophthalmic surgical workflow understanding","venue":null,"work_id":"5fa69f59-e0e3-45ce-83f6-8745152f6a54","year":2024},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":19,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:84db0114bc1de5360ee11605e53753d2ddd9feb367dc849c6bc5b91bbb3cb4a0","observation_id":"2735fd1d-0f0f-4a4c-a76d-2348acc336df","resolution":{"observed_at":"2026-05-26T13:17:49.910547Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"A new dataset and versatile multi-task surgical workflow analysis framework for thoracoscopic mitral valvuloplasty","venue":null,"work_id":"db2627f8-300c-47f5-92bc-d338234723a5","year":2025},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":21,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:85e5e8d00c5450b793122b75d214b1f42413170eb22b32648a7809407e456e71","observation_id":"8569740b-622e-44c3-ac3c-b35824ea44d9","resolution":{"observed_at":"2026-05-26T13:17:50.008378Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Lang, Luigi P","venue":null,"work_id":"06bf5224-f3a6-4b29-8649-5ed6a062e65a","year":2015},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":22,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:e6a71c396e1ee6122a62bc19308ba4fe38f7d98f91d595d2a758407629d9d400","observation_id":"df7d7a5c-411f-4533-88b7-7f0294074669","resolution":{"observed_at":"2026-05-26T13:17:50.064944Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"e l L Lavanchy, Sanat Ramesh, Diego Dall’Alba, Cristians Gonzalez, Paolo Fiorini, Beat P M \\","venue":null,"work_id":"5123b9ce-b6b7-41b0-9441-d6b6e365e521","year":2024},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":23,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:77772b62661ab1ee096c47a60257e9968cb312ad40e36a9aa9bdec8106a041fa","observation_id":"af8f6b2c-b691-4372-b6e0-76162ef5a6c2","resolution":{"observed_at":"2026-05-26T13:17:50.068892Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Galar-a large multi-label video capsule endoscopy dataset","venue":null,"work_id":"5bd38c2c-119d-4cd0-85c4-76b53f66ea5c","year":2025},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":24,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:b032950fd216cd52f591b97f15cbf03599acb2e85bd95efe122ed686a4ff3122","observation_id":"436aaeef-a3d2-498b-ae19-6c9e0914e19a","resolution":{"observed_at":"2026-05-26T13:17:50.084053Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Detecting moments and highlights in videos via natural language queries","venue":null,"work_id":"617e507f-0460-4945-9a75-2e9f65fc0026","year":2021},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":25,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:3ec8aeb2cb89c95357d2151a06557d105c06267afc069ae0d2d7054df6982d71","observation_id":"91a23ac5-3689-481b-9ff8-1b88c8632d2f","resolution":{"observed_at":"2026-05-26T13:17:50.061470Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Mvbench: A comprehensive multi-modal video understanding benchmark","venue":null,"work_id":"0136fda9-92ba-414e-a9c1-da164a698c63","year":2024},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":26,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:4809c0ac397e459a73b70062c5ca76dcc9371d3e9cd340d64c6c4e8711e872da","observation_id":"f30fa90e-da7a-4231-a130-10a5f5784162","resolution":{"observed_at":"2026-05-26T13:17:50.039990Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-05T05:20:43.332327Z","title":"Surgpub-video: A comprehensive surgical video framework for enhanced surgical intelligence in vision-language model","venue":null,"work_id":"f055f872-7a12-44ab-97cd-92e6b8b0681f","year":2026},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":28,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:fca32ce3b8be91d8a392d9e0ffcc5f00004b8ceb262ddfcd2412d34f0137ca4e","observation_id":"0a049a47-c8c3-4993-8df4-4fff880a1f8a","resolution":{"observed_at":"2026-05-26T13:17:50.026342Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Kimi K2.5 : Visual agentic intelligence","venue":null,"work_id":"91e71a07-d4d5-4018-b056-a177f11dee83","year":2026},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":30,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:fd0d7276eb1bdfa31ee1e1029377f4b4899fb8651b89d61ad9e8408c7a2d9bd6","observation_id":"6d50a534-01eb-4938-8243-4cc1e49f63a1","resolution":{"observed_at":"2026-05-26T13:17:50.036588Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Introducing GPT -5.4","venue":null,"work_id":"3b93e405-1662-4b32-9fe3-ed1f37695748","year":2026},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":31,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:4dd73d4900360a48258e08095e473bc209ee96f25af613234c33473c5dc45236","observation_id":"90f0868a-1205-4eee-b52b-5f921a9e18c9","resolution":{"observed_at":"2026-05-26T13:17:50.015573Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Introducing GPT-5.4 mini and nano","venue":null,"work_id":"4601919a-2159-4805-942d-328e1bda6e83","year":2026},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":32,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:f89db136f270801df494ce0f95aeb6777c9c851dae5e9e46a234b46dba979afc","observation_id":"470f0ce3-485c-403f-b5aa-c26c13b7391a","resolution":{"observed_at":"2026-05-26T13:17:50.033216Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Qwen3.5 : Towards native multimodal agents","venue":null,"work_id":"a033766f-4a1f-404b-b1c4-be1e244e6acd","year":2026},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":34,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:76813a05f38305838cc5ba32f074c84302ac07b1a8a480a771c7311a0ae8048b","observation_id":"9b3d469a-5d8c-4a5c-81b4-d98f790d7574","resolution":{"observed_at":"2026-05-26T13:17:50.050506Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Qwen3.6 : Advancing multimodal intelligence","venue":null,"work_id":"54cb807f-dacf-490d-91e2-fd551122cbe0","year":2026},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":35,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:debfd1ad84e128aa2c0578df5d1e96e41f97c656349d2bee44e4c6b86c7b2195","observation_id":"b6ae069b-b226-4a0d-844f-2764e7500246","resolution":{"observed_at":"2026-05-26T13:17:50.080383Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Rex, Joseph C","venue":null,"work_id":"186ec27f-c96b-4233-9af9-afc0e602312d","year":2024},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":36,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:7834bd294a4772fb7984c1da823b5f1f2d8f193d7cf5c0122b89c83d20bdc730","observation_id":"550525d8-3425-4b87-87d4-42c7efe71373","resolution":{"observed_at":"2026-05-26T13:17:50.087126Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"648046bd-02cf-45a4-b18d-9b33d730e39b","year":2022},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":37,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:d674a9fa802d8a52d670db207d3b2d4854cfbb2faa34415bf9753d78cda7e399","observation_id":"acef1349-8ae2-4bd2-8b42-db2413dc9e83","resolution":{"observed_at":"2026-05-26T13:17:50.012116Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Medgrpo: Multi-task reinforcement learning for heterogeneous medical video understanding","venue":null,"work_id":"a78a4bc7-034b-410a-8e10-27356400072b","year":2026},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":38,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:9edcb40a9862f16b2a77f9d861c3b531b0ad2db2c744d491a158c88789679549","observation_id":"0577351b-1846-49b2-ad82-dc305ac23921","resolution":{"observed_at":"2026-05-26T13:17:50.023004Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Adaptive keyframe sampling for long video understanding","venue":null,"work_id":"4222cc7b-eede-4721-a3db-883caf034c7e","year":2025},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":39,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:2dd6400019ad2f44b3862a847616db0a24b55fb40ed912cf6987b4cc461730f1","observation_id":"438f8796-466b-4cc1-949e-7d07169cc421","resolution":{"observed_at":"2026-05-26T13:17:50.029670Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Mimic-iv-echo-ext-mimicechoqa: A benchmark dataset for echocardiogram-based visual question answering","venue":null,"work_id":"39fdce29-c099-4fdf-b086-c6ec3ba11fa6","year":2025},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":41,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:5d6d326ec5ddd574959934c78c7dbf1e8f0db41b60b23743137617b02f09ca50","observation_id":"46c8d370-4f9c-4591-8304-160694a87846","resolution":{"observed_at":"2026-05-26T13:17:50.043501Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Gemini 3.1 pro: A smarter model for your most complex tasks","venue":null,"work_id":"d9d9fdcc-e2a3-41af-b54b-83ceaf8c25c7","year":2026},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":42,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:bfdffa08dbc840e1be63aea69285a47fd62c8b74b07cb678952673b8511b01e4","observation_id":"bc06ba78-eb71-4e1b-91e9-1367b8f3629b","resolution":{"observed_at":"2026-05-26T13:17:49.987435Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Lvbench: An extreme long video understanding benchmark","venue":null,"work_id":"55f0d6b2-fdc9-417a-80d4-00252682e539","year":2025},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":44,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:dd2b7ff319ff882355410dad41bb2886b74ee760baeb81016dfae6fb1a017cd9","observation_id":"4cf592f2-3312-4997-9792-e54622d60381","resolution":{"observed_at":"2026-05-26T13:17:49.990767Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Autolaparo: A new dataset of integrated multi-tasks for image-guided surgical automation in laparoscopic hysterectomy","venue":null,"work_id":"e9f8bc62-9e21-4ebf-b15c-da3fc61ca370","year":2022},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":46,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:f71128e1553a7a4298212e5b8a619d643c690dc98ed523207b60ebc09318351f","observation_id":"ef52dbea-d4d8-4cf1-b07a-cad1968a9b09","resolution":{"observed_at":"2026-05-26T13:17:50.002139Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Longvideobench: A benchmark for long-context interleaved video-language understanding","venue":null,"work_id":"785f0959-d7b4-4290-bd97-6368ad758917","year":2024},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":47,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:6c684208fa8e73977554ab9572814d0d037480cae1aba951837f9f77c6d2dd2d","observation_id":"9c7a4c55-4aec-4aa3-8b1a-2aff74e4e4d8","resolution":{"observed_at":"2026-05-26T13:17:49.966155Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2408.01800","last_updated":"2024-08-03T15:02:21Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-08-03T15:02:21Z","title":"MiniCPM-V: A GPT-4V Level MLLM on Your Phone","version":1},"cited_work":{"arxiv_id":"2408.01800","doi":null,"metadata_source":"pith","pith_arxiv_id":"2408.01800","snapshot_observed_at":"2026-07-10T11:37:03.161139Z","title":"MiniCPM-V: A GPT-4V Level MLLM on Your Phone","venue":"cs.CV","work_id":"0f06e436-0c76-4e3c-be5e-6168f6bc4336","year":2024},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":49,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"cited_paper":"/paper/2408.01800","citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:60e653a158a8b89cdac9f12b1f60ef409620f34f0d0679e3a5629378154a167d","observation_id":"b3ecea8b-ce60-4310-b7ab-6efa33b039b2","resolution":{"observed_at":"2026-05-11T19:16:07.486569Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-05T07:40:47.902764Z","title":"Proceedings of the IEEE/CVF Conference on Computer Vision and Pattern Recognition , pages=","venue":null,"work_id":"c7684330-0e20-4ced-aeea-ac696ea54900","year":null},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":54,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:ab941fa74b7f981dcec910f697a63c858429d021ce39e642badaaf1af40ea104","observation_id":"f62cf355-6bef-4d82-8735-97fd6c8a9661","resolution":{"observed_at":"2026-05-26T13:17:49.986798Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2508.18265","last_updated":"2025-08-27T14:39:45Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-08-25T17:58:17Z","title":"InternVL3.5: Advancing Open-Source Multimodal Models in Versatility, Reasoning, and Efficiency","version":2},"cited_work":{"arxiv_id":"2508.18265","doi":"10.48550/arxiv.2508.18265","metadata_source":"pith","pith_arxiv_id":"2508.18265","snapshot_observed_at":"2026-07-10T12:15:01.137692Z","title":"InternVL3.5: Advancing Open-Source Multimodal Models in Versatility, Reasoning, and Efficiency","venue":"cs.CV","work_id":"b8f5e260-fff5-444e-bcf5-2c42cfefd83d","year":2025},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":55,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"cited_paper":"/paper/2508.18265","citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:4ddd47cab4de8504c9f2ba88d5e5e46d8f1290fac15213a352035f73fdcc018b","observation_id":"019fdc59-210b-42f5-9ab9-a5c1eb01f229","resolution":{"observed_at":"2026-05-11T19:16:07.513468Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.10479","last_updated":"2025-04-19T03:47:21Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-04-14T17:59:25Z","title":"InternVL3: Exploring Advanced Training and Test-Time Recipes for Open-Source Multimodal Models","version":3},"cited_work":{"arxiv_id":"2504.10479","doi":"10.48550/arxiv.2504.10479","metadata_source":"pith","pith_arxiv_id":"2504.10479","snapshot_observed_at":"2026-07-11T03:17:51.831519Z","title":"InternVL3: Exploring Advanced Training and Test-Time Recipes for Open-Source Multimodal Models","venue":"cs.CV","work_id":"fe8637aa-12bc-4434-8d36-9f57b5eebcbe","year":2025},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":56,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"cited_paper":"/paper/2504.10479","citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:2c94f875ea22f0b7c6c88425c357d5da88bc81bda96bb6d143fca0ae7c5a9acc","observation_id":"0071f784-2caf-4bb8-a85a-b9944c00c5a5","resolution":{"observed_at":"2026-05-11T19:16:07.424343Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-05-20T07:54:09.017512+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-20T07:54:09.017512+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2409.12191","last_updated":"2024-10-03T15:54:49Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-09-18T17:59:32Z","title":"Qwen2-VL: Enhancing Vision-Language Model's Perception of the World at Any Resolution","version":2},"cited_work":{"arxiv_id":"2409.12191","doi":"10.48550/arxiv.2409.12191","metadata_source":"pith","pith_arxiv_id":"2409.12191","snapshot_observed_at":"2026-07-11T01:17:42.597062Z","title":"Qwen2-VL: Enhancing Vision-Language Model's Perception of the World at Any Resolution","venue":"cs.CV","work_id":"8abcfe4f-e0fb-44b7-9123-448fac95f90a","year":2024},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":57,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"cited_paper":"/paper/2409.12191","citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:7b98dc06beb9625142d931d3680058b4f2f1d6e57770df927d36e6c60aea8900","observation_id":"e21d934f-0528-4846-8e6a-91bda0f4ea25","resolution":{"observed_at":"2026-05-11T19:16:07.450668Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-07-11T02:19:33.884263+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T02:19:33.884263+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.13923","last_updated":"2025-02-19T18:00:14Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-02-19T18:00:14Z","title":"Qwen2.5-VL Technical Report","version":1},"cited_work":{"arxiv_id":"2502.13923","doi":"10.48550/arxiv.2502.13923","metadata_source":"pith","pith_arxiv_id":"2502.13923","snapshot_observed_at":"2026-07-11T01:17:45.013705Z","title":"Qwen2.5-VL Technical Report","venue":"cs.CV","work_id":"69dffacb-bfe8-442d-be86-48624c60426f","year":2025},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":58,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"cited_paper":"/paper/2502.13923","citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:0a23a33068f80ddabec93e936536e893a8a0174e96da169f0d7a4d4c6f427451","observation_id":"b5a39afa-532c-46f0-b469-d4724efd14d3","resolution":{"observed_at":"2026-05-11T19:16:07.416124Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-07-12T05:19:13.082554+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T05:19:13.082554+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.21631","last_updated":"2025-11-27T12:16:54Z","snapshot_observed_at":"2026-07-06T22:37:03.716474Z","submitted_at":"2025-11-26T17:59:08Z","title":"Qwen3-VL Technical Report","version":2},"cited_work":{"arxiv_id":"2511.21631","doi":"10.1016/j.neunet.2025.107777","metadata_source":"pith","pith_arxiv_id":"2511.21631","snapshot_observed_at":"2026-07-11T11:50:26.030339Z","title":"Qwen3-VL Technical Report","venue":"cs.CV","work_id":"1fe243aa-e3c0-4da6-b391-4cbcfc88d5c0","year":2025},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":59,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"cited_paper":"/paper/2511.21631","citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:94d2c4a83c134a17510b692279204ad9c1ea0b6dec8a13f7962f3fb6a3f1dc69","observation_id":"5c15d653-c77a-4879-878e-1d6a1ed7f8dd","resolution":{"observed_at":"2026-05-11T19:16:07.508619Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2507.05201","last_updated":"2026-04-06T19:50:20Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-07-07T17:01:44Z","title":"MedGemma Technical Report","version":4},"cited_work":{"arxiv_id":"2507.05201","doi":"10.48550/arxiv.2507.05201","metadata_source":"pith","pith_arxiv_id":"2507.05201","snapshot_observed_at":"2026-07-11T11:50:26.030339Z","title":"MedGemma Technical Report","venue":"cs.AI","work_id":"3d3f25c0-31e8-4859-bb3d-0d719b47a63d","year":2025},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":60,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"cited_paper":"/paper/2507.05201","citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:1f7c47fe9fd5a458400206cd529a9119a33856fc1260267b0bcd83e404b9d780","observation_id":"a0a72f57-7d51-48a1-8bb7-ed1803b01720","resolution":{"observed_at":"2026-05-11T19:16:07.492013Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2506.07044","last_updated":"2025-06-13T04:22:02Z","snapshot_observed_at":"2026-07-06T21:38:34.998414Z","submitted_at":"2025-06-08T08:47:30Z","title":"Lingshu: A Generalist Foundation Model for Unified Multimodal Medical Understanding and Reasoning","version":4},"cited_work":{"arxiv_id":"2506.07044","doi":"10.48550/arxiv.2506.07044","metadata_source":"pith","pith_arxiv_id":"2506.07044","snapshot_observed_at":"2026-07-10T12:15:01.137692Z","title":"Lingshu: A Generalist Foundation Model for Unified Multimodal Medical Understanding and Reasoning","venue":"cs.CL","work_id":"63908fd8-1967-4f5a-ae33-9d555390043d","year":2025},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":61,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"cited_paper":"/paper/2506.07044","citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:41cfbac417d0e89bd0f631a00f4a9250be981f7fbe9a58bff8981e5da4c0a4e2","observation_id":"2ed08202-a445-401a-910a-269f2ec7c7cd","resolution":{"observed_at":"2026-05-15T11:10:59.931236Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.19280","last_updated":"2024-09-30T06:45:16Z","snapshot_observed_at":"2026-07-06T18:38:00.379225Z","submitted_at":"2024-06-27T15:50:41Z","title":"HuatuoGPT-Vision, Towards Injecting Medical Visual Knowledge into Multimodal LLMs at Scale","version":4},"cited_work":{"arxiv_id":"2406.19280","doi":"10.48550/arxiv.2406.19280","metadata_source":"pith","pith_arxiv_id":"2406.19280","snapshot_observed_at":"2026-07-10T06:15:00.866473Z","title":"arXiv preprint arXiv:2406.19280 , year=","venue":"cs.CV","work_id":"dd32b8a1-4ad4-4155-b031-c317b565c6e7","year":2024},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":62,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"cited_paper":"/paper/2406.19280","citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:a8fd3e44c71940441367952ece8c8f7990b36634a452137feb7ad4b6af61c31c","observation_id":"2eb0ed3f-43a7-4ef2-a9cb-ba6d34e78ddb","resolution":{"observed_at":"2026-05-11T19:16:07.533998Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-05-25T22:53:37.302988+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-25T22:53:37.302988+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"e2f0371f-21e0-4de7-ba58-e60e71bf9b29","year":2025},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":63,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:b4b7ea64bb28c43efbdba784bc4b604473bf596d6d09efba5f62f967c8a38d98","observation_id":"69fa3833-5c51-42a5-b987-3410237bacfe","resolution":{"observed_at":"2026-05-26T13:17:50.004660Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Proceedings of the IEEE/CVF Conference on Computer Vision and Pattern Recognition , pages=","venue":null,"work_id":"3789cb2c-65d3-43fc-9f85-184fa7e30b44","year":2025},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":64,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:bc2481447a5ac515ce4083ebe196d4d776337f649759fd353c9ddd52c54d691c","observation_id":"19a12006-3b84-4e7e-b071-11e7cb984348","resolution":{"observed_at":"2026-05-26T13:17:49.965773Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2603.00512","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"arXiv preprint arXiv:2603.00512 , year=","venue":null,"work_id":"8fda1f20-a92e-4b09-b880-4f106cff8973","year":null},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":66,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:ef9a92f9d99c21f73f5784046d1c0d5b6cf171fd6a4f9a829a55f1240059977c","observation_id":"e0155bb6-715e-4429-a952-b2da743e96ed","resolution":{"observed_at":"2026-05-11T19:16:07.429415Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"American Journal of Gastroenterology , volume=","venue":null,"work_id":"9c1652bb-073e-48e9-bd15-fe000a937e5d","year":null},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":67,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:b9f29d6f4f02e4c6ae4d7b226c9997ea9b581cf048a43f1a525a4467a2b4f0b2","observation_id":"f1a9b58c-666c-44f1-b94f-1603ec438d54","resolution":{"observed_at":"2026-05-26T13:17:49.972052Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"New England Journal of Medicine , volume=","venue":null,"work_id":"752e1d2a-0013-444a-9227-5dc53968cf94","year":null},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":68,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:4d03e84fc508fa082c982bb0a45ddf8044650587807ea1a7dded389843d9e960","observation_id":"2e138424-c925-461f-80fb-dd668b94af6f","resolution":{"observed_at":"2026-05-26T13:17:49.977297Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Clinical Gastroenterology and Hepatology , volume=","venue":null,"work_id":"42f0b279-01d3-44ac-9311-3c44b8f7dbfa","year":null},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":69,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:ed6fb8aa7ba2ce0dbd795f4dfbe2754c4c138280e9271ed60cc3489148487cac","observation_id":"66c1444c-2725-4d8e-99c2-302ea202736b","resolution":{"observed_at":"2026-05-26T13:17:50.001531Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2604.21814","last_updated":"2026-04-23T16:07:51Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2026-04-23T16:07:51Z","title":"Divide-then-Diagnose: Weaving Clinician-Inspired Contexts for Ultra-Long Capsule Endoscopy Videos","version":1},"cited_work":{"arxiv_id":"2604.21814","doi":null,"metadata_source":"pith","pith_arxiv_id":"2604.21814","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Divide-then-Diagnose: Weaving Clinician-Inspired Contexts for Ultra-Long Capsule Endoscopy Videos","venue":"cs.CV","work_id":"51307048-7ab1-43b6-b5cb-65a0eb219f51","year":2026},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":70,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"cited_paper":"/paper/2604.21814","citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:840936d1d4b8d1b93546abc7a7e7c4da655c7aaa19e58f38ce690cd7af3b2803","observation_id":"5e5cfe94-d8c6-4be6-98bd-85fd6cad65fd","resolution":{"observed_at":"2026-05-11T19:16:07.475502Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Biomedical Optics Express , volume=","venue":null,"work_id":"15cf71db-a581-4bf5-8d2d-f525f8c582a2","year":null},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":71,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:a07f5ce4d3a8597c7cf6a101d03e5a719a31f6a6d170aab9d3e3e0d8a786a618","observation_id":"b30d5025-9271-4960-8ff4-1b633bc269d0","resolution":{"observed_at":"2026-05-26T13:17:50.057919Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"6d9934b3-8cca-49e6-94c8-b02771385606","year":null},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":72,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:6be5cb1e8b66839357c30bf84c7910d4875ee623af63fccba12bbaa6dea5fdee","observation_id":"aef39177-fa86-475b-8d55-11d277d32df9","resolution":{"observed_at":"2026-05-26T13:17:49.941288Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Journal of the American Society of Echocardiography , volume=","venue":null,"work_id":"77716c18-d2d2-4716-b244-dfe5ac99e474","year":null},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":73,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:1b9679823d1f49ca58b5f2bace3bd97bc049cb229480b508ebb97e323ddcb0f2","observation_id":"2911fadb-fbe2-40a9-8a8e-8f4b862fc0ef","resolution":{"observed_at":"2026-05-26T13:17:49.938600Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Acta Radiologica , volume=","venue":null,"work_id":"794785a9-53df-41b3-b9c1-c74e2e7add15","year":null},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":74,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:ff13e9e9e23cbf72a75d2660e2de7c00ba93ff9a9606595cd52cbff87b8b956a","observation_id":"5fed9393-475c-4c68-bfe2-5ee5341ebc0b","resolution":{"observed_at":"2026-05-26T13:17:49.951537Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"IEEE Journal of Biomedical and Health Informatics , volume=","venue":null,"work_id":"2be417d6-cf08-46a8-adbc-ca6c3cdf51f6","year":null},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":75,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:4da0d386fee5be34b566d7f5092cd1efc0a19923129dcac3490bbb9cde591d52","observation_id":"e9eeea78-7ff2-4b8e-9826-9581febb9da6","resolution":{"observed_at":"2026-05-26T13:17:50.005172Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"npj Digital Medicine , volume=","venue":null,"work_id":"dc4d17a6-5d74-47ac-8f04-546e9a5aeb74","year":null},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":76,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:59d18190debd5f9d8f77ccaecf60ef86b4c3e5f8b0b39a93f31518fe69dfa66b","observation_id":"8b786b6f-8085-42ea-9336-50f1551eb565","resolution":{"observed_at":"2026-05-26T13:17:50.054022Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Proceedings of the IEEE/CVF Conference on Computer Vision and Pattern Recognition , pages=","venue":null,"work_id":"2bba4d0b-b6f2-474d-ac7b-d2e970b8671f","year":2025},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":77,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:727cbf7a7855d9a0cae0c2cc541f5780724d9463221e50f73734dc53c8cd46c5","observation_id":"cc5d5c52-dfbc-4c47-8f58-ed3167bd43fa","resolution":{"observed_at":"2026-05-26T13:17:49.913560Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Proceedings of the IEEE/CVF Conference on Computer Vision and Pattern Recognition , pages=","venue":null,"work_id":"78162a81-0fff-46ee-a2ff-6adf67e3c6fb","year":null},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":78,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:f3076815c99b2c97b3e843de8f542979d93f3e84cbd55025bfd66571ad58e83e","observation_id":"91028cd8-6cd9-4428-936a-47fddf4ba213","resolution":{"observed_at":"2026-05-26T13:17:49.916710Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Advances in Neural Information Processing Systems , pages=","venue":null,"work_id":"af1894ff-23d6-4b40-9d24-6d30b07653c1","year":2024},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":79,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:2e1835f05e9d0cd2056ded4e5ab4905da3c1f8eb2ef9aa2124bb3c67a82b3bda","observation_id":"d4565e1a-02cc-4dd8-aa04-58045645717e","resolution":{"observed_at":"2026-05-26T13:17:50.076680Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Medical Image Computing and Computer Assisted Intervention -- MICCAI 2024 , pages=","venue":null,"work_id":"161d403a-de89-4a88-bf4a-71756e62dee6","year":2024},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":80,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:35ca7c74c882d973fdfce5e457c887cc0153fb6aa34c6d5036121839bda9a472","observation_id":"38cf701d-9669-4617-ad25-14d23916e5f6","resolution":{"observed_at":"2026-05-26T13:17:49.919692Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Computer Vision -- ECCV 2024 , pages=","venue":null,"work_id":"1cf72466-5118-4195-97da-943fa6086c5e","year":2024},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":81,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:edca6150e832f73586223f84a9295297aecc01c87651a9805af4289d17ea1281","observation_id":"bcae6083-f207-4b3b-b375-3781da64a295","resolution":{"observed_at":"2026-05-26T13:17:49.915427Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.14391","last_updated":"2025-05-22T00:02:09Z","snapshot_observed_at":"2026-07-06T21:11:58.671097Z","submitted_at":"2025-04-19T19:32:15Z","title":"How Well Can General Vision-Language Models Learn Medicine By Watching Public Educational Videos?","version":2},"cited_work":{"arxiv_id":"2504.14391","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2504.14391","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"arXiv preprint arXiv:2504.14391 , year=","venue":null,"work_id":"244ea0be-8d09-4c16-996a-b15b63b9bfc9","year":2025},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":82,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"cited_paper":"/paper/2504.14391","citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:472c1bd253fe1165022b024ef66d9e641fa5baaf154a5f109508c60417718a5c","observation_id":"187c4c83-cee0-4c58-adf3-e85d4228e7f3","resolution":{"observed_at":"2026-05-11T19:16:07.526906Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2603.06570","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-04T20:10:07.469860Z","title":"arXiv preprint arXiv:2603.06570 , year=","venue":null,"work_id":"9b0e7001-3f8c-4a25-b531-ff1feaff8e9e","year":2026},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":83,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:490c211418f9fe55f8d5cccef68f849904fca2fcb4a8c6090b673283e19af7fe","observation_id":"0554f377-4198-43b6-9270-2c061069cd4a","resolution":{"observed_at":"2026-05-11T19:16:07.444794Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Proceedings of the AAAI Conference on Artificial Intelligence , volume=","venue":null,"work_id":"b270aee9-a606-44d9-b44e-c4038d516bd8","year":2026},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":84,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:dd1744720dfacb88dfd0918bbae8ef61b7b71d6c9c9618e16fe1494185406be6","observation_id":"5a646395-0c8b-4d61-98ce-239c52630c6a","resolution":{"observed_at":"2026-05-26T13:17:49.925229Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2512.06581","last_updated":"2026-04-08T16:04:11Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-12-06T22:27:59Z","title":"MedGRPO: Multi-Task Reinforcement Learning for Heterogeneous Medical Video Understanding","version":4},"cited_work":{"arxiv_id":"2512.06581","doi":null,"metadata_source":"pith","pith_arxiv_id":"2512.06581","snapshot_observed_at":"2026-06-29T23:24:01.724884Z","title":"MedGRPO: Multi-Task Reinforcement Learning for Heterogeneous Medical Video Understanding","venue":"cs.CV","work_id":"369e9872-9603-4303-9d0e-579ce2204120","year":2025},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":85,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"cited_paper":"/paper/2512.06581","citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:0e8229dc55d48b403b697319d097e8cb558a4937e7d955e3675b5cba5c542986","observation_id":"c9370133-c00c-4d3a-8ec3-40a856bb46bf","resolution":{"observed_at":"2026-05-11T19:16:07.568594Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"PhysioNet , year=","venue":null,"work_id":"fdce3bb6-d810-43d4-9bfb-118217f64f90","year":null},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":86,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:4fb64027db9dbfb9ea43f9285d28809160cafd460574fd240244b81589fbc824","observation_id":"84fb6661-0a53-4bb8-9b00-4d0a6db3d4c6","resolution":{"observed_at":"2026-05-26T13:17:49.947580Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"MiniCPM-V: A","venue":null,"work_id":"98e9780f-bf4c-4322-b5ca-4635ffe1c27a","year":2024},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":87,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:acf1d7a65aaf842da41fbee49c788abfc0af84362a11d5848c6be73aebdc0b25","observation_id":"cb1e5f02-3f42-4bb0-8dcd-e7c51e226f38","resolution":{"observed_at":"2026-05-26T13:17:49.980048Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.02713","last_updated":"2025-08-01T16:40:14Z","snapshot_observed_at":"2026-07-06T19:27:11.995338Z","submitted_at":"2024-10-03T17:36:49Z","title":"LLaVA-Video: Video Instruction Tuning With Synthetic Data","version":3},"cited_work":{"arxiv_id":"2410.02713","doi":null,"metadata_source":"pith","pith_arxiv_id":"2410.02713","snapshot_observed_at":"2026-07-09T21:36:34.348434Z","title":"LLaVA-Video: Video Instruction Tuning With Synthetic Data","venue":"cs.CV","work_id":"e598f516-d992-449a-ab6d-6c788b3a1d7b","year":2024},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":88,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"cited_paper":"/paper/2410.02713","citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:75a513690adf4e00eabab039c4d000b9fc7b8eb18edb0793c77dea0b7850fbc6","observation_id":"a1207483-aae8-4b3e-a819-60c5bf9f5a17","resolution":{"observed_at":"2026-05-11T19:16:07.411617Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"2026 , howpublished=","venue":null,"work_id":"53237011-62d0-4c49-8b3c-e085cc6d8df7","year":2026},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":89,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:7b4fb93918e67b857d7aaa7e5e69a8617733074dadfbf8510786c61ac66a71d2","observation_id":"ea2ba43d-e694-4a2c-bacc-163bef11ee43","resolution":{"observed_at":"2026-05-26T13:17:49.882550Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"2026 , howpublished=","venue":null,"work_id":"3dbc0335-fc96-48cd-8762-9ac324f62d0f","year":2026},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":90,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:83d70b198ad2a909cb46a88507ee980c9718f0ef10c6c1638bb9785740187c5c","observation_id":"91fadc65-a633-4257-b8f8-563cae1d340c","resolution":{"observed_at":"2026-05-26T13:17:50.064538Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"We're Expanding Our","venue":null,"work_id":"8d9d9c99-3df4-47f8-940c-6dfe7d26adbe","year":2025},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":91,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:4842596565f1110ab70812561423dc058f94d9b9919098499d52ba73ee1ea344","observation_id":"d3e82b2f-7208-42c1-a92e-db6769e119c1","resolution":{"observed_at":"2026-05-26T13:17:49.896640Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-06T13:42:38.156435Z","title":null,"venue":null,"work_id":"9aa422ef-371a-4d66-910b-f90d626e3a37","year":2026},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":92,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:385b4bca6b3affd261e7ecd938239330172fa673e44f1c7e008d8f0648d75047","observation_id":"73fa5ec1-1c36-472f-a236-02bf9f1b6674","resolution":{"observed_at":"2026-05-26T13:17:49.876811Z","resolver_source":"raw_fallback","status":"parse_uncertain"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-11T02:37:54.519855Z","title":null,"venue":null,"work_id":"72647878-5eae-4741-b770-5602b7561189","year":2026},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":93,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:54cfee1e0e035f7fddec72415a51e605376dfc055849256083839c7c24de1123","observation_id":"82b5c1a8-d259-4ac2-88de-e2a36a45c683","resolution":{"observed_at":"2026-05-26T13:17:49.879550Z","resolver_source":"raw_fallback","status":"parse_uncertain"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2510.08668","doi":"10.48550/arxiv.2510.08668","metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-10T12:15:01.137692Z","title":"Jiang, Y","venue":null,"work_id":"7a8fdc22-20a2-4741-865e-cb4dfc77234d","year":2025},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":94,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:e91f4a18961ad873ffe2747860c883c8e25c009e410c59a698bbb494b2f7283b","observation_id":"4a176ee4-a216-495c-ba16-976d85638ee2","resolution":{"observed_at":"2026-05-11T19:16:07.461953Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.13106","last_updated":"2025-06-03T03:33:14Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-01-22T18:59:46Z","title":"VideoLLaMA 3: Frontier Multimodal Foundation Models for Image and Video Understanding","version":4},"cited_work":{"arxiv_id":"2501.13106","doi":"10.48550/arxiv.2501.13106","metadata_source":"pith","pith_arxiv_id":"2501.13106","snapshot_observed_at":"2026-07-10T06:15:00.866473Z","title":"VideoLLaMA 3: Frontier Multimodal Foundation Models for Image and Video Understanding","venue":"cs.CV","work_id":"38f52461-37fd-4266-bc46-9dea31be2824","year":2025},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":95,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"cited_paper":"/paper/2501.13106","citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:391b843cf34fb1b070cccf37019d6e7e5b47622e5737be7f727896ac617d8e8f","observation_id":"82a5614e-764b-46d7-bfb3-09c97d2a20a4","resolution":{"observed_at":"2026-05-11T19:16:07.439260Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.16852","last_updated":"2024-07-01T02:59:29Z","snapshot_observed_at":"2026-07-06T18:36:12.298104Z","submitted_at":"2024-06-24T17:58:06Z","title":"Long Context Transfer from Language to Vision","version":2},"cited_work":{"arxiv_id":"2406.16852","doi":"10.48550/arxiv.2406.16852","metadata_source":"pith","pith_arxiv_id":"2406.16852","snapshot_observed_at":"2026-07-10T06:15:00.866473Z","title":"Long Context Transfer from Language to Vision","venue":"cs.CV","work_id":"52f1b946-568f-4819-9d8a-a87296f8852d","year":2024},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":96,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"cited_paper":"/paper/2406.16852","citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:2a5e8335c2915beda0a96d5dd5db5a5a9007d9d61b1434366c7b3fb88eb30ea8","observation_id":"ace738c8-8441-410f-949c-d37a5b0e0ab4","resolution":{"observed_at":"2026-05-12T07:08:36.726593Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.00574","last_updated":"2025-07-13T16:21:16Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-12-31T18:01:23Z","title":"VideoChat-Flash: Hierarchical Compression for Long-Context Video Modeling","version":4},"cited_work":{"arxiv_id":"2501.00574","doi":null,"metadata_source":"pith","pith_arxiv_id":"2501.00574","snapshot_observed_at":"2026-07-04T16:49:57.265026Z","title":"VideoChat-Flash: Hierarchical Compression for Long-Context Video Modeling","venue":"cs.CV","work_id":"52ec7cb6-1ef6-4366-a43f-d4c13fce23ad","year":2024},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":97,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"cited_paper":"/paper/2501.00574","citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:8abf737693afa230b519413b1fa819e8033de3fd2ec9a9cbfee8095f902a4a49","observation_id":"52e12638-e132-44a5-8b65-f8bfc00e0044","resolution":{"observed_at":"2026-05-18T04:02:43.915683Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Advances in Neural Information Processing Systems , volume=","venue":null,"work_id":"ae755293-c2d8-4387-9ff4-a49ab10b498a","year":null},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":98,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:d7e85624f771eb5a62e9de46e2874ff30124d2701eb430fcd3f05ccb18c0aad1","observation_id":"dc519c26-77e9-4426-a602-ef0533594b9a","resolution":{"observed_at":"2026-05-26T13:17:49.874038Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.02438","last_updated":"2025-09-10T04:22:46Z","snapshot_observed_at":"2026-07-06T21:03:35.698923Z","submitted_at":"2025-04-03T09:55:09Z","title":"Scaling Video-Language Models to 10K Frames via Hierarchical Differential Distillation","version":5},"cited_work":{"arxiv_id":"2504.02438","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2504.02438","snapshot_observed_at":"2026-07-02T23:57:29.022097Z","title":"Scaling video-language models to 10k frames via hierarchical differential distillation","venue":null,"work_id":"0a6b0011-908b-4498-b4f3-1ba634df498f","year":2025},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":99,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"cited_paper":"/paper/2504.02438","citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:8dfaa4d4d52d5a7aafd3175c7cd236546afb4887cc95c47413d61fa6405b09d8","observation_id":"9b09393b-1b9e-482f-8bcb-fdd8bcb23fa0","resolution":{"observed_at":"2026-05-11T19:16:07.499176Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Proceedings of the IEEE/CVF International Conference on Computer Vision , pages=","venue":null,"work_id":"283c7846-9d25-4316-821f-7568e4dbd110","year":null},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":100,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:0959bd9b58fa2e736faabc3d2b4efe8ce75f870968f26f413dc87c661d6a810c","observation_id":"684cd3d3-6f79-425a-b97b-7290acecb609","resolution":{"observed_at":"2026-05-26T13:17:49.918449Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"2022 , eprint=","venue":null,"work_id":"e6419e28-c886-48fa-87b3-1e1f67a30715","year":2022},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":101,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:9df42b6af3e45f894de2409415a738dbab03f4469202cde44559a8157386ed97","observation_id":"ccc84196-5b42-48de-b067-6ba9ee596362","resolution":{"observed_at":"2026-05-26T13:17:49.892926Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Scientific Data , volume=","venue":null,"work_id":"5a373abd-34e6-4563-8600-afd9a7c98836","year":2025},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":102,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:a157ce36de91cee9c82d53d4b47701b7fa56b1b8617d808a2524e948b0b29100","observation_id":"25fbe252-25f3-4be1-b5f0-851631c9ade1","resolution":{"observed_at":"2026-05-26T13:17:49.931659Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Medical Image Analysis , pages=","venue":null,"work_id":"c0536f93-142a-43f9-9365-e06847fc1c05","year":2025},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":103,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:4fc91b6c4b922e71b89384e6c91e1c2ce690334e90e8b1e4869e19f10c287509","observation_id":"d10158d3-0f78-441c-bb92-87b5b6673533","resolution":{"observed_at":"2026-05-26T13:17:50.011762Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"International journal of computer assisted radiology and surgery , volume=","venue":null,"work_id":"6d286faa-1c0c-48bc-957e-7db833a9b9ea","year":2024},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":104,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:637add84b66222704e6940814e223796adffb9d9343964d36b7f0536ad86dd7c","observation_id":"da4a1ef2-d343-438b-a423-87c73cb583bd","resolution":{"observed_at":"2026-05-26T13:17:50.047466Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Scientific Data , volume=","venue":null,"work_id":"5aa9d441-f7b3-4f0f-bad6-65554b696e8c","year":2024},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":105,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:1995bbcdf62da76d162a7837e52e7bcb841d8123832da901ec237277e87c2278","observation_id":"74e3f639-bc76-438e-a768-410f60ace084","resolution":{"observed_at":"2026-05-26T13:17:49.983377Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Medical Image Analysis , pages=","venue":null,"work_id":"68dcc263-45d1-4067-a5fa-4484368970aa","year":2025},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":106,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:76477a7b96f6f98f7cba2ff50b6cd7d5e15c83589c312f59d41d877b0c7822de","observation_id":"bf9a7a3c-8f97-4116-ad64-f06ffb2ee507","resolution":{"observed_at":"2026-05-26T13:17:49.849405Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"American Journal of Obstetrics and Gynecology , volume=","venue":null,"work_id":"b4e42c5a-a026-4b1c-82f1-7079cb214926","year":2024},"citing_paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild","version":1},"reference_index":107,"source":"arxiv_source","source_observed_at":"2026-05-08T12:28:35.008604Z"},"links":{"citing_paper":"/paper/2605.06537"},"observation_digest":"sha256:03d3cbd81e48cd3b04fd00d2c1c0d882c84962fe6f69d1b7e696ce8ea16ab457","observation_id":"0585806c-8307-4504-88b9-f0f3ae88af75","resolution":{"observed_at":"2026-05-26T13:17:50.019074Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-01T06:32:01.292127+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2605.06537","last_updated":"2026-05-07T16:37:10Z","latest_version":1,"primary_category":"cs.CV","snapshot_observed_at":"2026-07-06T23:19:00.794334Z","submitted_at":"2026-05-07T16:37:10Z","title":"MedHorizon: Towards Long-context Medical Video Understanding in the Wild"},"reference_resolution":{"displayed":89,"state_counts":{"malformed_identifier":0,"metadata_mismatch":15,"parse_uncertain":2,"unresolved":3,"verified_exact":5,"verified_fuzzy":64},"total_outbound_references":89},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-01T06:32:01.292127+00:00","source":"crossref"},{"observed_at":"2026-08-01T06:31:58.492377+00:00","source":"retraction_watch"}],"thesis":"As of 1 August 2026, this Paper Citation Record lists 89 of 89 outbound references and 3 inbound Pith citation observations for arXiv:2605.06537."}