{"as_of":"2026-08-08T10:37:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:199450e6aac5d5f0be03310bd9e6e07b33b9c7e3b375bb268e95ef83c9ce1011","coverage":[{"denominator":104,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":100,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T14:08:15.295893Z","state":"measured"},{"denominator":100,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":100,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-08T06:32:00.761636+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2505.19959/citation-record","integrity":"/paper/2505.19959/integrity","json":"/paper/2505.19959/citation-record.json","paper":"/paper/2505.19959"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:08:05.030910Z","title":null,"venue":null,"work_id":null,"year":2010},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":1,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:05.030910Z"},"links":{"citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:b11f6e3ada160164aa5655c2c8278be1aa21e93f229365c9249a297d0b92e510","observation_id":"36f009a4-20ce-4197-8907-d9602dbea4bc","resolution":{"observed_at":"2026-08-07T14:08:05.030910Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2404.11018","last_updated":"2024-10-17T17:45:09Z","snapshot_observed_at":"2026-07-06T18:01:18.745347Z","submitted_at":"2024-04-17T02:49:26Z","title":"Many-Shot In-Context Learning","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.11018","snapshot_observed_at":"2026-08-07T14:08:05.136655Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":2,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:05.136655Z"},"links":{"cited_paper":"/paper/2404.11018","citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:dadb316caa6c18bd03f03a50aa197112107cbe46693cae757cbb8fc57b374054","observation_id":"23c732eb-df23-457f-a0a0-5fa177053e3f","resolution":{"observed_at":"2026-08-07T14:08:05.136655Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2303.09752","last_updated":"2023-10-24T00:51:49Z","snapshot_observed_at":"2026-07-06T15:04:31.778399Z","submitted_at":"2023-03-17T03:28:17Z","title":"CoLT5: Faster Long-Range Transformers with Conditional Computation","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2303.09752","snapshot_observed_at":"2026-08-07T14:08:05.236860Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":3,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:05.236860Z"},"links":{"cited_paper":"/paper/2303.09752","citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:3dd27a4b6d6009418243c4e5c224056e63fb2b20d31f97c937bf410203399b17","observation_id":"be7399e1-51bd-41b4-8c07-2c21a5f53ef5","resolution":{"observed_at":"2026-08-07T14:08:05.236860Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2307.11088","last_updated":"2023-10-04T10:04:25Z","snapshot_observed_at":"2026-07-06T15:56:38.019661Z","submitted_at":"2023-07-20T17:59:41Z","title":"L-Eval: Instituting Standardized Evaluation for Long Context Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2307.11088","snapshot_observed_at":"2026-08-07T14:08:05.311060Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":4,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:05.311060Z"},"links":{"cited_paper":"/paper/2307.11088","citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:d2a305dff370908eded5fcb1f7f386f250462ed9bbd9cddd369ee3a986856aa6","observation_id":"c2baec95-7bfa-4e92-aac5-fcc38b2ae39b","resolution":{"observed_at":"2026-08-07T14:08:05.311060Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:08:05.422287Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":5,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:05.422287Z"},"links":{"citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:51b8ebb6fd6ce8935a6434753ccac36b819c752b8462f72c4327238108f3c3cd","observation_id":"57510041-f2d9-4931-8326-eefd7badf959","resolution":{"observed_at":"2026-08-07T14:08:05.422287Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:08:05.547022Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":6,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:05.547022Z"},"links":{"citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:ce00da3ae945d127a999bebfdec4b74f9152b0a3726e3c0f99da207b6f4fccdd","observation_id":"a77b48b4-e293-44d5-a58a-7b39933eb371","resolution":{"observed_at":"2026-08-07T14:08:05.547022Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:08:05.631782Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":7,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:05.631782Z"},"links":{"citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:954f2b3f6f8e8e26942e7b2aa45f16236a233b341ce5844d4dfd0873f83be038","observation_id":"8ad5788d-8e53-46f1-8c28-dd8014a251cf","resolution":{"observed_at":"2026-08-07T14:08:05.631782Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:08:05.758842Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:05.758842Z"},"links":{"citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:29943503a4bad0108b75e1057607ec15b61f9eb12800ae92ad4325eae9f7697e","observation_id":"15d1f274-6d41-41b4-83ee-7d18abbe0fbf","resolution":{"observed_at":"2026-08-07T14:08:05.758842Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:08:05.914253Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":9,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:05.914253Z"},"links":{"citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:00ddc2e93444dc851c598acda1bb6a7c6d39cfdf7a3d60572232b2dabcf43253","observation_id":"50e0fadc-4431-4484-bc04-b0d3bed83d4b","resolution":{"observed_at":"2026-08-07T14:08:05.914253Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2412.15204","last_updated":"2025-01-03T11:44:51Z","snapshot_observed_at":"2026-08-04T21:07:04.238345Z","submitted_at":"2024-12-19T18:59:17Z","title":"LongBench v2: Towards Deeper Understanding and Reasoning on Realistic Long-context Multitasks","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.15204","snapshot_observed_at":"2026-08-07T14:08:06.032342Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":10,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:06.032342Z"},"links":{"cited_paper":"/paper/2412.15204","citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:2eb0c96206a157aa45cb0e92c88d427f358545ab37ac312c033af792f91bce17","observation_id":"f8d27438-057f-41e9-b951-2da0fda1244f","resolution":{"observed_at":"2026-08-07T14:08:06.032342Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2408.07055","last_updated":"2024-08-13T17:46:12Z","snapshot_observed_at":"2026-07-06T19:00:17.992934Z","submitted_at":"2024-08-13T17:46:12Z","title":"LongWriter: Unleashing 10,000+ Word Generation from Long Context LLMs","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2408.07055","snapshot_observed_at":"2026-08-07T14:08:06.170959Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:06.170959Z"},"links":{"cited_paper":"/paper/2408.07055","citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:777ffb2912e6b35e345cea9cb1ae4901e387eee8e2e3e3127377cb5ba043df51","observation_id":"a97493f8-f3ab-442b-aef1-73bcfab04b73","resolution":{"observed_at":"2026-08-07T14:08:06.170959Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2004.05150","last_updated":"2020-12-02T17:52:35Z","snapshot_observed_at":"2026-07-31T17:17:17.205582Z","submitted_at":"2020-04-10T17:54:09Z","title":"Longformer: The Long-Document Transformer","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2004.05150","snapshot_observed_at":"2026-08-07T14:08:06.251314Z","title":null,"venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":12,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:06.251314Z"},"links":{"cited_paper":"/paper/2004.05150","citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:b75b789c7b94e8af18f70f133bf7fb47aa5dbbb527589b05e4b4c05aa5e72185","observation_id":"da1b6433-0708-49d1-ac74-2f0e9a5d2d47","resolution":{"observed_at":"2026-08-07T14:08:06.251314Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.11612","last_updated":"2024-06-17T14:58:29Z","snapshot_observed_at":"2026-07-06T18:32:15.404516Z","submitted_at":"2024-06-17T14:58:29Z","title":"Long Code Arena: a Set of Benchmarks for Long-Context Code Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.11612","snapshot_observed_at":"2026-08-07T14:08:06.406839Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":13,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:06.406839Z"},"links":{"cited_paper":"/paper/2406.11612","citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:fbcfefd6c455cef94bc61b3989945f42d3cf0611836d6e5857c235c2a7dd9194","observation_id":"6c337982-4c99-441e-ad16-0052d5a4d5e5","resolution":{"observed_at":"2026-08-07T14:08:06.406839Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:08:06.483335Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":14,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:06.483335Z"},"links":{"citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:cbb1e86cc6af191ef81d558bfc9cda077bd261756e3c77d484c048809ba07962","observation_id":"d480e881-d9dc-405a-bd8f-dcd5fcb33e43","resolution":{"observed_at":"2026-08-07T14:08:06.483335Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:08:06.579957Z","title":null,"venue":null,"work_id":null,"year":2016},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":15,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:06.579957Z"},"links":{"citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:bfc199e581bbdf3b1f3c6b88d01bac3b0731a2ffbf0fde1d72a72050835af348","observation_id":"ef648103-5316-44ca-ade2-f5e1d8d48d69","resolution":{"observed_at":"2026-08-07T14:08:06.579957Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2306.15595","last_updated":"2023-06-28T04:26:05Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-06-27T16:26:26Z","title":"Extending Context Window of Large Language Models via Positional Interpolation","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2306.15595","snapshot_observed_at":"2026-08-07T14:08:06.693696Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":16,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:06.693696Z"},"links":{"cited_paper":"/paper/2306.15595","citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:1d3f53a2cc937afb0f725a6c9e7435b2be40e9e592b0a18d4c0dd8f1d65cdf8f","observation_id":"a21c9707-03d2-45e9-9ce7-ee8405ceb5d2","resolution":{"observed_at":"2026-08-07T14:08:06.693696Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1904.10509","last_updated":"2019-04-23T19:29:47Z","snapshot_observed_at":"2026-08-06T08:05:35.311510Z","submitted_at":"2019-04-23T19:29:47Z","title":"Generating Long Sequences with Sparse Transformers","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1904.10509","snapshot_observed_at":"2026-08-07T14:08:06.811278Z","title":null,"venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":17,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:06.811278Z"},"links":{"cited_paper":"/paper/1904.10509","citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:8191c01f7f7a504d38813233885cbc8bd243ecef60fa0bfe15461a262497c029","observation_id":"f114dd67-bc14-47f4-9b9b-2efed519dbee","resolution":{"observed_at":"2026-08-07T14:08:06.811278Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:08:06.950136Z","title":null,"venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":18,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:06.950136Z"},"links":{"citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:e2e58f3f0ca84f1a2e6f6e5ecaada5ce4c2f94e5ec0413d84036a05103bb8d62","observation_id":"87bf290e-3bba-4cdd-915a-d99f0a966450","resolution":{"observed_at":"2026-08-07T14:08:06.950136Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:08:07.053346Z","title":null,"venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":19,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:07.053346Z"},"links":{"citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:29415b32a069f79ea953218669884499e56235662aec6e84dc876679a1379080","observation_id":"30eab7a4-7492-4031-bed5-ef7c0eb7fe14","resolution":{"observed_at":"2026-08-07T14:08:07.053346Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2307.02486","last_updated":"2023-07-19T12:25:35Z","snapshot_observed_at":"2026-08-05T18:04:01.030198Z","submitted_at":"2023-07-05T17:59:38Z","title":"LongNet: Scaling Transformers to 1,000,000,000 Tokens","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2307.02486","snapshot_observed_at":"2026-08-07T14:08:07.209849Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":20,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:07.209849Z"},"links":{"cited_paper":"/paper/2307.02486","citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:5f55686116bae205bdaf19e022b70f9a91450f64168a4537e0316ae0ddb56803","observation_id":"969928fe-81f9-49eb-b8bc-0c3df11bde47","resolution":{"observed_at":"2026-08-07T14:08:07.209849Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:08:07.392559Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":21,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:07.392559Z"},"links":{"citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:e17579dc69934be545115b50ad3bded41216c8568be9150547c766dba051e429","observation_id":"48bc94f3-5fca-44e4-a814-661f321d2550","resolution":{"observed_at":"2026-08-07T14:08:07.392559Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:08:07.563125Z","title":null,"venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":22,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:07.563125Z"},"links":{"citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:19f004959b2049e7f3f9179f79ff3128694cc1b50ff82e8335408c85a85da34b","observation_id":"ee3cac01-7388-4313-b766-a267c752b592","resolution":{"observed_at":"2026-08-07T14:08:07.563125Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.21783","last_updated":"2024-11-23T23:27:33Z","snapshot_observed_at":"2026-07-06T18:55:11.576666Z","submitted_at":"2024-07-31T17:54:27Z","title":"The Llama 3 Herd of Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.21783","snapshot_observed_at":"2026-08-07T14:08:07.748603Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":23,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:07.748603Z"},"links":{"cited_paper":"/paper/2407.21783","citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:565c1975eb5431525b25ba55c0099719e747d71cde9b148ee1090036ae6f970d","observation_id":"71ed5744-e0af-4a3c-89f1-84b3d0d16d90","resolution":{"observed_at":"2026-08-07T14:08:07.748603Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:08:07.919917Z","title":null,"venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":24,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:07.919917Z"},"links":{"citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:8e6f25b660c3f0482eaa3bf86b876b69b9c881b85f88c4e4acc938c235a5b697","observation_id":"3fc3a51a-5ee8-42c5-911f-2ed7f3fb8c3a","resolution":{"observed_at":"2026-08-07T14:08:07.919917Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:08:07.977882Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":25,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:07.977882Z"},"links":{"citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:1938c40a69b040fdcaa3a9fedc23383487c5d6d9004d6a2e8fcdf224923f9d11","observation_id":"35fda01c-0384-43e1-9911-c5e4a41c165e","resolution":{"observed_at":"2026-08-07T14:08:07.977882Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:08:08.055716Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":26,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:08.055716Z"},"links":{"citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:5eae666635cb688162b7e3b00c1f372280eac4711cb2396cb7af3e86559f3bee","observation_id":"290cb31b-c8bf-4b7d-8540-cc10c75b5866","resolution":{"observed_at":"2026-08-07T14:08:08.055716Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:08:08.151042Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":27,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:08.151042Z"},"links":{"citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:cbabfb526fa82c5823061b008c081b9d066ed13fc3f606684e8f0287942de8a7","observation_id":"5d2e31bd-15cd-4e1c-92cc-f1a66cb79221","resolution":{"observed_at":"2026-08-07T14:08:08.151042Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:08:08.209939Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":28,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:08.209939Z"},"links":{"citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:fd55096e6a288cf542f106dc7d08e37346cb11e892a39fa958d9880b8d1090b6","observation_id":"228eb581-238d-4c91-8c31-b939e8cc29ae","resolution":{"observed_at":"2026-08-07T14:08:08.209939Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.12793","last_updated":"2024-07-30T03:58:11Z","snapshot_observed_at":"2026-08-07T13:56:34.167869Z","submitted_at":"2024-06-18T16:58:21Z","title":"ChatGLM: A Family of Large Language Models from GLM-130B to GLM-4 All Tools","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.12793","snapshot_observed_at":"2026-08-07T14:08:08.324456Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":29,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:08.324456Z"},"links":{"cited_paper":"/paper/2406.12793","citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:b30e3e4630aa9879333b360490ef32ee9e18a53e72120f5b2ac4f32e6ca4f69e","observation_id":"1371f73f-d71c-40a7-9fde-f9f17b0a8472","resolution":{"observed_at":"2026-08-07T14:08:08.324456Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:08:08.411868Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":30,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:08.411868Z"},"links":{"citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:e17c82efde933b15e501e7364bd45e9518b1018cbf7c2325968075ea05caf4b4","observation_id":"850e93ce-2b77-44cf-aa43-18b9cf0bf1ad","resolution":{"observed_at":"2026-08-07T14:08:08.411868Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:08:08.503186Z","title":null,"venue":null,"work_id":null,"year":2003},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":31,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:08.503186Z"},"links":{"citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:70ff9cb76aa827c6fb9702b3d6895f9c820786586a6169d329225efb373ea1fb","observation_id":"d120db7c-28d0-462b-9768-f6dfac044ba0","resolution":{"observed_at":"2026-08-07T14:08:08.503186Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:08:08.667419Z","title":null,"venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":32,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:08.667419Z"},"links":{"citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:d069de7df9552db4222579a19f407d497fca1e5eb235cb34f640363bd305d286","observation_id":"dd5cd532-6371-4231-ae3f-b43f55b1f4a8","resolution":{"observed_at":"2026-08-07T14:08:08.667419Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:08:08.771896Z","title":null,"venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":33,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:08.771896Z"},"links":{"citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:6925c27d1e55a369d63295454e1c95c7e40d7eb609962934ad8efa89921ebd44","observation_id":"a93c96bf-87cf-4b59-bcac-06be6648e577","resolution":{"observed_at":"2026-08-07T14:08:08.771896Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2404.06654","last_updated":"2024-08-06T21:48:58Z","snapshot_observed_at":"2026-07-06T17:58:00.820879Z","submitted_at":"2024-04-09T23:41:27Z","title":"RULER: What's the Real Context Size of Your Long-Context Language Models?","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.06654","snapshot_observed_at":"2026-08-07T14:08:08.874292Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":34,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:08.874292Z"},"links":{"cited_paper":"/paper/2404.06654","citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:64d3cf78aca79ac61e6deb0cfe7c1f897ba823d40c46360b11131d59bc0715f1","observation_id":"3bbfd306-4f72-42d7-9f66-93a9ef997fe7","resolution":{"observed_at":"2026-08-07T14:08:08.874292Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:08:08.993850Z","title":null,"venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":35,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:08.993850Z"},"links":{"citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:ad2656ae77a2022d573587a7761ce5603ac9d5a945c057894c1ee73947208771","observation_id":"18ee87ac-c286-4071-a107-ecf4c743d4a1","resolution":{"observed_at":"2026-08-07T14:08:08.993850Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2207.07858","last_updated":"2022-07-16T07:08:59Z","snapshot_observed_at":"2026-07-06T13:32:01.087734Z","submitted_at":"2022-07-16T07:08:59Z","title":"The Lottery Ticket Hypothesis for Self-attention in Convolutional Neural Network","version":1},"cited_work":{"arxiv_id":"2207.07858","doi":null,"metadata_source":"pith","pith_arxiv_id":"2207.07858","snapshot_observed_at":"2026-08-07T14:08:17.320763Z","title":"The Lottery Ticket Hypothesis for Self-attention in Convolutional Neural Network","venue":"cs.CV","work_id":"dd9709d7-aaed-4e55-9d70-bcc35db11730","year":2022},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":36,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:09.088736Z"},"links":{"cited_paper":"/paper/2207.07858","citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:05c3c5ea015eaaa8d619fee39a3432e3095f04199dd1c6f4eef3abd164517a74","observation_id":"c0252c68-8494-4cdf-939c-61a954ea9b28","resolution":{"observed_at":"2026-08-07T14:08:17.423537Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:08:09.188607Z","title":null,"venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":37,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:09.188607Z"},"links":{"citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:65c17d01be7e2fa0c3ba93645e9b0f9a3429bc1bdd141258f808e999c4b41be2","observation_id":"914562a3-c7ab-45f8-847b-cab42b0ce325","resolution":{"observed_at":"2026-08-07T14:08:09.188607Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:08:09.291430Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":38,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:09.291430Z"},"links":{"citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:388d030d888a12189e11a9dcade788e226ef2b887c71c9d923465b23620ab224","observation_id":"9ad75f55-6ae2-463c-be1d-9ddd0bd3bd08","resolution":{"observed_at":"2026-08-07T14:08:09.291430Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.02834","last_updated":"2024-06-23T08:45:33Z","snapshot_observed_at":"2026-08-07T23:10:12.222808Z","submitted_at":"2024-02-05T09:44:49Z","title":"Shortened LLaMA: Depth Pruning for Large Language Models with Comparison of Retraining Methods","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.02834","snapshot_observed_at":"2026-08-07T14:08:09.402382Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":39,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:09.402382Z"},"links":{"cited_paper":"/paper/2402.02834","citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:c7eed173ffd66ef988e3627e01bdc549adc2d4f1cf83f310dc8db5690ad6a33b","observation_id":"6556d149-7077-46c5-8a50-0d204d2d8db7","resolution":{"observed_at":"2026-08-07T14:08:09.402382Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.12844","last_updated":"2025-02-20T16:20:11Z","snapshot_observed_at":"2026-08-05T16:38:33.638599Z","submitted_at":"2024-07-04T17:57:38Z","title":"metabench -- A Sparse Benchmark of Reasoning and Knowledge in Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.12844","snapshot_observed_at":"2026-08-07T14:08:09.475795Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":40,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:09.475795Z"},"links":{"cited_paper":"/paper/2407.12844","citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:2db9ff95d5efa57363d4d655f826311eba353b9f09935dca1a88c6ffbc0a10e0","observation_id":"bd35b406-d089-423a-ade4-f0b92e37ae66","resolution":{"observed_at":"2026-08-07T14:08:09.475795Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:08:09.589789Z","title":null,"venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":41,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:09.589789Z"},"links":{"citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:7dac7a947829c53541d575265d36fb16679c5099682a41e7306bd203f5771f18","observation_id":"911bf4f1-4f37-422b-9add-79148f797967","resolution":{"observed_at":"2026-08-07T14:08:09.589789Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:08:09.675627Z","title":null,"venue":null,"work_id":null,"year":2002},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":42,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:09.675627Z"},"links":{"citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:6887b6808e66fd831707d196bb6875028eee67514ad4e91d1362ddeb3bfa1978","observation_id":"11090f03-d740-49d1-81f1-c182a6ebdf46","resolution":{"observed_at":"2026-08-07T14:08:09.675627Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.12941","last_updated":"2025-01-24T19:23:15Z","snapshot_observed_at":"2026-08-05T01:31:35.222989Z","submitted_at":"2024-09-19T17:52:07Z","title":"Fact, Fetch, and Reason: A Unified Evaluation of Retrieval-Augmented Generation","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.12941","snapshot_observed_at":"2026-08-07T14:08:09.772386Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":43,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:09.772386Z"},"links":{"cited_paper":"/paper/2409.12941","citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:c32420a4cbb07fe46a68526d030126fb643a041a5deccdc70654a566f506f111","observation_id":"331578aa-4376-4345-9ac5-68ab5bb9eb12","resolution":{"observed_at":"2026-08-07T14:08:09.772386Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.10149","last_updated":"2024-11-06T14:50:40Z","snapshot_observed_at":"2026-08-06T20:22:33.555366Z","submitted_at":"2024-06-14T16:00:29Z","title":"BABILong: Testing the Limits of LLMs with Long Context Reasoning-in-a-Haystack","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.10149","snapshot_observed_at":"2026-08-07T14:08:09.869247Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":44,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:09.869247Z"},"links":{"cited_paper":"/paper/2406.10149","citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:93999de613c04e703f104448dce8f1eb8adfdf4731027c30767c3b874ce863cc","observation_id":"54cc1b36-120f-4524-a2db-2901e5781a05","resolution":{"observed_at":"2026-08-07T14:08:09.869247Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:08:09.987230Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":45,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:09.987230Z"},"links":{"citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:9dfcc4d46d1af39ac2abe9a980adc341975e7a86e648c304d1f3d8ce03578a0e","observation_id":"d7d61f63-e1c4-4a1c-90cf-e3d62cc8d642","resolution":{"observed_at":"2026-08-07T14:08:09.987230Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:08:10.067995Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":46,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:10.067995Z"},"links":{"citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:e99e623d2fc28a26145787343cffd17a2cb2b6823c01f29536570218e94ede98","observation_id":"1df3c3c2-46c9-4818-88cd-40eaf7374158","resolution":{"observed_at":"2026-08-07T14:08:10.067995Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:08:10.178485Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":47,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:10.178485Z"},"links":{"citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:62463602390e9a7e049aa41569a10e890da44d12bb39f4eaa95c30875462e016","observation_id":"fe12a0cf-eee7-4c20-9f4f-1df80de5642f","resolution":{"observed_at":"2026-08-07T14:08:10.178485Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:08:22.643683Z","title":"Gonzalez, Ion Stoica, Xuezhe Ma, and Hao Zhang","venue":null,"work_id":"2e5449ad-3a58-4379-8f9a-563d177c3c53","year":2023},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":48,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:10.263096Z"},"links":{"citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:18687ab865e1ce2a5f4a37f2fce965a70c72b11442a435b3268150ab4af5c713","observation_id":"b3767e00-56b7-4996-8b68-1c6bc3ed3450","resolution":{"observed_at":"2026-08-07T14:08:22.753146Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2311.04939","last_updated":"2024-09-06T05:06:51Z","snapshot_observed_at":"2026-07-06T16:44:55.494764Z","submitted_at":"2023-11-08T01:45:37Z","title":"LooGLE: Can Long-Context Language Models Understand Long Contexts?","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2311.04939","snapshot_observed_at":"2026-08-07T14:08:10.341641Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":49,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:10.341641Z"},"links":{"cited_paper":"/paper/2311.04939","citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:027605d341aea8e3dadd81956d0310739a672fd1a3f2f362c90cd6a1da57581e","observation_id":"c563a4cc-dee0-4af1-a8ad-1ab6b3cadab3","resolution":{"observed_at":"2026-08-07T14:08:10.341641Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:08:10.457924Z","title":null,"venue":null,"work_id":null,"year":2002},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":50,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:10.457924Z"},"links":{"citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:da5c50d0606dbf8976a8c3eba365d964e3d4054283ab059c7d217bf53ec10ce9","observation_id":"590d1170-c01e-4af8-b63f-1409c92f6ab2","resolution":{"observed_at":"2026-08-07T14:08:10.457924Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:08:22.410282Z","title":null,"venue":null,"work_id":"d71b4b20-4358-4eb3-9595-57b0fbb8715e","year":2020},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":51,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:10.527332Z"},"links":{"citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:bacf0b1eafd9b854d61e572c81e0e2ceedef72143595ec57ab8ffa9eb7ae5a24","observation_id":"713bb8db-7235-4b2d-b246-e24730378138","resolution":{"observed_at":"2026-08-07T14:08:22.489352Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2304.13343","last_updated":"2025-03-18T02:16:56Z","snapshot_observed_at":"2026-08-01T15:59:14.413948Z","submitted_at":"2023-04-26T07:25:31Z","title":"SCM: Enhancing Large Language Model with Self-Controlled Memory Framework","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2304.13343","snapshot_observed_at":"2026-08-07T14:08:10.635750Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":52,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:10.635750Z"},"links":{"cited_paper":"/paper/2304.13343","citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:fa30ca602e10860e61ed41e9e8a6c1ecc16b4fce3d2aee5d2489d531ceb50e4f","observation_id":"c6c3adf7-ece6-4e28-9505-baf7c31fdd2b","resolution":{"observed_at":"2026-08-07T14:08:10.635750Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:08:22.218098Z","title":null,"venue":null,"work_id":"638617fb-eddc-4d42-9bf9-2a97b2b33a71","year":2024},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":53,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:10.773877Z"},"links":{"citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:1577c444d6a49cb9678d2a4d4d8a2796c9b366fa438acc6f5792cb8f057514c9","observation_id":"bd8e2a89-9280-4813-9505-72fee745fe3f","resolution":{"observed_at":"2026-08-07T14:08:22.270695Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.03091","last_updated":"2023-10-04T01:13:49Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-06-05T17:59:41Z","title":"RepoBench: Benchmarking Repository-Level Code Auto-Completion Systems","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2306.03091","snapshot_observed_at":"2026-08-07T14:08:10.859817Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":54,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:10.859817Z"},"links":{"cited_paper":"/paper/2306.03091","citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:dfaa830981dcb7fea1e618572e0a91a9ee3817fc0164fa9184fb1f08fcfdf34e","observation_id":"ffb61a20-d502-4cc3-b553-175d347860e5","resolution":{"observed_at":"2026-08-07T14:08:10.859817Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:08:22.027430Z","title":null,"venue":null,"work_id":"50b6ce3e-3e72-4244-9e2c-a78830d82509","year":2024},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":55,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:10.962003Z"},"links":{"citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:387a8c189e6ecb399fb13149ecd4625e09eea5950e61a63415e85e9687175a78","observation_id":"975fa213-e651-4bc7-b6d8-794521052eeb","resolution":{"observed_at":"2026-08-07T14:08:22.107262Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:08:21.812511Z","title":null,"venue":null,"work_id":"0ae936ef-a6b1-434f-8c83-63535afc8243","year":2019},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":56,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:11.066520Z"},"links":{"citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:98140b6875ccdf88cba965d5421bf0aa00e976cc7b236c80d0f8ab7551b9b1d3","observation_id":"659b3a56-e17f-4311-927b-5ac91875216d","resolution":{"observed_at":"2026-08-07T14:08:21.911882Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:08:11.164535Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":57,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:11.164535Z"},"links":{"citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:463fe95e162325659db66b4ea36ad11d139f496aa68e3775a7a9acda5310056e","observation_id":"d863895b-7ab1-4774-b2ab-5dfff1cce93c","resolution":{"observed_at":"2026-08-07T14:08:11.164535Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:08:11.263738Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":58,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:11.263738Z"},"links":{"citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:b701564a55bd5abc1281a3bd3f72e19d0228d1eb3cb6674abf5b68443fdef436","observation_id":"eb788b76-9f08-4ba5-9f02-2c97a112fcf8","resolution":{"observed_at":"2026-08-07T14:08:11.263738Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:08:21.566672Z","title":null,"venue":null,"work_id":"ec84d619-50bc-45cc-a0f4-97d9487096e7","year":2024},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":59,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:11.366868Z"},"links":{"citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:fc675347ad4178d66f0f26c95fc0bd5c8831ecf0a88e1c95a1031c3a232be585","observation_id":"c1995f35-4992-40ef-822b-5f93b1301f17","resolution":{"observed_at":"2026-08-07T14:08:21.648741Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2303.06349","last_updated":"2023-03-11T08:53:11Z","snapshot_observed_at":"2026-08-07T23:56:00.613985Z","submitted_at":"2023-03-11T08:53:11Z","title":"Resurrecting Recurrent Neural Networks for Long Sequences","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2303.06349","snapshot_observed_at":"2026-08-07T14:08:11.477758Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":60,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:11.477758Z"},"links":{"cited_paper":"/paper/2303.06349","citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:4c74163e0bab3cce49729e3fe969a2a2a746a96c5803293780a168714917f682","observation_id":"caab8c05-60db-45b6-96f1-b1f50c6e58b0","resolution":{"observed_at":"2026-08-07T14:08:11.477758Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.03563","last_updated":"2024-09-05T14:19:45Z","snapshot_observed_at":"2026-07-06T19:11:01.164926Z","submitted_at":"2024-09-05T14:19:45Z","title":"100 instances is all you need: predicting the success of a new LLM on unseen data by testing on a few instances","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.03563","snapshot_observed_at":"2026-08-07T14:08:11.562127Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":61,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:11.562127Z"},"links":{"cited_paper":"/paper/2409.03563","citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:dc7a93eb2e866955de9ccd1291c2409d71b181ce40349e972166ac9c9da88237","observation_id":"9a27f10b-e3da-4e11-8b1f-2892246834e9","resolution":{"observed_at":"2026-08-07T14:08:11.562127Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:08:21.365209Z","title":null,"venue":null,"work_id":"633b031e-be37-4e36-bb06-c9bd96e3d54b","year":2022},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":62,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:11.671826Z"},"links":{"citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:bf5f1b18d4a7ab66c905adaee335365c5237ded00dc067812bb2eec5e90f1355","observation_id":"969dc72c-4365-49a3-b14e-2f720a7bf920","resolution":{"observed_at":"2026-08-07T14:08:21.453040Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:08:21.202648Z","title":null,"venue":null,"work_id":"18e29e33-9de6-47a3-b3b8-9a284e4152f2","year":2024},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":63,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:11.758897Z"},"links":{"citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:c88e0f69a410ec59f2c7bb6aadb0a0d95d7eb5ce4ceaa07a6fe614e289c922d1","observation_id":"b3755203-6e5a-4b64-8275-f03103895d99","resolution":{"observed_at":"2026-08-07T14:08:21.272801Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:08:11.873071Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":64,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:11.873071Z"},"links":{"citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:e723a0cb973b6cacb80361ad83c98a336b97aeb1403d93e96b846a92eb932f36","observation_id":"fa28b1c8-1a01-4999-9b0c-182ebfd681f9","resolution":{"observed_at":"2026-08-07T14:08:11.873071Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.16191","last_updated":"2024-09-24T15:38:11Z","snapshot_observed_at":"2026-08-03T19:24:54.460507Z","submitted_at":"2024-09-24T15:38:11Z","title":"HelloBench: Evaluating Long Text Generation Capabilities of Large Language Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.16191","snapshot_observed_at":"2026-08-07T14:08:11.951710Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":65,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:11.951710Z"},"links":{"cited_paper":"/paper/2409.16191","citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:b945b0652baa26cb09a0437c83dacd11e4475bedb4f3151875174ee8d0b4d331","observation_id":"87ee0065-c944-4130-8978-cd0b559da56b","resolution":{"observed_at":"2026-08-07T14:08:11.951710Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:08:12.059470Z","title":null,"venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":66,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:12.059470Z"},"links":{"citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:abf35dc723861c75082ec2f3a0a5423bdcb6c3bc8569e815627e20dfb40961b6","observation_id":"d256a96e-bfba-418d-93af-3e13acaa1efa","resolution":{"observed_at":"2026-08-07T14:08:12.059470Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2403.05530","last_updated":"2024-12-16T17:39:39Z","snapshot_observed_at":"2026-07-06T17:41:42.995949Z","submitted_at":"2024-03-08T18:54:20Z","title":"Gemini 1.5: Unlocking multimodal understanding across millions of tokens of context","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.05530","snapshot_observed_at":"2026-08-07T14:08:12.159404Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":67,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:12.159404Z"},"links":{"cited_paper":"/paper/2403.05530","citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:e9883ff33411ce0061c83af2d9e9fdbcb3862bdfd1f8b36506487c62ebfde224","observation_id":"fa8ad007-d480-4cb0-ae77-afc02aae811a","resolution":{"observed_at":"2026-08-07T14:08:12.159404Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:08:12.268435Z","title":null,"venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":68,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:12.268435Z"},"links":{"citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:a2e42ac7d970a3812ebb8d21480a6cba0d2985e255fad0ac043339e7f28596c1","observation_id":"1165895d-a84c-4ea6-b074-b71bd1f189f3","resolution":{"observed_at":"2026-08-07T14:08:12.268435Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2301.04272","last_updated":"2023-09-26T04:43:31Z","snapshot_observed_at":"2026-07-06T14:40:12.494082Z","submitted_at":"2023-01-11T02:25:10Z","title":"Data Distillation: A Survey","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2301.04272","snapshot_observed_at":"2026-08-07T14:08:12.368424Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":69,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:12.368424Z"},"links":{"cited_paper":"/paper/2301.04272","citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:abd12f90bc4ff23387fea258098464f7527ad912cf3ba8dd9f34dc580bdef07b","observation_id":"8eac7892-6b1c-44d1-86da-d5d1554946a3","resolution":{"observed_at":"2026-08-07T14:08:12.368424Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.14196","last_updated":"2023-12-17T17:05:09Z","snapshot_observed_at":"2026-08-03T17:17:47.899275Z","submitted_at":"2023-05-23T16:15:31Z","title":"ZeroSCROLLS: A Zero-Shot Benchmark for Long Text Understanding","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.14196","snapshot_observed_at":"2026-08-07T14:08:12.475103Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":70,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:12.475103Z"},"links":{"cited_paper":"/paper/2305.14196","citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:bccd75cd19326e22ca58569483f62be14b6fba0501f3043efa6cbfbab6f833ed","observation_id":"0dc304fa-0f89-40e6-8639-7fcad892a0ca","resolution":{"observed_at":"2026-08-07T14:08:12.475103Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:08:12.587074Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":71,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:12.587074Z"},"links":{"citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:9e1d245977df7f96e92ea52516b178c49fe2b1192ae81f441e80ba3c7e6c3faf","observation_id":"c079bdd4-900f-4443-a405-df9200804bb8","resolution":{"observed_at":"2026-08-07T14:08:12.587074Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2403.11802","last_updated":"2024-12-24T01:41:28Z","snapshot_observed_at":"2026-07-06T17:46:20.935124Z","submitted_at":"2024-03-18T14:01:45Z","title":"Counting-Stars: A Multi-evidence, Position-aware, and Scalable Benchmark for Evaluating Long-Context Large Language Models","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.11802","snapshot_observed_at":"2026-08-07T14:08:12.726053Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":72,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:12.726053Z"},"links":{"cited_paper":"/paper/2403.11802","citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:d478bc37ecd3ca0f8593e76acf405b870b6971822942d6e4b0513b1ae98e923d","observation_id":"9ccd3baa-49a3-4c3f-aabd-9f5bb6097bfc","resolution":{"observed_at":"2026-08-07T14:08:12.726053Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:08:12.802069Z","title":null,"venue":null,"work_id":null,"year":1961},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":73,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:12.802069Z"},"links":{"citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:90ded66ee5df89237e7c7dc0a594bdb4c7398b46bf944edd41455414cdb6733e","observation_id":"ed4a417c-0849-49bf-b923-b2c7fdfd3bb8","resolution":{"observed_at":"2026-08-07T14:08:12.802069Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2212.10554","last_updated":"2022-12-20T18:56:20Z","snapshot_observed_at":"2026-08-05T12:58:36.921230Z","submitted_at":"2022-12-20T18:56:20Z","title":"A Length-Extrapolatable Transformer","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2212.10554","snapshot_observed_at":"2026-08-07T14:08:12.883259Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":74,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:12.883259Z"},"links":{"cited_paper":"/paper/2212.10554","citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:9da7bf6628d6e430088ae6e837280e8202263e62546e9d32c7f5cbf0e8736425","observation_id":"613f2a18-8206-4807-8b1f-bcc3779a936f","resolution":{"observed_at":"2026-08-07T14:08:12.883259Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:08:12.956018Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":75,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:12.956018Z"},"links":{"citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:a21f8e27d2e2166d95648c83e1a9300c3fae407a93f2a1c941eff3c70d24e1ba","observation_id":"f4d2d20d-9037-4145-acfa-88d1d4f1784c","resolution":{"observed_at":"2026-08-07T14:08:12.956018Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:08:20.756573Z","title":null,"venue":null,"work_id":"83405447-3a79-48af-91cb-ebadd558f3fe","year":2021},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":76,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:13.052388Z"},"links":{"citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:c214e4cdcb42e3b05884a573cfa53aebcc64e4b7f6dec16b839cabb4cb893ae5","observation_id":"76d8bcea-53d0-4f88-befa-c629a28f5fe6","resolution":{"observed_at":"2026-08-07T14:08:20.887675Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:08:13.148130Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":77,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:13.148130Z"},"links":{"citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:bf246f19517536df7f5a3dfb5f3859d8b729b322a55733a78301c84a85a50d3d","observation_id":"b6092f5c-4adb-4a28-b4fa-9857e8a9b1cc","resolution":{"observed_at":"2026-08-07T14:08:13.148130Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:08:13.223968Z","title":null,"venue":null,"work_id":null,"year":2008},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":78,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:13.223968Z"},"links":{"citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:830cbaa04aacb033b85ba4f4d2c2c7af7cd9294f00cc55c1d8d9344cca48e4f5","observation_id":"ea2f2ad2-1e87-43fc-8744-ef7668cff118","resolution":{"observed_at":"2026-08-07T14:08:13.223968Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.12640","last_updated":"2024-09-20T00:47:33Z","snapshot_observed_at":"2026-07-06T19:18:04.618039Z","submitted_at":"2024-09-19T10:38:01Z","title":"Michelangelo: Long Context Evaluations Beyond Haystacks via Latent Structure Queries","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.12640","snapshot_observed_at":"2026-08-07T14:08:13.293848Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":79,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:13.293848Z"},"links":{"cited_paper":"/paper/2409.12640","citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:c9468c9cbcdc0acfa520c86c3758560fa6bf4cac34045013a536cffdacab58b5","observation_id":"053e82e7-0cc7-4800-afee-e20eb51302f8","resolution":{"observed_at":"2026-08-07T14:08:13.293848Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:08:20.514294Z","title":null,"venue":null,"work_id":"06c7cc4e-2c77-4f21-9b36-74803ffa144c","year":2022},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":80,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:13.395062Z"},"links":{"citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:308ee4ed6b554148fae798e037ce9acb983b717f15daf364ff8d3ef367b10a57","observation_id":"cdda6206-0c00-4e11-8cc3-e73651ab7dcf","resolution":{"observed_at":"2026-08-07T14:08:20.572057Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:08:20.387967Z","title":null,"venue":null,"work_id":"9264027d-e866-4e8f-8472-b21c40e4bd80","year":2024},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":81,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:13.528436Z"},"links":{"citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:c175971f7e101330efe1e8aa9f6736f6dd2230affbdc0f0a3b316013bf0a184b","observation_id":"13713aa4-6407-45e6-98bc-a76cd77816c0","resolution":{"observed_at":"2026-08-07T14:08:20.446913Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2006.04768","last_updated":"2020-06-14T08:15:54Z","snapshot_observed_at":"2026-07-06T09:27:03.809621Z","submitted_at":"2020-06-08T17:37:52Z","title":"Linformer: Self-Attention with Linear Complexity","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2006.04768","snapshot_observed_at":"2026-08-07T14:08:13.613513Z","title":null,"venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":82,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:13.613513Z"},"links":{"cited_paper":"/paper/2006.04768","citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:00a6fab8230e743e08d85435bb497b2726a3263c441ae27ca82404df27052399","observation_id":"978c295b-9eda-4a7b-becb-fcc701d835e6","resolution":{"observed_at":"2026-08-07T14:08:13.613513Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.02076","last_updated":"2025-01-23T00:52:08Z","snapshot_observed_at":"2026-08-07T03:51:37.675554Z","submitted_at":"2024-09-03T17:25:54Z","title":"LongGenBench: Benchmarking Long-Form Generation in Long Context LLMs","version":7},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.02076","snapshot_observed_at":"2026-08-07T14:08:13.719847Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":83,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:13.719847Z"},"links":{"cited_paper":"/paper/2409.02076","citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:84762496464226a2d63e99a80370796d5efd8134fc93274b800462870a95bc15","observation_id":"f084daa4-80b3-4738-be30-c30b2fe271c1","resolution":{"observed_at":"2026-08-07T14:08:13.719847Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:08:13.808547Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":84,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:13.808547Z"},"links":{"citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:dc13183cb01c8972661f857cd0607070cfcc9341005def4cc5251d2caa34c9cb","observation_id":"3d31721f-a409-4053-b2e7-3e9d34cf48d4","resolution":{"observed_at":"2026-08-07T14:08:13.808547Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:08:20.221456Z","title":null,"venue":null,"work_id":"017002fe-fda9-4ba2-8291-8830015430d6","year":2024},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":85,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:13.883028Z"},"links":{"citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:d593965a22e96b7e052e77f6a2d08f7103904d7b1b0fdff1957972d8ef698c74","observation_id":"b26290f0-e434-4f2d-acf4-8ea8706ac880","resolution":{"observed_at":"2026-08-07T14:08:20.274164Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:08:20.085202Z","title":null,"venue":null,"work_id":"0d58f78d-b953-4b2b-8bb0-f5fdda965bee","year":2024},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":86,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:13.960049Z"},"links":{"citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:f1a32e5dd9e6c2da908310293c638603a065ba7fe02f8643b89e630a63de27d5","observation_id":"020da7cb-4fad-45c7-937a-e813364e8d57","resolution":{"observed_at":"2026-08-07T14:08:20.146089Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.11187","last_updated":"2024-10-15T01:58:58Z","snapshot_observed_at":"2026-08-02T23:45:33.761426Z","submitted_at":"2024-02-17T04:16:30Z","title":"LaCo: Large Language Model Pruning via Layer Collapse","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.11187","snapshot_observed_at":"2026-08-07T14:08:14.032835Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":87,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:14.032835Z"},"links":{"cited_paper":"/paper/2402.11187","citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:bcb048bca57879571ad4d6078a83fc3e60aa4626bcbe41b1aea3c29fccd15521","observation_id":"4a039955-dfb5-49d8-adb2-f6367fb18e70","resolution":{"observed_at":"2026-08-07T14:08:14.032835Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:08:14.143534Z","title":null,"venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":88,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:14.143534Z"},"links":{"citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:12b87839d7faa04f5195159ac35604cbeb5ed072c16bf2ca3e02009ec09ad5f4","observation_id":"45e51e4a-1647-4678-8a1b-94cbc619a263","resolution":{"observed_at":"2026-08-07T14:08:14.143534Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.02694","last_updated":"2025-03-06T18:41:54Z","snapshot_observed_at":"2026-08-06T10:41:12.245173Z","submitted_at":"2024-10-03T17:20:11Z","title":"HELMET: How to Evaluate Long-Context Language Models Effectively and Thoroughly","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.02694","snapshot_observed_at":"2026-08-07T14:08:14.241741Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":89,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:14.241741Z"},"links":{"cited_paper":"/paper/2410.02694","citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:60e8498a147f777a1da533e9d764be02b829e2e0919e1957cd21572df0704a4c","observation_id":"27e9d54c-f2ab-49f2-9bfe-d1246d76f661","resolution":{"observed_at":"2026-08-07T14:08:14.241741Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:08:19.880081Z","title":null,"venue":null,"work_id":"83760980-7c63-43f3-829b-519021f3856e","year":2023},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":90,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:14.344622Z"},"links":{"citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:d9fcf8c5b01a879aabcbac7c2da5245a72de4829e6ee7204c209a74a4122f1f1","observation_id":"dc2e0918-8420-4d97-a763-0657d119a533","resolution":{"observed_at":"2026-08-07T14:08:19.931372Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:08:14.438299Z","title":null,"venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":91,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:14.438299Z"},"links":{"citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:de177c9d00c7a0af665156581393f5d42d4d7bbb2ccb98b89036e7c6414aa48e","observation_id":"0ae1b696-6ad6-4084-bd2c-ade7194d028f","resolution":{"observed_at":"2026-08-07T14:08:14.438299Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:08:14.538802Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":92,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:14.538802Z"},"links":{"citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:c021d10b0f1547c2e62f6941913f3a89c84a79acc7560594541407ff74325a0b","observation_id":"99d0467f-c0f3-4048-8fab-2207ed4bae3a","resolution":{"observed_at":"2026-08-07T14:08:14.538802Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.02897","last_updated":"2024-09-10T07:43:19Z","snapshot_observed_at":"2026-07-06T19:10:29.880696Z","submitted_at":"2024-09-04T17:41:19Z","title":"LongCite: Enabling LLMs to Generate Fine-grained Citations in Long-context QA","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.02897","snapshot_observed_at":"2026-08-07T14:08:14.621446Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":93,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:14.621446Z"},"links":{"cited_paper":"/paper/2409.02897","citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:e1275ee0f297a7d657ccafe22dafa03280e6d67f49eb37682359fb07598154d5","observation_id":"21837170-dad2-4678-a1ed-726b05d60ae2","resolution":{"observed_at":"2026-08-07T14:08:14.621446Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:08:14.728257Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":94,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:14.728257Z"},"links":{"citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:715f2c8cc76ccedebf58b0f78ad5fe3ec0a99be9f1ba36118425ae19de572b7f","observation_id":"0f0d6dee-93dd-40dd-be67-08aa5684ffc0","resolution":{"observed_at":"2026-08-07T14:08:14.728257Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:08:14.797283Z","title":null,"venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":95,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:14.797283Z"},"links":{"citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:f296383121a41ed6669d963e8479902715ca6cb9b2370e4958199d0ba896f7fd","observation_id":"12101dee-2e4c-4a83-a7a1-c593944293ed","resolution":{"observed_at":"2026-08-07T14:08:14.797283Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:08:19.582651Z","title":null,"venue":null,"work_id":"8452e7a4-c4d0-431e-ab7b-d191db1a12b6","year":2024},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":96,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:14.886811Z"},"links":{"citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:eec054cce4a6af89855d657886d8c8fc1ca1bb9ef8eef89188fe1daf5bb0a594","observation_id":"2a3f47dc-e8a8-4940-8793-ba35e51c9adb","resolution":{"observed_at":"2026-08-07T14:08:19.651989Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:08:19.365373Z","title":null,"venue":null,"work_id":"2e08db2f-88f3-434a-b087-7c13f431e971","year":2024},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":97,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:15.012338Z"},"links":{"citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:f3cde06f1aa0485c934863c2735882e9e838668d9c8e093efe11814cc2017d4e","observation_id":"4b0cde7d-d13a-4a9a-b3f5-ad962fd395bc","resolution":{"observed_at":"2026-08-07T14:08:19.472961Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:08:19.182636Z","title":null,"venue":null,"work_id":"038b5458-4447-4be9-8b2e-45dc225d6869","year":2023},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":98,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:15.115835Z"},"links":{"citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:9f429e60652caf3f739f1ea31853d9fe43103563d37ebcad003e48fffe928647","observation_id":"57925309-3ef2-48f2-9190-7e7f9d81cf31","resolution":{"observed_at":"2026-08-07T14:08:19.259176Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:08:18.970008Z","title":null,"venue":null,"work_id":"1a736613-1fc9-48d6-8090-0bed10e9dfb3","year":2022},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":99,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:15.195497Z"},"links":{"citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:aedb8445a8bc6823bcf7ea4c074b9788e91c7f63977c706dd174db3b41ceb5b0","observation_id":"0045c5c7-f944-4994-8438-0ece18e4e05c","resolution":{"observed_at":"2026-08-07T14:08:19.054089Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T14:08:18.647577Z","title":null,"venue":null,"work_id":"c006b93c-212d-4bd0-946b-afd04fc82b5c","year":2023},"citing_paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models","version":2},"reference_index":100,"source":"arxiv_source","source_observed_at":"2026-08-07T14:08:15.295893Z"},"links":{"citing_paper":"/paper/2505.19959"},"observation_digest":"sha256:71837125972c1e119d67365f2dec0bc36db1d42651bf8187a8f0abdf87067650","observation_id":"d7c3ac62-bf70-45f9-aae4-548b22167c97","resolution":{"observed_at":"2026-08-07T14:08:18.760719Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2505.19959","last_updated":"2025-07-30T16:46:12Z","latest_version":2,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-07T23:10:33.455433Z","submitted_at":"2025-05-26T13:21:18Z","title":"MiniLongBench: The Low-cost Long Context Understanding Benchmark for Large Language Models"},"reference_resolution":{"displayed":100,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":98,"verified_exact":1,"verified_fuzzy":1},"total_outbound_references":104},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"thesis":"As of 8 August 2026, this Paper Citation Record lists 100 of 104 outbound references and 0 inbound Pith citation observations for arXiv:2505.19959."}