{"as_of":"2026-08-12T08:18:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:906b25e6629e52e4992de4bda36482dc0b1e035f4372653b792f94e2ed4f7d21","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":18,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":18,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-12T06:34:41.77262+00:00","state":"measured"},{"denominator":18,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":18,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T13:44:50.687143Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":1,"source":"arxiv_reference","source_observed_at":"2026-08-05T02:28:24.338817Z","state":"measured"}],"external_citation_measurements":[{"count":2,"observed_at":"2026-08-05T02:28:24.338817Z","source":"arxiv_reference"}],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2401.04658","last_updated":"2024-01-15T14:57:29Z","snapshot_observed_at":"2026-08-10T13:02:55.737137Z","submitted_at":"2024-01-09T16:27:28Z","title":"Lightning Attention-2: A Free Lunch for Handling Unlimited Sequence Lengths in Large Language Models","version":2},"cited_work":{"arxiv_id":"2401.04658","doi":"10.48550/arxiv.2401.04658","metadata_source":"arxiv_reference","pith_arxiv_id":"2401.04658","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Lightning attention-2: A free lunch for handling unlimited sequence lengths in large language models.arXiv preprint arXiv:2401.04658","venue":"arXiv (Cornell University)","work_id":"02c34d0b-5e0a-48e5-9c0d-9d5474fb4c00","year":2024},"citing_paper":{"arxiv_id":"2410.13846","last_updated":"2026-05-18T05:12:14Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-10-17T17:58:14Z","title":"LightTransfer: Your Long-Context LLM is Secretly a Hybrid Model with Effortless Adaptation","version":3},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-05-23T18:31:35.391674Z"},"links":{"cited_paper":"/paper/2401.04658","citing_paper":"/paper/2410.13846"},"observation_digest":"sha256:bb6247e3a86e4a981310302327d844066a91a6e43f99256fb27ff3cdcc20c757","observation_id":"c2f4e469-cf29-4ec6-81df-2559fb36797b","resolution":{"observed_at":"2026-05-23T18:33:19.483122Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2401.04658","last_updated":"2024-01-15T14:57:29Z","snapshot_observed_at":"2026-08-10T13:02:55.737137Z","submitted_at":"2024-01-09T16:27:28Z","title":"Lightning Attention-2: A Free Lunch for Handling Unlimited Sequence Lengths in Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.04658","snapshot_observed_at":"2026-08-07T13:44:50.687143Z","title":"Lightning attention-2: A free lunch for handling unlim- ited sequence lengths in large language models","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2505.21136","last_updated":"2025-06-06T07:47:22Z","snapshot_observed_at":"2026-08-10T06:09:22.229605Z","submitted_at":"2025-05-27T12:50:36Z","title":"SageAttention2++: A More Efficient Implementation of SageAttention2","version":3},"reference_index":2016,"source":"pdf_text","source_observed_at":"2026-08-07T13:44:50.687143Z"},"links":{"cited_paper":"/paper/2401.04658","citing_paper":"/paper/2505.21136"},"observation_digest":"sha256:58623fce240f1924b13804340bff5c1616547d7c10480b1f199cb69ee5cb1739","observation_id":"34b89f3c-8c5e-482f-a8a0-c935734b2667","resolution":{"observed_at":"2026-08-07T13:44:50.687143Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2401.04658","last_updated":"2024-01-15T14:57:29Z","snapshot_observed_at":"2026-08-10T13:02:55.737137Z","submitted_at":"2024-01-09T16:27:28Z","title":"Lightning Attention-2: A Free Lunch for Handling Unlimited Sequence Lengths in Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.04658","snapshot_observed_at":"2026-08-06T20:29:53.778431Z","title":"Lightning attention-2: A free lunch for handling unlimited sequence lengths in large language models","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.02591","last_updated":"2025-07-23T07:25:27Z","snapshot_observed_at":"2026-08-09T23:16:14.999128Z","submitted_at":"2025-07-03T12:55:16Z","title":"AuroraLong: Bringing RNNs Back to Efficient Open-Ended Video Understanding","version":3},"reference_index":69,"source":"pdf_text","source_observed_at":"2026-08-06T20:29:53.778431Z"},"links":{"cited_paper":"/paper/2401.04658","citing_paper":"/paper/2507.02591"},"observation_digest":"sha256:3258f13daf2e5a93e1bda01ce790e00a7e23c7655aef2988a742c589c0a7bc2d","observation_id":"210f87e3-5250-4913-835f-ce77de68d1b6","resolution":{"observed_at":"2026-08-06T20:29:53.778431Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2401.04658","last_updated":"2024-01-15T14:57:29Z","snapshot_observed_at":"2026-08-10T13:02:55.737137Z","submitted_at":"2024-01-09T16:27:28Z","title":"Lightning Attention-2: A Free Lunch for Handling Unlimited Sequence Lengths in Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.04658","snapshot_observed_at":"2026-08-06T14:59:20.274728Z","title":"arXiv preprint arXiv:2401.04658 (2024)","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.17245","last_updated":"2025-07-23T06:29:38Z","snapshot_observed_at":"2026-08-06T19:10:03.852056Z","submitted_at":"2025-07-23T06:29:38Z","title":"DistrAttention: An Efficient and Flexible Self-Attention Mechanism on Modern GPUs","version":1},"reference_index":2024,"source":"pdf_text","source_observed_at":"2026-08-06T14:59:20.274728Z"},"links":{"cited_paper":"/paper/2401.04658","citing_paper":"/paper/2507.17245"},"observation_digest":"sha256:e1d595e7e41e159bbbdc17f914bcd0f9e462b8d82a4f44800b1c567bc6ac3bb9","observation_id":"4417498f-24f2-49bf-9455-b66670c7696b","resolution":{"observed_at":"2026-08-06T14:59:20.274728Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2401.04658","last_updated":"2024-01-15T14:57:29Z","snapshot_observed_at":"2026-08-10T13:02:55.737137Z","submitted_at":"2024-01-09T16:27:28Z","title":"Lightning Attention-2: A Free Lunch for Handling Unlimited Sequence Lengths in Large Language Models","version":2},"cited_work":{"arxiv_id":"2401.04658","doi":"10.48550/arxiv.2401.04658","metadata_source":"arxiv_reference","pith_arxiv_id":"2401.04658","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Lightning attention-2: A free lunch for handling unlimited sequence lengths in large language models.arXiv preprint arXiv:2401.04658","venue":"arXiv (Cornell University)","work_id":"02c34d0b-5e0a-48e5-9c0d-9d5474fb4c00","year":2024},"citing_paper":{"arxiv_id":"2510.04800","last_updated":"2026-04-21T13:16:20Z","snapshot_observed_at":"2026-08-02T20:03:37.497118Z","submitted_at":"2025-10-06T13:30:07Z","title":"Hybrid Architectures for Language Models: Systematic Analysis and Design Insights","version":3},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-05-18T10:18:04.431436Z"},"links":{"cited_paper":"/paper/2401.04658","citing_paper":"/paper/2510.04800"},"observation_digest":"sha256:a927ba640fd33641c5453c4c6dd4c6a07547deba5f32d4ce48a4edbdb2175fcd","observation_id":"1ab9bc81-4ff8-4e94-9e55-4dc7078412fd","resolution":{"observed_at":"2026-05-18T10:21:15.061486Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2401.04658","last_updated":"2024-01-15T14:57:29Z","snapshot_observed_at":"2026-08-10T13:02:55.737137Z","submitted_at":"2024-01-09T16:27:28Z","title":"Lightning Attention-2: A Free Lunch for Handling Unlimited Sequence Lengths in Large Language Models","version":2},"cited_work":{"arxiv_id":"2401.04658","doi":"10.48550/arxiv.2401.04658","metadata_source":"arxiv_reference","pith_arxiv_id":"2401.04658","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Lightning attention-2: A free lunch for handling unlimited sequence lengths in large language models.arXiv preprint arXiv:2401.04658","venue":"arXiv (Cornell University)","work_id":"02c34d0b-5e0a-48e5-9c0d-9d5474fb4c00","year":2024},"citing_paper":{"arxiv_id":"2510.27258","last_updated":"2026-05-14T15:35:59Z","snapshot_observed_at":"2026-08-08T14:18:17.260680Z","submitted_at":"2025-10-31T07:54:37Z","title":"Higher-order Linear Attention","version":3},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-05-18T03:05:35.823369Z"},"links":{"cited_paper":"/paper/2401.04658","citing_paper":"/paper/2510.27258"},"observation_digest":"sha256:f71ac11a4d37ff284162a81202820ed859360a5e7944caa76c596a1190af2d30","observation_id":"dd382d08-9c0c-4a57-b0f0-d3738fb1b59a","resolution":{"observed_at":"2026-05-18T03:05:47.382651Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2401.04658","last_updated":"2024-01-15T14:57:29Z","snapshot_observed_at":"2026-08-10T13:02:55.737137Z","submitted_at":"2024-01-09T16:27:28Z","title":"Lightning Attention-2: A Free Lunch for Handling Unlimited Sequence Lengths in Large Language Models","version":2},"cited_work":{"arxiv_id":"2401.04658","doi":"10.48550/arxiv.2401.04658","metadata_source":"arxiv_reference","pith_arxiv_id":"2401.04658","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Lightning attention-2: A free lunch for handling unlimited sequence lengths in large language models.arXiv preprint arXiv:2401.04658","venue":"arXiv (Cornell University)","work_id":"02c34d0b-5e0a-48e5-9c0d-9d5474fb4c00","year":2024},"citing_paper":{"arxiv_id":"2512.13564","last_updated":"2026-01-13T09:33:57Z","snapshot_observed_at":"2026-08-06T08:27:07.254588Z","submitted_at":"2025-12-15T17:22:34Z","title":"Memory in the Age of AI Agents","version":2},"reference_index":7,"source":"arxiv_source","source_observed_at":"2026-05-11T18:18:19.911342Z"},"links":{"cited_paper":"/paper/2401.04658","citing_paper":"/paper/2512.13564"},"observation_digest":"sha256:d5b81464f4a87d5fcc45309fbc4612f7c9a5e35bee5f9f40aca1baab5267668a","observation_id":"bbf7b08b-9b95-402f-94d8-f88925d89b30","resolution":{"observed_at":"2026-05-11T18:18:20.163987Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2401.04658","last_updated":"2024-01-15T14:57:29Z","snapshot_observed_at":"2026-08-10T13:02:55.737137Z","submitted_at":"2024-01-09T16:27:28Z","title":"Lightning Attention-2: A Free Lunch for Handling Unlimited Sequence Lengths in Large Language Models","version":2},"cited_work":{"arxiv_id":"2401.04658","doi":"10.48550/arxiv.2401.04658","metadata_source":"arxiv_reference","pith_arxiv_id":"2401.04658","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Lightning attention-2: A free lunch for handling unlimited sequence lengths in large language models.arXiv preprint arXiv:2401.04658","venue":"arXiv (Cornell University)","work_id":"02c34d0b-5e0a-48e5-9c0d-9d5474fb4c00","year":2024},"citing_paper":{"arxiv_id":"2604.10064","last_updated":"2026-04-11T07:06:33Z","snapshot_observed_at":"2026-08-02T07:52:32.788608Z","submitted_at":"2026-04-11T07:06:33Z","title":"On The Application of Linear Attention in Multimodal Transformers","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-05-10T17:12:49.695614Z"},"links":{"cited_paper":"/paper/2401.04658","citing_paper":"/paper/2604.10064"},"observation_digest":"sha256:eb92a7cbe4e1812f6ebf55393707b6c583a48803e1584ff45d91095c4bd4bc8f","observation_id":"3065605d-6d66-4b14-b76f-0514b71f8d99","resolution":{"observed_at":"2026-05-11T07:20:59.526297Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2401.04658","last_updated":"2024-01-15T14:57:29Z","snapshot_observed_at":"2026-08-10T13:02:55.737137Z","submitted_at":"2024-01-09T16:27:28Z","title":"Lightning Attention-2: A Free Lunch for Handling Unlimited Sequence Lengths in Large Language Models","version":2},"cited_work":{"arxiv_id":"2401.04658","doi":"10.48550/arxiv.2401.04658","metadata_source":"arxiv_reference","pith_arxiv_id":"2401.04658","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Lightning attention-2: A free lunch for handling unlimited sequence lengths in large language models.arXiv preprint arXiv:2401.04658","venue":"arXiv (Cornell University)","work_id":"02c34d0b-5e0a-48e5-9c0d-9d5474fb4c00","year":2024},"citing_paper":{"arxiv_id":"2604.24715","last_updated":"2026-04-27T17:23:37Z","snapshot_observed_at":"2026-08-11T11:54:11.489569Z","submitted_at":"2026-04-27T17:23:37Z","title":"Long-Context Aware Upcycling: A New Frontier for Hybrid LLM Scaling","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-05-08T03:39:37.485602Z"},"links":{"cited_paper":"/paper/2401.04658","citing_paper":"/paper/2604.24715"},"observation_digest":"sha256:80feeb5867da34eee9012c0714f1c696c76c8c3e5bc299bbc98340ca0f172af3","observation_id":"18579422-e733-4de1-88c4-2a17a8c7df5b","resolution":{"observed_at":"2026-05-11T22:01:10.671459Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2401.04658","last_updated":"2024-01-15T14:57:29Z","snapshot_observed_at":"2026-08-10T13:02:55.737137Z","submitted_at":"2024-01-09T16:27:28Z","title":"Lightning Attention-2: A Free Lunch for Handling Unlimited Sequence Lengths in Large Language Models","version":2},"cited_work":{"arxiv_id":"2401.04658","doi":"10.48550/arxiv.2401.04658","metadata_source":"arxiv_reference","pith_arxiv_id":"2401.04658","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Lightning attention-2: A free lunch for handling unlimited sequence lengths in large language models.arXiv preprint arXiv:2401.04658","venue":"arXiv (Cornell University)","work_id":"02c34d0b-5e0a-48e5-9c0d-9d5474fb4c00","year":2024},"citing_paper":{"arxiv_id":"2605.06946","last_updated":"2026-05-07T21:05:28Z","snapshot_observed_at":"2026-08-04T02:08:48.810630Z","submitted_at":"2026-05-07T21:05:28Z","title":"Adaptive Memory Decay for Log-Linear Attention","version":1},"reference_index":14,"source":"arxiv_source","source_observed_at":"2026-05-11T01:02:56.848785Z"},"links":{"cited_paper":"/paper/2401.04658","citing_paper":"/paper/2605.06946"},"observation_digest":"sha256:33b7e0ec3365b29fddcb2b1ab3f83a29db8506b0b5be26d1fda44f023f27aac9","observation_id":"682016e4-3e7f-45a6-8a2f-8379c78b9b59","resolution":{"observed_at":"2026-05-11T04:50:56.380655Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2401.04658","last_updated":"2024-01-15T14:57:29Z","snapshot_observed_at":"2026-08-10T13:02:55.737137Z","submitted_at":"2024-01-09T16:27:28Z","title":"Lightning Attention-2: A Free Lunch for Handling Unlimited Sequence Lengths in Large Language Models","version":2},"cited_work":{"arxiv_id":"2401.04658","doi":"10.48550/arxiv.2401.04658","metadata_source":"arxiv_reference","pith_arxiv_id":"2401.04658","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Lightning attention-2: A free lunch for handling unlimited sequence lengths in large language models.arXiv preprint arXiv:2401.04658","venue":"arXiv (Cornell University)","work_id":"02c34d0b-5e0a-48e5-9c0d-9d5474fb4c00","year":2024},"citing_paper":{"arxiv_id":"2605.23081","last_updated":"2026-05-21T22:28:27Z","snapshot_observed_at":"2026-08-04T11:05:52.350024Z","submitted_at":"2026-05-21T22:28:27Z","title":"ThriftAttention: Selective Mixed Precision for Long-Context FP4 Attention","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-05-25T05:28:40.888215Z"},"links":{"cited_paper":"/paper/2401.04658","citing_paper":"/paper/2605.23081"},"observation_digest":"sha256:939a18360502b918f817784604b3d6fb167e86ff812215ac5f144440b4980ca6","observation_id":"fa0b417a-bda2-4b71-ba44-4fc0ae69d3f4","resolution":{"observed_at":"2026-05-25T05:30:22.840962Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2401.04658","last_updated":"2024-01-15T14:57:29Z","snapshot_observed_at":"2026-08-10T13:02:55.737137Z","submitted_at":"2024-01-09T16:27:28Z","title":"Lightning Attention-2: A Free Lunch for Handling Unlimited Sequence Lengths in Large Language Models","version":2},"cited_work":{"arxiv_id":"2401.04658","doi":"10.48550/arxiv.2401.04658","metadata_source":"arxiv_reference","pith_arxiv_id":"2401.04658","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Lightning attention-2: A free lunch for handling unlimited sequence lengths in large language models.arXiv preprint arXiv:2401.04658","venue":"arXiv (Cornell University)","work_id":"02c34d0b-5e0a-48e5-9c0d-9d5474fb4c00","year":2024},"citing_paper":{"arxiv_id":"2606.03825","last_updated":"2026-06-02T16:07:55Z","snapshot_observed_at":"2026-08-03T23:42:45.221248Z","submitted_at":"2026-06-02T16:07:55Z","title":"Dynamic Short Convolutions Improve Transformers","version":1},"reference_index":87,"source":"arxiv_source","source_observed_at":"2026-06-28T10:48:50.103004Z"},"links":{"cited_paper":"/paper/2401.04658","citing_paper":"/paper/2606.03825"},"observation_digest":"sha256:16633ab85078454c030bb0662f26df50c65c4cf30f1c5ddbb0658ab4a25af46e","observation_id":"ecbae65f-9ade-4f20-a27c-0b70ebf09afc","resolution":{"observed_at":"2026-07-02T02:36:27.012767Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2401.04658","last_updated":"2024-01-15T14:57:29Z","snapshot_observed_at":"2026-08-10T13:02:55.737137Z","submitted_at":"2024-01-09T16:27:28Z","title":"Lightning Attention-2: A Free Lunch for Handling Unlimited Sequence Lengths in Large Language Models","version":2},"cited_work":{"arxiv_id":"2401.04658","doi":"10.48550/arxiv.2401.04658","metadata_source":"arxiv_reference","pith_arxiv_id":"2401.04658","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Lightning attention-2: A free lunch for handling unlimited sequence lengths in large language models.arXiv preprint arXiv:2401.04658","venue":"arXiv (Cornell University)","work_id":"02c34d0b-5e0a-48e5-9c0d-9d5474fb4c00","year":2024},"citing_paper":{"arxiv_id":"2606.11634","last_updated":"2026-06-10T03:56:03Z","snapshot_observed_at":"2026-08-02T13:37:43.737297Z","submitted_at":"2026-06-10T03:56:03Z","title":"Architecture-Aware Reinforcement Learning Makes Sliding-Window Attention Competitive in Math Reasoning","version":1},"reference_index":24,"source":"arxiv_source","source_observed_at":"2026-06-27T10:18:54.163862Z"},"links":{"cited_paper":"/paper/2401.04658","citing_paper":"/paper/2606.11634"},"observation_digest":"sha256:7493066b02fae8c1cf9f63609525c60ae86190b5c0a5d05f01943ef3a78e54ba","observation_id":"d7fabafa-d1f5-400e-9d54-9a99635e3c7f","resolution":{"observed_at":"2026-07-03T09:47:59.878110Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2401.04658","last_updated":"2024-01-15T14:57:29Z","snapshot_observed_at":"2026-08-10T13:02:55.737137Z","submitted_at":"2024-01-09T16:27:28Z","title":"Lightning Attention-2: A Free Lunch for Handling Unlimited Sequence Lengths in Large Language Models","version":2},"cited_work":{"arxiv_id":"2401.04658","doi":"10.48550/arxiv.2401.04658","metadata_source":"arxiv_reference","pith_arxiv_id":"2401.04658","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Lightning attention-2: A free lunch for handling unlimited sequence lengths in large language models.arXiv preprint arXiv:2401.04658","venue":"arXiv (Cornell University)","work_id":"02c34d0b-5e0a-48e5-9c0d-9d5474fb4c00","year":2024},"citing_paper":{"arxiv_id":"2606.19853","last_updated":"2026-06-18T07:01:14Z","snapshot_observed_at":"2026-08-08T03:27:10.046575Z","submitted_at":"2026-06-18T07:01:14Z","title":"Physics-Informed Neural Network with Squeeze-Excitation-like Attention","version":1},"reference_index":57,"source":"pdf_text","source_observed_at":"2026-06-26T18:02:39.721547Z"},"links":{"cited_paper":"/paper/2401.04658","citing_paper":"/paper/2606.19853"},"observation_digest":"sha256:8997580ae0ee1f074665c7810fbc61434c2094a3f8366f6ab2c7e2d687a1c604","observation_id":"90107221-b4ec-47be-946d-3037a2b38362","resolution":{"observed_at":"2026-07-04T03:29:30.199130Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2401.04658","last_updated":"2024-01-15T14:57:29Z","snapshot_observed_at":"2026-08-10T13:02:55.737137Z","submitted_at":"2024-01-09T16:27:28Z","title":"Lightning Attention-2: A Free Lunch for Handling Unlimited Sequence Lengths in Large Language Models","version":2},"cited_work":{"arxiv_id":"2401.04658","doi":"10.48550/arxiv.2401.04658","metadata_source":"arxiv_reference","pith_arxiv_id":"2401.04658","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Lightning attention-2: A free lunch for handling unlimited sequence lengths in large language models.arXiv preprint arXiv:2401.04658","venue":"arXiv (Cornell University)","work_id":"02c34d0b-5e0a-48e5-9c0d-9d5474fb4c00","year":2024},"citing_paper":{"arxiv_id":"2606.30562","last_updated":"2026-06-29T17:02:34Z","snapshot_observed_at":"2026-07-07T00:04:25.385438Z","submitted_at":"2026-06-29T17:02:34Z","title":"Morphing into Hybrid Attention Models","version":1},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-06-30T05:56:51.447893Z"},"links":{"cited_paper":"/paper/2401.04658","citing_paper":"/paper/2606.30562"},"observation_digest":"sha256:915cfcbeccce026cba5274049b3ff27323a2772dfe1615d7364d65b054591317","observation_id":"375ebb43-7efb-4d36-8eec-854e9478d337","resolution":{"observed_at":"2026-06-30T08:44:28.072509Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2401.04658","last_updated":"2024-01-15T14:57:29Z","snapshot_observed_at":"2026-08-10T13:02:55.737137Z","submitted_at":"2024-01-09T16:27:28Z","title":"Lightning Attention-2: A Free Lunch for Handling Unlimited Sequence Lengths in Large Language Models","version":2},"cited_work":{"arxiv_id":"2401.04658","doi":"10.48550/arxiv.2401.04658","metadata_source":"arxiv_reference","pith_arxiv_id":"2401.04658","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Lightning attention-2: A free lunch for handling unlimited sequence lengths in large language models.arXiv preprint arXiv:2401.04658","venue":"arXiv (Cornell University)","work_id":"02c34d0b-5e0a-48e5-9c0d-9d5474fb4c00","year":2024},"citing_paper":{"arxiv_id":"2607.01299","last_updated":"2026-07-12T15:14:05Z","snapshot_observed_at":"2026-08-06T01:10:35.623891Z","submitted_at":"2026-07-01T14:03:56Z","title":"HYPIC: Accelerating Hybrid-Attention LLM Serving with Position-Independent Caching","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-07-03T18:53:12.126023Z"},"links":{"cited_paper":"/paper/2401.04658","citing_paper":"/paper/2607.01299"},"observation_digest":"sha256:ae36e4ad4a8d1c6913cda60a7abd5ccf2c572e349c219ab7544b759881715b73","observation_id":"0ba81ed1-a556-4cdd-b9e0-1437a75629f7","resolution":{"observed_at":"2026-07-03T18:58:50.719338Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-12T06:34:41.77262+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2401.04658","last_updated":"2024-01-15T14:57:29Z","snapshot_observed_at":"2026-08-10T13:02:55.737137Z","submitted_at":"2024-01-09T16:27:28Z","title":"Lightning Attention-2: A Free Lunch for Handling Unlimited Sequence Lengths in Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.04658","snapshot_observed_at":"2026-07-14T16:49:27.301911Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.01299","last_updated":"2026-07-12T15:14:05Z","snapshot_observed_at":"2026-08-06T01:10:35.623891Z","submitted_at":"2026-07-01T14:03:56Z","title":"HYPIC: Accelerating Hybrid-Attention LLM Serving with Position-Independent Caching","version":2},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-07-14T16:49:27.301911Z"},"links":{"cited_paper":"/paper/2401.04658","citing_paper":"/paper/2607.01299"},"observation_digest":"sha256:e3c4d8d3bb2c50ffa0f1bfc0dbf9abcea5f5cb46a07a0034b861407a2522e766","observation_id":"4988e43d-14e9-4bfe-8888-c123aa8877ca","resolution":{"observed_at":"2026-07-14T16:49:27.301911Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2401.04658","last_updated":"2024-01-15T14:57:29Z","snapshot_observed_at":"2026-08-10T13:02:55.737137Z","submitted_at":"2024-01-09T16:27:28Z","title":"Lightning Attention-2: A Free Lunch for Handling Unlimited Sequence Lengths in Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2401.04658","snapshot_observed_at":"2026-08-02T06:37:38.357406Z","title":"Lightning attention-2: A free lunch for handling unlimited sequence lengths in large language models","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.12395","last_updated":"2026-07-16T03:43:34Z","snapshot_observed_at":"2026-08-11T12:30:37.966398Z","submitted_at":"2026-07-14T06:14:55Z","title":"Ring-Zero: Scaling Zero RL to a Trillion Parameters for Emergent Reasoning","version":2},"reference_index":66,"source":"arxiv_source","source_observed_at":"2026-08-02T06:37:38.357406Z"},"links":{"cited_paper":"/paper/2401.04658","citing_paper":"/paper/2607.12395"},"observation_digest":"sha256:1c6dcb79f113ab66d51bad5d684db338f6ef73a1afde7b1d124f863bca433bb7","observation_id":"36ba925c-029b-4b53-afd1-c6e80ab1e360","resolution":{"observed_at":"2026-08-02T06:37:38.357406Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2401.04658/citation-record","integrity":"/paper/2401.04658/integrity","json":"/paper/2401.04658/citation-record.json","paper":"/paper/2401.04658"},"outbound":[],"paper":{"arxiv_id":"2401.04658","last_updated":"2024-01-15T14:57:29Z","latest_version":2,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-10T13:02:55.737137Z","submitted_at":"2024-01-09T16:27:28Z","title":"Lightning Attention-2: A Free Lunch for Handling Unlimited Sequence Lengths in Large Language Models"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-12T06:34:41.77262+00:00","source":"crossref"},{"observed_at":"2026-08-12T06:34:36.333875+00:00","source":"retraction_watch"}],"thesis":"As of 12 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 18 inbound Pith citation observations for arXiv:2401.04658."}