{"as_of":"2026-08-05T14:52:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:4169e1750b71a2f87916d5a6263288d98398069dfee79219cf6b8a949070e948","coverage":[{"denominator":33,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":33,"source":"paper_references, paper_reference_links","source_observed_at":"2026-05-11T02:27:55.991919Z","state":"measured"},{"denominator":35,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":35,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-05T06:32:48.257954+00:00","state":"measured"},{"denominator":2,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":2,"source":"paper_references, paper_reference_links","source_observed_at":"2026-07-11T01:55:09.658053Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"pith","source_observed_at":"2026-07-11T01:57:51.710279Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2605.07363","last_updated":"2026-05-08T07:19:34Z","snapshot_observed_at":"2026-07-06T23:19:46.018236Z","submitted_at":"2026-05-08T07:19:34Z","title":"MISA: Mixture of Indexer Sparse Attention for Long-Context LLM Inference","version":1},"cited_work":{"arxiv_id":"2605.07363","doi":null,"metadata_source":"pith","pith_arxiv_id":"2605.07363","snapshot_observed_at":"2026-07-11T01:57:51.710279Z","title":"MISA: Mixture of Indexer Sparse Attention for Long-Context LLM Inference","venue":"cs.LG","work_id":"21df7960-841c-4aaf-8937-b93c56be78d6","year":2026},"citing_paper":{"arxiv_id":"2607.05876","last_updated":"2026-07-08T02:47:41Z","snapshot_observed_at":"2026-08-05T07:46:13.093303Z","submitted_at":"2026-07-07T06:11:54Z","title":"Think Before You Grid-Search: Floor-First Triage for LLM Serving","version":1},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-07-08T22:38:12.637901Z"},"links":{"cited_paper":"/paper/2605.07363","citing_paper":"/paper/2607.05876"},"observation_digest":"sha256:0efe51a5a7ce15bdf66404b4df1dedb2c6b66e2c6c1350d8db3a9a0e81ae644d","observation_id":"345229f4-aeeb-4d4d-a3f1-4401f8c9087d","resolution":{"observed_at":"2026-07-08T22:45:40.112264Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2605.07363","last_updated":"2026-05-08T07:19:34Z","snapshot_observed_at":"2026-07-06T23:19:46.018236Z","submitted_at":"2026-05-08T07:19:34Z","title":"MISA: Mixture of Indexer Sparse Attention for Long-Context LLM Inference","version":1},"cited_work":{"arxiv_id":"2605.07363","doi":null,"metadata_source":"pith","pith_arxiv_id":"2605.07363","snapshot_observed_at":"2026-07-11T01:57:51.710279Z","title":"MISA: Mixture of Indexer Sparse Attention for Long-Context LLM Inference","venue":"cs.LG","work_id":"21df7960-841c-4aaf-8937-b93c56be78d6","year":2026},"citing_paper":{"arxiv_id":"2607.05876","last_updated":"2026-07-08T02:47:41Z","snapshot_observed_at":"2026-08-05T07:46:13.093303Z","submitted_at":"2026-07-07T06:11:54Z","title":"Think Before You Grid-Search: Floor-First Triage for LLM Serving","version":2},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-07-11T01:55:09.658053Z"},"links":{"cited_paper":"/paper/2605.07363","citing_paper":"/paper/2607.05876"},"observation_digest":"sha256:5142978a96891bddfce94daa6ac33ce1f498cef344a668884de20147ccf0ba42","observation_id":"0ee00fb2-e902-4d49-9563-c6e5fda339f6","resolution":{"observed_at":"2026-07-11T01:57:51.736154Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2605.07363/citation-record","integrity":"/paper/2605.07363/integrity","json":"/paper/2605.07363/citation-record.json","paper":"/paper/2605.07363"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Introducing Claude Opus 4.7","venue":null,"work_id":"54340b24-c285-4b1c-9593-47631555c346","year":2026},"citing_paper":{"arxiv_id":"2605.07363","last_updated":"2026-05-08T07:19:34Z","snapshot_observed_at":"2026-07-06T23:19:46.018236Z","submitted_at":"2026-05-08T07:19:34Z","title":"MISA: Mixture of Indexer Sparse Attention for Long-Context LLM Inference","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-05-11T02:27:55.991919Z"},"links":{"citing_paper":"/paper/2605.07363"},"observation_digest":"sha256:4e0144bf77b3e4780721f1f05c5696b8aa08c0443d0836e7b39b884ddd6f338c","observation_id":"a9a36d0f-9678-4daa-8a90-41162c9ea171","resolution":{"observed_at":"2026-05-14T12:35:15.410226Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2603.12201","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-02T13:36:59.567865Z","title":"Indexcache: Accelerating sparse attention via cross-layer index reuse","venue":null,"work_id":"a3414223-70e8-45f8-95f6-3113f1b1facb","year":2026},"citing_paper":{"arxiv_id":"2605.07363","last_updated":"2026-05-08T07:19:34Z","snapshot_observed_at":"2026-07-06T23:19:46.018236Z","submitted_at":"2026-05-08T07:19:34Z","title":"MISA: Mixture of Indexer Sparse Attention for Long-Context LLM Inference","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-05-11T02:27:55.991919Z"},"links":{"citing_paper":"/paper/2605.07363"},"observation_digest":"sha256:8e07dc7ac33dac1cd6b2a35e351ca33e9b8871d88ce9dc789c1c1b32cdee95a9","observation_id":"f522e26b-f321-4175-bd7c-f49df9d713cd","resolution":{"observed_at":"2026-05-11T03:30:56.563898Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2004.05150","last_updated":"2020-12-02T17:52:35Z","snapshot_observed_at":"2026-07-31T17:17:17.205582Z","submitted_at":"2020-04-10T17:54:09Z","title":"Longformer: The Long-Document Transformer","version":2},"cited_work":{"arxiv_id":"2004.05150","doi":"10.48550/arxiv.2004.05150","metadata_source":"pith","pith_arxiv_id":"2004.05150","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Longformer: The Long-Document Transformer","venue":"cs.CL","work_id":"abea7a44-6668-4de7-aab6-f53a6e5aa088","year":2020},"citing_paper":{"arxiv_id":"2605.07363","last_updated":"2026-05-08T07:19:34Z","snapshot_observed_at":"2026-07-06T23:19:46.018236Z","submitted_at":"2026-05-08T07:19:34Z","title":"MISA: Mixture of Indexer Sparse Attention for Long-Context LLM Inference","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-05-11T02:27:55.991919Z"},"links":{"cited_paper":"/paper/2004.05150","citing_paper":"/paper/2605.07363"},"observation_digest":"sha256:acdbfdcbb4d4a69a9035414ae4fe9f7dd2eb800f2775c53a00afe415c4672957","observation_id":"c12e082b-a273-4c74-b016-0b9ccbbd83d3","resolution":{"observed_at":"2026-05-11T03:30:56.451892Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-12T21:49:59.161233+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T21:49:59.161233+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.16179","last_updated":"2024-12-18T17:36:36Z","snapshot_observed_at":"2026-07-06T19:37:11.636987Z","submitted_at":"2024-10-21T16:44:51Z","title":"MagicPIG: LSH Sampling for Efficient LLM Generation","version":4},"cited_work":{"arxiv_id":"2410.16179","doi":"10.48550/arxiv.2410.16179","metadata_source":"arxiv_reference","pith_arxiv_id":"2410.16179","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"arXiv preprint arXiv:2410.16179 , year=","venue":"arXiv (Cornell University)","work_id":"24103510-199d-40aa-ae36-22602d72d97f","year":2021},"citing_paper":{"arxiv_id":"2605.07363","last_updated":"2026-05-08T07:19:34Z","snapshot_observed_at":"2026-07-06T23:19:46.018236Z","submitted_at":"2026-05-08T07:19:34Z","title":"MISA: Mixture of Indexer Sparse Attention for Long-Context LLM Inference","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-05-11T02:27:55.991919Z"},"links":{"cited_paper":"/paper/2410.16179","citing_paper":"/paper/2605.07363"},"observation_digest":"sha256:4b552fb0527467db31348e7a5b16fa788f305186df8d626f7655c984b451d45c","observation_id":"0e5b6112-8330-45fc-9d69-023112795595","resolution":{"observed_at":"2026-05-11T03:30:56.411973Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1904.10509","last_updated":"2019-04-23T19:29:47Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2019-04-23T19:29:47Z","title":"Generating Long Sequences with Sparse Transformers","version":1},"cited_work":{"arxiv_id":"1904.10509","doi":"10.48550/arxiv.1904.10509","metadata_source":"pith","pith_arxiv_id":"1904.10509","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Generating Long Sequences with Sparse Transformers","venue":"cs.LG","work_id":"c5b81688-45ee-4a9a-b095-e6290f45cb6c","year":2019},"citing_paper":{"arxiv_id":"2605.07363","last_updated":"2026-05-08T07:19:34Z","snapshot_observed_at":"2026-07-06T23:19:46.018236Z","submitted_at":"2026-05-08T07:19:34Z","title":"MISA: Mixture of Indexer Sparse Attention for Long-Context LLM Inference","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-05-11T02:27:55.991919Z"},"links":{"cited_paper":"/paper/1904.10509","citing_paper":"/paper/2605.07363"},"observation_digest":"sha256:587f062a18403f2bbae742c247834f72f1f0951c71bf1d004a5c21ad2ddfee70","observation_id":"ce913117-6874-4b6f-a249-33bef6449870","resolution":{"observed_at":"2026-05-11T03:30:56.502407Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-12T21:49:59.71041+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T21:49:59.71041+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2401.06066","last_updated":"2024-01-11T17:31:42Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-01-11T17:31:42Z","title":"DeepSeekMoE: Towards Ultimate Expert Specialization in Mixture-of-Experts Language Models","version":1},"cited_work":{"arxiv_id":"2401.06066","doi":"10.48550/arxiv.2401.06066","metadata_source":"pith","pith_arxiv_id":"2401.06066","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"DeepSeekMoE: Towards Ultimate Expert Specialization in Mixture-of-Experts Language Models","venue":"cs.CL","work_id":"a9888d6d-bf47-4324-9834-7cc12ac3a78c","year":2024},"citing_paper":{"arxiv_id":"2605.07363","last_updated":"2026-05-08T07:19:34Z","snapshot_observed_at":"2026-07-06T23:19:46.018236Z","submitted_at":"2026-05-08T07:19:34Z","title":"MISA: Mixture of Indexer Sparse Attention for Long-Context LLM Inference","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-05-11T02:27:55.991919Z"},"links":{"cited_paper":"/paper/2401.06066","citing_paper":"/paper/2605.07363"},"observation_digest":"sha256:0bbff9ab69507a5893a97e59dfd8d8e31273cb0c22f45be5e56f37df746432a4","observation_id":"dae1dbdb-ce78-47b5-a5cf-a418b972445a","resolution":{"observed_at":"2026-05-11T22:50:20.439308Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2512.02556","last_updated":"2025-12-02T09:25:14Z","snapshot_observed_at":"2026-07-31T23:49:25.878472Z","submitted_at":"2025-12-02T09:25:14Z","title":"DeepSeek-V3.2: Pushing the Frontier of Open Large Language Models","version":1},"cited_work":{"arxiv_id":"2512.02556","doi":"10.18653/v1/d18-1512","metadata_source":"pith","pith_arxiv_id":"2512.02556","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"DeepSeek-V3.2: Pushing the Frontier of Open Large Language Models","venue":"cs.CL","work_id":"07c85cc5-4086-4abc-823b-6d0f4ff784d0","year":2025},"citing_paper":{"arxiv_id":"2605.07363","last_updated":"2026-05-08T07:19:34Z","snapshot_observed_at":"2026-07-06T23:19:46.018236Z","submitted_at":"2026-05-08T07:19:34Z","title":"MISA: Mixture of Indexer Sparse Attention for Long-Context LLM Inference","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-05-11T02:27:55.991919Z"},"links":{"cited_paper":"/paper/2512.02556","citing_paper":"/paper/2605.07363"},"observation_digest":"sha256:b75cd4dcabf74cbf7f4b9d1a901ee79d18671fd01224cde513ee02a2aa960fa1","observation_id":"c507ef6a-6acd-40cb-853c-6bde391abdb0","resolution":{"observed_at":"2026-05-11T03:30:56.394117Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"DeepSeek-V4: Towards highly efficient million-token context.Technical Report, DeepSeek","venue":null,"work_id":"77290725-69e9-4b4d-a8ea-d0ce1dff7988","year":2026},"citing_paper":{"arxiv_id":"2605.07363","last_updated":"2026-05-08T07:19:34Z","snapshot_observed_at":"2026-07-06T23:19:46.018236Z","submitted_at":"2026-05-08T07:19:34Z","title":"MISA: Mixture of Indexer Sparse Attention for Long-Context LLM Inference","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-05-11T02:27:55.991919Z"},"links":{"citing_paper":"/paper/2605.07363"},"observation_digest":"sha256:e064e4e9885589b3b23b008f90f505e16b187620622e473b6bab291de1b995f9","observation_id":"986ec7ee-9894-4022-acd4-a628d0cde128","resolution":{"observed_at":"2026-05-14T12:35:15.391428Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2101.03961","last_updated":"2022-06-16T20:36:07Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2021-01-11T16:11:52Z","title":"Switch Transformers: Scaling to Trillion Parameter Models with Simple and Efficient Sparsity","version":3},"cited_work":{"arxiv_id":"2101.03961","doi":"10.1214/18-ejs1395","metadata_source":"pith","pith_arxiv_id":"2101.03961","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Switch Transformers: Scaling to Trillion Parameter Models with Simple and Efficient Sparsity","venue":"cs.LG","work_id":"f43c4955-a965-4897-a11b-c4b25d2aeaa8","year":2021},"citing_paper":{"arxiv_id":"2605.07363","last_updated":"2026-05-08T07:19:34Z","snapshot_observed_at":"2026-07-06T23:19:46.018236Z","submitted_at":"2026-05-08T07:19:34Z","title":"MISA: Mixture of Indexer Sparse Attention for Long-Context LLM Inference","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-05-11T02:27:55.991919Z"},"links":{"cited_paper":"/paper/2101.03961","citing_paper":"/paper/2605.07363"},"observation_digest":"sha256:9df432fbf67f0675945cfda13da4728c832fc933c20ff0d5cac94c14d6d8f2f3","observation_id":"c2a0c222-1184-49d8-8db6-b038287e8c2a","resolution":{"observed_at":"2026-05-12T23:57:11.134962Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.13276","last_updated":"2025-02-17T02:24:47Z","snapshot_observed_at":"2026-08-05T05:19:17.395325Z","submitted_at":"2024-10-17T07:07:09Z","title":"SeerAttention: Learning Intrinsic Sparse Attention in Your LLMs","version":4},"cited_work":{"arxiv_id":"2410.13276","doi":null,"metadata_source":"pith","pith_arxiv_id":"2410.13276","snapshot_observed_at":"2026-07-10T01:36:44.065347Z","title":"Seerattention: Learning intrinsic sparse attention in your llms.arXiv preprint arXiv:2410.13276","venue":"cs.CL","work_id":"7e69f9a0-9472-4c02-8f79-806ba81f8531","year":2024},"citing_paper":{"arxiv_id":"2605.07363","last_updated":"2026-05-08T07:19:34Z","snapshot_observed_at":"2026-07-06T23:19:46.018236Z","submitted_at":"2026-05-08T07:19:34Z","title":"MISA: Mixture of Indexer Sparse Attention for Long-Context LLM Inference","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-05-11T02:27:55.991919Z"},"links":{"cited_paper":"/paper/2410.13276","citing_paper":"/paper/2605.07363"},"observation_digest":"sha256:3ac23c7afc6ea827ff4b29db73d3df6e8833b65d52a522a57c7a2a4930ca3203","observation_id":"e91831c1-8e6e-4b27-b64d-7763b5db5d4b","resolution":{"observed_at":"2026-05-11T03:30:56.540017Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Gemini 3: A new era of intelligence","venue":null,"work_id":"2370052f-3b6e-475c-8382-559a60591250","year":2026},"citing_paper":{"arxiv_id":"2605.07363","last_updated":"2026-05-08T07:19:34Z","snapshot_observed_at":"2026-07-06T23:19:46.018236Z","submitted_at":"2026-05-08T07:19:34Z","title":"MISA: Mixture of Indexer Sparse Attention for Long-Context LLM Inference","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-05-11T02:27:55.991919Z"},"links":{"citing_paper":"/paper/2605.07363"},"observation_digest":"sha256:4a4e0af1756660d23a757e573ae8dd2eb3158cb7545819260eebdc3b53c7fe05","observation_id":"fc0c9dd7-c851-415c-b97a-e4a13808ef54","resolution":{"observed_at":"2026-05-14T12:35:15.403225Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2401.04088","last_updated":"2024-01-08T18:47:34Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-01-08T18:47:34Z","title":"Mixtral of Experts","version":1},"cited_work":{"arxiv_id":"2401.04088","doi":"10.48550/arxiv.2401.04088","metadata_source":"pith","pith_arxiv_id":"2401.04088","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Mixtral of Experts","venue":"cs.LG","work_id":"0de8c352-9daa-4e1e-8c7b-3d0dec69f369","year":2024},"citing_paper":{"arxiv_id":"2605.07363","last_updated":"2026-05-08T07:19:34Z","snapshot_observed_at":"2026-07-06T23:19:46.018236Z","submitted_at":"2026-05-08T07:19:34Z","title":"MISA: Mixture of Indexer Sparse Attention for Long-Context LLM Inference","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-05-11T02:27:55.991919Z"},"links":{"cited_paper":"/paper/2401.04088","citing_paper":"/paper/2605.07363"},"observation_digest":"sha256:4cdd9e172b904974eef756c6ee1c489209ecea495e6c09e1c52fa45c5204e69f","observation_id":"0c7df7b1-847f-4004-950b-729561a0263b","resolution":{"observed_at":"2026-05-11T03:30:56.546230Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-09T08:48:39.110013+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-09T08:48:39.110013+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2410.11842","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-04T03:49:29.882562Z","title":"Moh: Multi-head attention as mixture-of-head attention","venue":null,"work_id":"13a82af3-d0e4-45fd-acc6-1abdc537d0ee","year":2024},"citing_paper":{"arxiv_id":"2605.07363","last_updated":"2026-05-08T07:19:34Z","snapshot_observed_at":"2026-07-06T23:19:46.018236Z","submitted_at":"2026-05-08T07:19:34Z","title":"MISA: Mixture of Indexer Sparse Attention for Long-Context LLM Inference","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-05-11T02:27:55.991919Z"},"links":{"citing_paper":"/paper/2605.07363"},"observation_digest":"sha256:717a4d57a4fcbe1542580565ff5554e53e27b5eff360d72db20caf0cfcc2d5f0","observation_id":"8f4ef40e-8dc8-4034-98c6-001b3517bcc0","resolution":{"observed_at":"2026-05-11T03:30:56.463733Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2006.16668","last_updated":"2020-06-30T10:42:02Z","snapshot_observed_at":"2026-07-06T09:33:58.857566Z","submitted_at":"2020-06-30T10:42:02Z","title":"GShard: Scaling Giant Models with Conditional Computation and Automatic Sharding","version":1},"cited_work":{"arxiv_id":"2006.16668","doi":"10.48550/arxiv.2006.16668","metadata_source":"pith","pith_arxiv_id":"2006.16668","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"GShard: Scaling Giant Models with Conditional Computation and Automatic Sharding","venue":"cs.CL","work_id":"52b3c9a6-2a27-45a7-ba2b-ebe4b5bb5a5f","year":2020},"citing_paper":{"arxiv_id":"2605.07363","last_updated":"2026-05-08T07:19:34Z","snapshot_observed_at":"2026-07-06T23:19:46.018236Z","submitted_at":"2026-05-08T07:19:34Z","title":"MISA: Mixture of Indexer Sparse Attention for Long-Context LLM Inference","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-05-11T02:27:55.991919Z"},"links":{"cited_paper":"/paper/2006.16668","citing_paper":"/paper/2605.07363"},"observation_digest":"sha256:2ff95d52159fc7a985bf6253d83646bf627100720a7d900b9e925150f0491432","observation_id":"37daca28-299e-4d2c-b503-4c22c27e29c6","resolution":{"observed_at":"2026-05-11T03:30:56.551559Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.08313","last_updated":"2025-01-14T18:50:05Z","snapshot_observed_at":"2026-07-06T20:21:04.675084Z","submitted_at":"2025-01-14T18:50:05Z","title":"MiniMax-01: Scaling Foundation Models with Lightning Attention","version":1},"cited_work":{"arxiv_id":"2501.08313","doi":"10.48550/arxiv.2501.08313","metadata_source":"pith","pith_arxiv_id":"2501.08313","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"MiniMax-01: Scaling Foundation Models with Lightning Attention","venue":"cs.CL","work_id":"137e6ed4-1c1d-40bd-801e-9d2baea0bd0c","year":2025},"citing_paper":{"arxiv_id":"2605.07363","last_updated":"2026-05-08T07:19:34Z","snapshot_observed_at":"2026-07-06T23:19:46.018236Z","submitted_at":"2026-05-08T07:19:34Z","title":"MISA: Mixture of Indexer Sparse Attention for Long-Context LLM Inference","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-05-11T02:27:55.991919Z"},"links":{"cited_paper":"/paper/2501.08313","citing_paper":"/paper/2605.07363"},"observation_digest":"sha256:39c730e3c7d031b90f05729e2cc62b4eb1857147f7bff2b2c82cb94999d3d77b","observation_id":"1d8eaca9-34ec-4343-ad59-a6f847d3af5e","resolution":{"observed_at":"2026-05-16T06:26:38.921226Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2404.14469","last_updated":"2024-06-17T03:01:58Z","snapshot_observed_at":"2026-07-06T18:03:55.981890Z","submitted_at":"2024-04-22T17:42:58Z","title":"SnapKV: LLM Knows What You are Looking for Before Generation","version":2},"cited_work":{"arxiv_id":"2404.14469","doi":"10.48550/arxiv.2404.14469","metadata_source":"pith","pith_arxiv_id":"2404.14469","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SnapKV: LLM Knows What You are Looking for Before Generation","venue":"cs.CL","work_id":"4afe6cb0-dba4-42b6-b553-e685e5730d61","year":2024},"citing_paper":{"arxiv_id":"2605.07363","last_updated":"2026-05-08T07:19:34Z","snapshot_observed_at":"2026-07-06T23:19:46.018236Z","submitted_at":"2026-05-08T07:19:34Z","title":"MISA: Mixture of Indexer Sparse Attention for Long-Context LLM Inference","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-05-11T02:27:55.991919Z"},"links":{"cited_paper":"/paper/2404.14469","citing_paper":"/paper/2605.07363"},"observation_digest":"sha256:656e4a59dd18aedb8d808377a6a37c9116f3cb79274971a71b827ec15a1e6a42","observation_id":"7a1255c1-0344-4ec4-b9e1-8c09edd69192","resolution":{"observed_at":"2026-05-13T12:57:43.256882Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-05-25T13:53:27.134462+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-25T13:53:27.134462+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.13189","last_updated":"2025-02-18T14:06:05Z","snapshot_observed_at":"2026-07-06T20:38:48.725605Z","submitted_at":"2025-02-18T14:06:05Z","title":"MoBA: Mixture of Block Attention for Long-Context LLMs","version":1},"cited_work":{"arxiv_id":"2502.13189","doi":"10.48550/arxiv.2502.13189","metadata_source":"pith","pith_arxiv_id":"2502.13189","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"MoBA: Mixture of Block Attention for Long-Context LLMs","venue":"cs.LG","work_id":"de16df2d-eb7f-429d-82c7-8e4bf5183193","year":2025},"citing_paper":{"arxiv_id":"2605.07363","last_updated":"2026-05-08T07:19:34Z","snapshot_observed_at":"2026-07-06T23:19:46.018236Z","submitted_at":"2026-05-08T07:19:34Z","title":"MISA: Mixture of Indexer Sparse Attention for Long-Context LLM Inference","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-05-11T02:27:55.991919Z"},"links":{"cited_paper":"/paper/2502.13189","citing_paper":"/paper/2605.07363"},"observation_digest":"sha256:fa3fc62dfb1da610a48ee827d5c4d359a685026fb04b5564b3e53cfaf12c4d71","observation_id":"b0897f00-1ee1-4dca-8059-d00e17609658","resolution":{"observed_at":"2026-05-16T06:15:46.352230Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Kimi K2: Open agentic intelligence","venue":null,"work_id":"8916c70e-a419-457b-80d6-e0fd773d84fe","year":2025},"citing_paper":{"arxiv_id":"2605.07363","last_updated":"2026-05-08T07:19:34Z","snapshot_observed_at":"2026-07-06T23:19:46.018236Z","submitted_at":"2026-05-08T07:19:34Z","title":"MISA: Mixture of Indexer Sparse Attention for Long-Context LLM Inference","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-05-11T02:27:55.991919Z"},"links":{"citing_paper":"/paper/2605.07363"},"observation_digest":"sha256:c9d0cececd6057f9d6a8f9d3e7cd5ff21bcb05ca0bb8a776b548fb6e162e9ac0","observation_id":"dcf4542a-904c-4e80-8d22-9a7547b698ea","resolution":{"observed_at":"2026-05-14T12:35:15.395974Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Introducing GPT-5.5","venue":null,"work_id":"50acbef5-8d85-4ac9-b441-4a4cc364b288","year":2026},"citing_paper":{"arxiv_id":"2605.07363","last_updated":"2026-05-08T07:19:34Z","snapshot_observed_at":"2026-07-06T23:19:46.018236Z","submitted_at":"2026-05-08T07:19:34Z","title":"MISA: Mixture of Indexer Sparse Attention for Long-Context LLM Inference","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-05-11T02:27:55.991919Z"},"links":{"citing_paper":"/paper/2605.07363"},"observation_digest":"sha256:12b6c037bf340f2a2fe112ab634ef535fa67e5a13c761b4440be96b1951e8ed5","observation_id":"3238e8bf-6309-4218-890b-477adb94e268","resolution":{"observed_at":"2026-05-14T12:35:15.413675Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1701.06538","last_updated":"2017-01-23T18:10:00Z","snapshot_observed_at":"2026-07-06T05:27:13.416519Z","submitted_at":"2017-01-23T18:10:00Z","title":"Outrageously Large Neural Networks: The Sparsely-Gated Mixture-of-Experts Layer","version":1},"cited_work":{"arxiv_id":"1701.06538","doi":"10.48550/arxiv.1701.06538","metadata_source":"pith","pith_arxiv_id":"1701.06538","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Outrageously Large Neural Networks: The Sparsely-Gated Mixture-of-Experts Layer","venue":"cs.LG","work_id":"2c6b3f6d-54e4-4df7-baa7-475a490799af","year":2017},"citing_paper":{"arxiv_id":"2605.07363","last_updated":"2026-05-08T07:19:34Z","snapshot_observed_at":"2026-07-06T23:19:46.018236Z","submitted_at":"2026-05-08T07:19:34Z","title":"MISA: Mixture of Indexer Sparse Attention for Long-Context LLM Inference","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-05-11T02:27:55.991919Z"},"links":{"cited_paper":"/paper/1701.06538","citing_paper":"/paper/2605.07363"},"observation_digest":"sha256:1f23912f7a0dc92b193a5d23d486749303c157e19f89059583e86456fe82bace","observation_id":"aa952ab4-9b27-4adf-9db9-c79e3bf76b62","resolution":{"observed_at":"2026-05-11T03:30:56.508389Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-01T08:08:24.174744+00:00","source":"crossref_status_cache"},{"observed_at":"2026-08-01T08:08:24.174744+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.10774","last_updated":"2024-08-26T21:01:02Z","snapshot_observed_at":"2026-07-06T18:31:36.504766Z","submitted_at":"2024-06-16T01:33:02Z","title":"Quest: Query-Aware Sparsity for Efficient Long-Context LLM Inference","version":2},"cited_work":{"arxiv_id":"2406.10774","doi":"10.48550/arxiv.2406.10774","metadata_source":"pith","pith_arxiv_id":"2406.10774","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Quest: Query-Aware Sparsity for Efficient Long-Context LLM Inference","venue":"cs.CL","work_id":"2cad64c9-e2d5-42b6-8db9-03fafde4bcb0","year":2024},"citing_paper":{"arxiv_id":"2605.07363","last_updated":"2026-05-08T07:19:34Z","snapshot_observed_at":"2026-07-06T23:19:46.018236Z","submitted_at":"2026-05-08T07:19:34Z","title":"MISA: Mixture of Indexer Sparse Attention for Long-Context LLM Inference","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-05-11T02:27:55.991919Z"},"links":{"cited_paper":"/paper/2406.10774","citing_paper":"/paper/2605.07363"},"observation_digest":"sha256:f80ace92cd67761cd116319b8e6c080290331fcd9e1b25673bc62868392ac85f","observation_id":"cb396611-294a-4a83-a9a1-aad706f1076a","resolution":{"observed_at":"2026-05-15T14:12:21.110645Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.04617","last_updated":"2024-05-28T12:05:12Z","snapshot_observed_at":"2026-07-06T17:26:39.109878Z","submitted_at":"2024-02-07T06:50:42Z","title":"InfLLM: Training-Free Long-Context Extrapolation for LLMs with an Efficient Context Memory","version":2},"cited_work":{"arxiv_id":"2402.04617","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2402.04617","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Infllm: Unveiling the intrinsic capacity of llms for under- standing extremely long sequences with training-free memory","venue":null,"work_id":"45ccf018-82ed-4c2e-b31c-ae31d8346349","year":2024},"citing_paper":{"arxiv_id":"2605.07363","last_updated":"2026-05-08T07:19:34Z","snapshot_observed_at":"2026-07-06T23:19:46.018236Z","submitted_at":"2026-05-08T07:19:34Z","title":"MISA: Mixture of Indexer Sparse Attention for Long-Context LLM Inference","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-05-11T02:27:55.991919Z"},"links":{"cited_paper":"/paper/2402.04617","citing_paper":"/paper/2605.07363"},"observation_digest":"sha256:99672a9fafefd3c8fc481bca0a962ac7b500f29460787fc1d3e905566930dc12","observation_id":"6833d093-6f8f-40ec-9c39-9bccd792a586","resolution":{"observed_at":"2026-05-11T03:30:56.375873Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Efficient streaming language models with attention sinks","venue":null,"work_id":"70be4e3c-c780-471e-9011-933007101a87","year":null},"citing_paper":{"arxiv_id":"2605.07363","last_updated":"2026-05-08T07:19:34Z","snapshot_observed_at":"2026-07-06T23:19:46.018236Z","submitted_at":"2026-05-08T07:19:34Z","title":"MISA: Mixture of Indexer Sparse Attention for Long-Context LLM Inference","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-05-11T02:27:55.991919Z"},"links":{"citing_paper":"/paper/2605.07363"},"observation_digest":"sha256:639290bf850a46a6942f435cc0d3845bd84530826435fb7c65747a7da17ff043","observation_id":"cb0ba58d-e7dd-40d7-b256-443f559dfb3e","resolution":{"observed_at":"2026-05-14T12:35:15.399691Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2309.17453","last_updated":"2024-04-07T00:56:53Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-09-29T17:59:56Z","title":"Efficient Streaming Language Models with Attention Sinks","version":4},"cited_work":{"arxiv_id":"2309.17453","doi":"10.48550/arxiv.2309.17453","metadata_source":"pith","pith_arxiv_id":"2309.17453","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Efficient Streaming Language Models with Attention Sinks","venue":"cs.CL","work_id":"a8d25452-c237-48c9-88a4-682717c3979a","year":2023},"citing_paper":{"arxiv_id":"2605.07363","last_updated":"2026-05-08T07:19:34Z","snapshot_observed_at":"2026-07-06T23:19:46.018236Z","submitted_at":"2026-05-08T07:19:34Z","title":"MISA: Mixture of Indexer Sparse Attention for Long-Context LLM Inference","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-05-11T02:27:55.991919Z"},"links":{"cited_paper":"/paper/2309.17453","citing_paper":"/paper/2605.07363"},"observation_digest":"sha256:be5c6cd4f9bc1ff642bf84e3804ee5914107f493612bfc925da97013d0e5e33c","observation_id":"0fa2ef93-ac8c-404e-85a5-eb73bc872a92","resolution":{"observed_at":"2026-05-11T03:30:56.468603Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2603.28458","last_updated":"2026-04-06T09:47:34Z","snapshot_observed_at":"2026-08-02T11:27:36.874895Z","submitted_at":"2026-03-30T13:59:51Z","title":"HISA: Efficient Hierarchical Indexing for Fine-Grained Sparse Attention","version":3},"cited_work":{"arxiv_id":"2603.28458","doi":null,"metadata_source":"pith","pith_arxiv_id":"2603.28458","snapshot_observed_at":"2026-07-02T07:56:47.852979Z","title":"HISA: Efficient Hierarchical Indexing for Fine-Grained Sparse Attention","venue":"cs.LG","work_id":"a459e50a-fa2c-4805-82dc-20bd0de8b999","year":2026},"citing_paper":{"arxiv_id":"2605.07363","last_updated":"2026-05-08T07:19:34Z","snapshot_observed_at":"2026-07-06T23:19:46.018236Z","submitted_at":"2026-05-08T07:19:34Z","title":"MISA: Mixture of Indexer Sparse Attention for Long-Context LLM Inference","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-05-11T02:27:55.991919Z"},"links":{"cited_paper":"/paper/2603.28458","citing_paper":"/paper/2605.07363"},"observation_digest":"sha256:d04bd7f71b23f013a9d43c69e17b879de98c384a664f61729d233a961fec0589","observation_id":"268c7aea-e243-4a3f-b23a-3a5d63664b4a","resolution":{"observed_at":"2026-05-11T03:30:56.369152Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.09388","last_updated":"2025-05-14T13:41:34Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-05-14T13:41:34Z","title":"Qwen3 Technical Report","version":1},"cited_work":{"arxiv_id":"2505.09388","doi":"10.1016/j.aiopen.2022.12","metadata_source":"pith","pith_arxiv_id":"2505.09388","snapshot_observed_at":"2026-07-11T11:50:26.030339Z","title":"Qwen3 Technical Report","venue":"cs.CL","work_id":"25a4e30c-1232-48e7-9925-02fa12ba7c9e","year":2025},"citing_paper":{"arxiv_id":"2605.07363","last_updated":"2026-05-08T07:19:34Z","snapshot_observed_at":"2026-07-06T23:19:46.018236Z","submitted_at":"2026-05-08T07:19:34Z","title":"MISA: Mixture of Indexer Sparse Attention for Long-Context LLM Inference","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-05-11T02:27:55.991919Z"},"links":{"cited_paper":"/paper/2505.09388","citing_paper":"/paper/2605.07363"},"observation_digest":"sha256:81cc9b503e9212b4015cf24b52737707d69b327f704800f5c656638f16e0d8fe","observation_id":"b8e3921d-dff3-4abe-a338-4610f3da69cc","resolution":{"observed_at":"2026-05-11T03:30:56.436638Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.11089","last_updated":"2025-02-27T09:01:21Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-02-16T11:53:44Z","title":"Native Sparse Attention: Hardware-Aligned and Natively Trainable Sparse Attention","version":2},"cited_work":{"arxiv_id":"2502.11089","doi":null,"metadata_source":"pith","pith_arxiv_id":"2502.11089","snapshot_observed_at":"2026-07-10T01:36:44.142264Z","title":"Native Sparse Attention: Hardware-Aligned and Natively Trainable Sparse Attention","venue":"cs.CL","work_id":"c74d35ec-94d3-48f5-a291-20cfad7c04a3","year":2025},"citing_paper":{"arxiv_id":"2605.07363","last_updated":"2026-05-08T07:19:34Z","snapshot_observed_at":"2026-07-06T23:19:46.018236Z","submitted_at":"2026-05-08T07:19:34Z","title":"MISA: Mixture of Indexer Sparse Attention for Long-Context LLM Inference","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-05-11T02:27:55.991919Z"},"links":{"cited_paper":"/paper/2502.11089","citing_paper":"/paper/2605.07363"},"observation_digest":"sha256:cd4090cc321a5e3e0f466e4e4a4128cba448560eaf09cfd6c1ab86b61aa03508","observation_id":"e76698ae-4700-42e5-a47c-e70ba82b6cdf","resolution":{"observed_at":"2026-05-16T23:46:30.294205Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Big bird: Transformers for longer sequences","venue":null,"work_id":"4369d25c-a9c0-4d5d-9ebf-8b18ef6e872e","year":null},"citing_paper":{"arxiv_id":"2605.07363","last_updated":"2026-05-08T07:19:34Z","snapshot_observed_at":"2026-07-06T23:19:46.018236Z","submitted_at":"2026-05-08T07:19:34Z","title":"MISA: Mixture of Indexer Sparse Attention for Long-Context LLM Inference","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-05-11T02:27:55.991919Z"},"links":{"citing_paper":"/paper/2605.07363"},"observation_digest":"sha256:5300e6351c72627b090a272cce4e17176dddb7b49f4f22236de10090b8f588a2","observation_id":"127cb2ff-e5ae-4b35-9284-e3e57ad64398","resolution":{"observed_at":"2026-05-14T12:35:15.387596Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2007.14062","last_updated":"2021-01-08T07:41:50Z","snapshot_observed_at":"2026-07-06T09:42:23.296738Z","submitted_at":"2020-07-28T08:34:04Z","title":"Big Bird: Transformers for Longer Sequences","version":2},"cited_work":{"arxiv_id":"2007.14062","doi":"10.48550/arxiv.2007.14062","metadata_source":"pith","pith_arxiv_id":"2007.14062","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Big Bird: Transformers for Longer Sequences","venue":"cs.LG","work_id":"605bd800-a1a3-4bcc-b188-604145af1773","year":2020},"citing_paper":{"arxiv_id":"2605.07363","last_updated":"2026-05-08T07:19:34Z","snapshot_observed_at":"2026-07-06T23:19:46.018236Z","submitted_at":"2026-05-08T07:19:34Z","title":"MISA: Mixture of Indexer Sparse Attention for Long-Context LLM Inference","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-05-11T02:27:55.991919Z"},"links":{"cited_paper":"/paper/2007.14062","citing_paper":"/paper/2605.07363"},"observation_digest":"sha256:a163dce21f4da2de1da30cb831a00edf97050050850b4fd5d146a9d361581aea","observation_id":"4d2c8331-d083-41d8-9e09-46961bc62ef4","resolution":{"observed_at":"2026-05-17T01:54:01.139247Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-07-12T21:50:04.522608+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-12T21:50:04.522608+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2602.15763","last_updated":"2026-02-24T10:44:44Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2026-02-17T17:50:56Z","title":"GLM-5: from Vibe Coding to Agentic Engineering","version":2},"cited_work":{"arxiv_id":"2602.15763","doi":"10.48550/arxiv.2602.15763","metadata_source":"pith","pith_arxiv_id":"2602.15763","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"GLM-5: from Vibe Coding to Agentic Engineering","venue":"cs.LG","work_id":"ad29b1a2-bf77-46b3-9ead-fb62b1d2c6fe","year":2026},"citing_paper":{"arxiv_id":"2605.07363","last_updated":"2026-05-08T07:19:34Z","snapshot_observed_at":"2026-07-06T23:19:46.018236Z","submitted_at":"2026-05-08T07:19:34Z","title":"MISA: Mixture of Indexer Sparse Attention for Long-Context LLM Inference","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-05-11T02:27:55.991919Z"},"links":{"cited_paper":"/paper/2602.15763","citing_paper":"/paper/2605.07363"},"observation_digest":"sha256:49c3c82508eda898262cc83678169312958bd7e1cc5377b351fc70ca2fac67ed","observation_id":"20b0fdd4-87e8-463d-a5aa-c36474994de4","resolution":{"observed_at":"2026-05-11T05:46:41.418110Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2210.05144","last_updated":"2022-10-11T04:54:05Z","snapshot_observed_at":"2026-07-06T14:03:25.739373Z","submitted_at":"2022-10-11T04:54:05Z","title":"Mixture of Attention Heads: Selecting Attention Heads Per Token","version":1},"cited_work":{"arxiv_id":"2210.05144","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2210.05144","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Mixture of attention heads: Selecting attention heads per token","venue":null,"work_id":"da26b16b-0e32-4d2f-9972-8676585f3c3d","year":2022},"citing_paper":{"arxiv_id":"2605.07363","last_updated":"2026-05-08T07:19:34Z","snapshot_observed_at":"2026-07-06T23:19:46.018236Z","submitted_at":"2026-05-08T07:19:34Z","title":"MISA: Mixture of Indexer Sparse Attention for Long-Context LLM Inference","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-05-11T02:27:55.991919Z"},"links":{"cited_paper":"/paper/2210.05144","citing_paper":"/paper/2605.07363"},"observation_digest":"sha256:0cb69acea9aaeabafdf865a7591be789085be591c44389c1d97f982ae8591c03","observation_id":"899edbde-80e4-45f3-8cb6-41f16005532d","resolution":{"observed_at":"2026-05-11T03:30:56.432036Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.14048","last_updated":"2023-12-18T19:10:00Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-06-24T20:11:14Z","title":"H$_2$O: Heavy-Hitter Oracle for Efficient Generative Inference of Large Language Models","version":3},"cited_work":{"arxiv_id":"2306.14048","doi":"10.48550/arxiv.2306.14048","metadata_source":"pith","pith_arxiv_id":"2306.14048","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"H$_2$O: Heavy-Hitter Oracle for Efficient Generative Inference of Large Language Models","venue":"cs.LG","work_id":"1ee58eaa-8aa0-4dec-8e67-6e079ad5fc13","year":2023},"citing_paper":{"arxiv_id":"2605.07363","last_updated":"2026-05-08T07:19:34Z","snapshot_observed_at":"2026-07-06T23:19:46.018236Z","submitted_at":"2026-05-08T07:19:34Z","title":"MISA: Mixture of Indexer Sparse Attention for Long-Context LLM Inference","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-05-11T02:27:55.991919Z"},"links":{"cited_paper":"/paper/2306.14048","citing_paper":"/paper/2605.07363"},"observation_digest":"sha256:5398ab9d0cfc2f32cd1f024ee2c3dbd8ae428758f5a289cd64b7427284e34afe","observation_id":"d1d313f6-5cca-4500-9a0b-e416e8a8d1de","resolution":{"observed_at":"2026-05-17T18:00:50.699870Z","resolver_source":"arxiv_id","status":"malformed_identifier"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Guidelines: • The answer [N/A] means that the paper does not involve crowdsourcing nor research with human subjects","venue":null,"work_id":"b983599a-d4ec-486c-af56-25267266a520","year":null},"citing_paper":{"arxiv_id":"2605.07363","last_updated":"2026-05-08T07:19:34Z","snapshot_observed_at":"2026-07-06T23:19:46.018236Z","submitted_at":"2026-05-08T07:19:34Z","title":"MISA: Mixture of Indexer Sparse Attention for Long-Context LLM Inference","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-05-11T02:27:55.991919Z"},"links":{"citing_paper":"/paper/2605.07363"},"observation_digest":"sha256:b361f6e9876628a484447ac481d782ebf862e67b1d80a5653d176c8093ea20f4","observation_id":"225359fc-4b72-4096-84cd-7b627b5aa297","resolution":{"observed_at":"2026-05-14T12:35:15.406628Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2605.07363","last_updated":"2026-05-08T07:19:34Z","latest_version":1,"primary_category":"cs.LG","snapshot_observed_at":"2026-07-06T23:19:46.018236Z","submitted_at":"2026-05-08T07:19:34Z","title":"MISA: Mixture of Indexer Sparse Attention for Long-Context LLM Inference"},"reference_resolution":{"displayed":33,"state_counts":{"malformed_identifier":1,"metadata_mismatch":6,"parse_uncertain":0,"unresolved":0,"verified_exact":18,"verified_fuzzy":8},"total_outbound_references":33},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"thesis":"As of 5 August 2026, this Paper Citation Record lists 33 of 33 outbound references and 2 inbound Pith citation observations for arXiv:2605.07363."}