{"as_of":"2026-08-05T12:27:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:4b16bcde1f4eba0779660df2689186004c33856218716ab4ca4b6c43f2937ba0","coverage":[{"denominator":35,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":35,"source":"paper_references, paper_reference_links","source_observed_at":"2026-07-31T11:04:04.591687Z","state":"measured"},{"denominator":35,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":35,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-05T06:32:48.257954+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2607.24593/citation-record","integrity":"/paper/2607.24593/integrity","json":"/paper/2607.24593/citation-record.json","paper":"/paper/2607.24593"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2512.02556","last_updated":"2025-12-02T09:25:14Z","snapshot_observed_at":"2026-07-31T23:49:25.878472Z","submitted_at":"2025-12-02T09:25:14Z","title":"DeepSeek-V3.2: Pushing the Frontier of Open Large Language Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2512.02556","snapshot_observed_at":"2026-07-31T11:04:04.513448Z","title":"arXiv preprint arXiv:2512.02556 , year =","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.24593","last_updated":"2026-07-27T15:58:07Z","snapshot_observed_at":"2026-08-04T16:51:01.691381Z","submitted_at":"2026-07-27T15:58:07Z","title":"PIVOT: Efficient Query-Group Indexing for Token-Level Sparse Attention","version":1},"reference_index":1,"source":"arxiv_source","source_observed_at":"2026-07-31T11:04:04.513448Z"},"links":{"cited_paper":"/paper/2512.02556","citing_paper":"/paper/2607.24593"},"observation_digest":"sha256:a26ea860d45d54f966be7ee99752de404c1716f5980b476175168f17420da2c2","observation_id":"f5b24f9d-cac6-4615-9f85-94c6b31ec424","resolution":{"observed_at":"2026-07-31T11:04:04.513448Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2603.28458","last_updated":"2026-04-06T09:47:34Z","snapshot_observed_at":"2026-08-02T11:27:36.874895Z","submitted_at":"2026-03-30T13:59:51Z","title":"HISA: Efficient Hierarchical Indexing for Fine-Grained Sparse Attention","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2603.28458","snapshot_observed_at":"2026-07-31T11:04:04.517513Z","title":"arXiv preprint arXiv:2603.28458 , year =","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.24593","last_updated":"2026-07-27T15:58:07Z","snapshot_observed_at":"2026-08-04T16:51:01.691381Z","submitted_at":"2026-07-27T15:58:07Z","title":"PIVOT: Efficient Query-Group Indexing for Token-Level Sparse Attention","version":1},"reference_index":2,"source":"arxiv_source","source_observed_at":"2026-07-31T11:04:04.517513Z"},"links":{"cited_paper":"/paper/2603.28458","citing_paper":"/paper/2607.24593"},"observation_digest":"sha256:a7da63f6ee567f77d447b665871cc21fabeb549581cf71da8e4736b198433d51","observation_id":"d697f38f-c1f0-4ce2-adf2-895147bca37b","resolution":{"observed_at":"2026-07-31T11:04:04.517513Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T11:04:04.520125Z","title":"arXiv preprint , year =","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.24593","last_updated":"2026-07-27T15:58:07Z","snapshot_observed_at":"2026-08-04T16:51:01.691381Z","submitted_at":"2026-07-27T15:58:07Z","title":"PIVOT: Efficient Query-Group Indexing for Token-Level Sparse Attention","version":1},"reference_index":3,"source":"arxiv_source","source_observed_at":"2026-07-31T11:04:04.520125Z"},"links":{"citing_paper":"/paper/2607.24593"},"observation_digest":"sha256:daefb0c09685d87404cfa85283327e21e7df17cafca3eaf44ca8c4dd7a5f249d","observation_id":"e001cbac-c7c8-49de-813f-5d1f8739b938","resolution":{"observed_at":"2026-07-31T11:04:04.520125Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T11:04:04.522372Z","title":"arXiv preprint arXiv:2603.12201 , year =","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.24593","last_updated":"2026-07-27T15:58:07Z","snapshot_observed_at":"2026-08-04T16:51:01.691381Z","submitted_at":"2026-07-27T15:58:07Z","title":"PIVOT: Efficient Query-Group Indexing for Token-Level Sparse Attention","version":1},"reference_index":4,"source":"arxiv_source","source_observed_at":"2026-07-31T11:04:04.522372Z"},"links":{"citing_paper":"/paper/2607.24593"},"observation_digest":"sha256:9c257a370f4d17208c5f3a43feb116ce2d9eec90bd35c3da21b235628d0f8553","observation_id":"5850cd51-bd3c-4300-8a02-6fa625e34ac2","resolution":{"observed_at":"2026-07-31T11:04:04.522372Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T11:04:04.524693Z","title":"Advances in Neural Information Processing Systems , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.24593","last_updated":"2026-07-27T15:58:07Z","snapshot_observed_at":"2026-08-04T16:51:01.691381Z","submitted_at":"2026-07-27T15:58:07Z","title":"PIVOT: Efficient Query-Group Indexing for Token-Level Sparse Attention","version":1},"reference_index":5,"source":"arxiv_source","source_observed_at":"2026-07-31T11:04:04.524693Z"},"links":{"citing_paper":"/paper/2607.24593"},"observation_digest":"sha256:dcdaa353ad904a57a764583b3cd7021cb2e44c21b907e98eb978e54da07eb589","observation_id":"11414566-a211-4752-acaa-656bd0a0f56b","resolution":{"observed_at":"2026-07-31T11:04:04.524693Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.20766","last_updated":"2025-02-28T06:34:53Z","snapshot_observed_at":"2026-07-06T20:44:13.353359Z","submitted_at":"2025-02-28T06:34:53Z","title":"FlexPrefill: A Context-Aware Sparse Attention Mechanism for Efficient Long-Sequence Inference","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.20766","snapshot_observed_at":"2026-07-31T11:04:04.526969Z","title":"arXiv preprint arXiv:2502.20766 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.24593","last_updated":"2026-07-27T15:58:07Z","snapshot_observed_at":"2026-08-04T16:51:01.691381Z","submitted_at":"2026-07-27T15:58:07Z","title":"PIVOT: Efficient Query-Group Indexing for Token-Level Sparse Attention","version":1},"reference_index":6,"source":"arxiv_source","source_observed_at":"2026-07-31T11:04:04.526969Z"},"links":{"cited_paper":"/paper/2502.20766","citing_paper":"/paper/2607.24593"},"observation_digest":"sha256:77c872a8904d41ca9035cfc717625b23d403a7d17830fed64ec109c8c1cd0e2c","observation_id":"615534be-7122-4418-8206-4f466bc122a5","resolution":{"observed_at":"2026-07-31T11:04:04.526969Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2503.16428","last_updated":"2025-03-20T17:59:58Z","snapshot_observed_at":"2026-07-06T20:56:14.017923Z","submitted_at":"2025-03-20T17:59:58Z","title":"XAttention: Block Sparse Attention with Antidiagonal Scoring","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.16428","snapshot_observed_at":"2026-07-31T11:04:04.529750Z","title":"arXiv preprint arXiv:2503.16428 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.24593","last_updated":"2026-07-27T15:58:07Z","snapshot_observed_at":"2026-08-04T16:51:01.691381Z","submitted_at":"2026-07-27T15:58:07Z","title":"PIVOT: Efficient Query-Group Indexing for Token-Level Sparse Attention","version":1},"reference_index":7,"source":"arxiv_source","source_observed_at":"2026-07-31T11:04:04.529750Z"},"links":{"cited_paper":"/paper/2503.16428","citing_paper":"/paper/2607.24593"},"observation_digest":"sha256:900eb3d2e18135f128a04c3d5c73237ebc7fbbd4926f29390a06d20806847889","observation_id":"51fcd71d-befb-4202-ac04-25aed211ea33","resolution":{"observed_at":"2026-07-31T11:04:04.529750Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2502.13189","last_updated":"2025-02-18T14:06:05Z","snapshot_observed_at":"2026-07-06T20:38:48.725605Z","submitted_at":"2025-02-18T14:06:05Z","title":"MoBA: Mixture of Block Attention for Long-Context LLMs","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2502.13189","snapshot_observed_at":"2026-07-31T11:04:04.532117Z","title":"URL https://arxiv","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.24593","last_updated":"2026-07-27T15:58:07Z","snapshot_observed_at":"2026-08-04T16:51:01.691381Z","submitted_at":"2026-07-27T15:58:07Z","title":"PIVOT: Efficient Query-Group Indexing for Token-Level Sparse Attention","version":1},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-07-31T11:04:04.532117Z"},"links":{"cited_paper":"/paper/2502.13189","citing_paper":"/paper/2607.24593"},"observation_digest":"sha256:2081a92dbfe269ed2b5d766680d83ab5ecdcbce33c58fa29e96881b366ad0a48","observation_id":"59aba606-4570-4574-8eb9-a9b4061065c8","resolution":{"observed_at":"2026-07-31T11:04:04.532117Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T11:04:04.534536Z","title":"Proceedings of the 63rd Annual Meeting of the Association for Computational Linguistics (Volume 1: Long Papers) , pages=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.24593","last_updated":"2026-07-27T15:58:07Z","snapshot_observed_at":"2026-08-04T16:51:01.691381Z","submitted_at":"2026-07-27T15:58:07Z","title":"PIVOT: Efficient Query-Group Indexing for Token-Level Sparse Attention","version":1},"reference_index":9,"source":"arxiv_source","source_observed_at":"2026-07-31T11:04:04.534536Z"},"links":{"citing_paper":"/paper/2607.24593"},"observation_digest":"sha256:4c987fb1178b63d36985135624dcc714f7277169588cc9adc45199adb93baff6","observation_id":"838a464e-1e50-4607-8d35-f5a80d2fa7b8","resolution":{"observed_at":"2026-07-31T11:04:04.534536Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T11:04:04.536683Z","title":"arXiv preprint arXiv:2603.06274 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.24593","last_updated":"2026-07-27T15:58:07Z","snapshot_observed_at":"2026-08-04T16:51:01.691381Z","submitted_at":"2026-07-27T15:58:07Z","title":"PIVOT: Efficient Query-Group Indexing for Token-Level Sparse Attention","version":1},"reference_index":10,"source":"arxiv_source","source_observed_at":"2026-07-31T11:04:04.536683Z"},"links":{"citing_paper":"/paper/2607.24593"},"observation_digest":"sha256:19a0cef67b8144beb9a839de1066b53a9e7885dcf2b6384b313b6c2dde3be028","observation_id":"1f3360d4-b79e-4188-8e35-ceb5c3872d29","resolution":{"observed_at":"2026-07-31T11:04:04.536683Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.14057","last_updated":"2024-07-19T06:34:45Z","snapshot_observed_at":"2026-07-06T18:48:49.546589Z","submitted_at":"2024-07-19T06:34:45Z","title":"LazyLLM: Dynamic Token Pruning for Efficient Long Context LLM Inference","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.14057","snapshot_observed_at":"2026-07-31T11:04:04.538729Z","title":"arXiv preprint arXiv:2407.14057 , year =","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.24593","last_updated":"2026-07-27T15:58:07Z","snapshot_observed_at":"2026-08-04T16:51:01.691381Z","submitted_at":"2026-07-27T15:58:07Z","title":"PIVOT: Efficient Query-Group Indexing for Token-Level Sparse Attention","version":1},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-07-31T11:04:04.538729Z"},"links":{"cited_paper":"/paper/2407.14057","citing_paper":"/paper/2607.24593"},"observation_digest":"sha256:5292cac18a492199e489dec582cb407633b3471fa6ac3bc5add4e1c4617ad14a","observation_id":"ec0f7d09-cb78-4a40-a4ec-8285c62a8ff6","resolution":{"observed_at":"2026-07-31T11:04:04.538729Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T11:04:04.541061Z","title":"Advances in Neural Information Processing Systems , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.24593","last_updated":"2026-07-27T15:58:07Z","snapshot_observed_at":"2026-08-04T16:51:01.691381Z","submitted_at":"2026-07-27T15:58:07Z","title":"PIVOT: Efficient Query-Group Indexing for Token-Level Sparse Attention","version":1},"reference_index":12,"source":"arxiv_source","source_observed_at":"2026-07-31T11:04:04.541061Z"},"links":{"citing_paper":"/paper/2607.24593"},"observation_digest":"sha256:0283e1999175a5210e07839dd25e4a43bc1eb20301fca41b73427e79c869cb2d","observation_id":"60aa39bc-da1b-49f3-bad3-f940605cb2e3","resolution":{"observed_at":"2026-07-31T11:04:04.541061Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T11:04:04.543095Z","title":"Proceedings of the 2024 Conference on Empirical Methods in Natural Language Processing , pages=","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.24593","last_updated":"2026-07-27T15:58:07Z","snapshot_observed_at":"2026-08-04T16:51:01.691381Z","submitted_at":"2026-07-27T15:58:07Z","title":"PIVOT: Efficient Query-Group Indexing for Token-Level Sparse Attention","version":1},"reference_index":13,"source":"arxiv_source","source_observed_at":"2026-07-31T11:04:04.543095Z"},"links":{"citing_paper":"/paper/2607.24593"},"observation_digest":"sha256:09cd16484c39addf03ae23675276108c661cfe66e2f7c452267f133f7ff55543","observation_id":"bd86e0ae-07d9-4d56-941c-d3917dae9a68","resolution":{"observed_at":"2026-07-31T11:04:04.543095Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T11:04:04.546121Z","title":"2024 , eprint=","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.24593","last_updated":"2026-07-27T15:58:07Z","snapshot_observed_at":"2026-08-04T16:51:01.691381Z","submitted_at":"2026-07-27T15:58:07Z","title":"PIVOT: Efficient Query-Group Indexing for Token-Level Sparse Attention","version":1},"reference_index":14,"source":"arxiv_source","source_observed_at":"2026-07-31T11:04:04.546121Z"},"links":{"citing_paper":"/paper/2607.24593"},"observation_digest":"sha256:489093de79d0ac492965ede159302ace6c9e7cd6a262e0a1d37ffc8b6a0316f0","observation_id":"fc6ae7ef-02eb-440a-8d69-f720f25b3684","resolution":{"observed_at":"2026-07-31T11:04:04.546121Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2602.15763","last_updated":"2026-02-24T10:44:44Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2026-02-17T17:50:56Z","title":"GLM-5: from Vibe Coding to Agentic Engineering","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2602.15763","snapshot_observed_at":"2026-07-31T11:04:04.548222Z","title":"arXiv preprint arXiv:2602.15763 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.24593","last_updated":"2026-07-27T15:58:07Z","snapshot_observed_at":"2026-08-04T16:51:01.691381Z","submitted_at":"2026-07-27T15:58:07Z","title":"PIVOT: Efficient Query-Group Indexing for Token-Level Sparse Attention","version":1},"reference_index":15,"source":"arxiv_source","source_observed_at":"2026-07-31T11:04:04.548222Z"},"links":{"cited_paper":"/paper/2602.15763","citing_paper":"/paper/2607.24593"},"observation_digest":"sha256:acfbc22ae250c093d30e928c62ce4617ce75173bd510006d04abafb8894e03e5","observation_id":"529fc8b3-111b-436a-8bbe-32566a58c635","resolution":{"observed_at":"2026-07-31T11:04:04.548222Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T11:04:04.550453Z","title":"Proceedings of the 62nd annual meeting of the association for computational linguistics (volume 1: Long papers) , pages=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.24593","last_updated":"2026-07-27T15:58:07Z","snapshot_observed_at":"2026-08-04T16:51:01.691381Z","submitted_at":"2026-07-27T15:58:07Z","title":"PIVOT: Efficient Query-Group Indexing for Token-Level Sparse Attention","version":1},"reference_index":16,"source":"arxiv_source","source_observed_at":"2026-07-31T11:04:04.550453Z"},"links":{"citing_paper":"/paper/2607.24593"},"observation_digest":"sha256:3869d735e480c63e8445e1055a2171a463ed8c1d122a2cb2bbde2759aa9873b0","observation_id":"6fa48edc-3045-492d-bd38-e27021e1eafe","resolution":{"observed_at":"2026-07-31T11:04:04.550453Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T11:04:04.552295Z","title":"Annual Meeting of the Association for Computational Linguistics (ACL) , year =","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.24593","last_updated":"2026-07-27T15:58:07Z","snapshot_observed_at":"2026-08-04T16:51:01.691381Z","submitted_at":"2026-07-27T15:58:07Z","title":"PIVOT: Efficient Query-Group Indexing for Token-Level Sparse Attention","version":1},"reference_index":17,"source":"arxiv_source","source_observed_at":"2026-07-31T11:04:04.552295Z"},"links":{"citing_paper":"/paper/2607.24593"},"observation_digest":"sha256:42e52b0592a889f0cdbf91e9516b5af59e1aa3fb6d2e94c86e5ba803497e8980","observation_id":"fd82159f-171e-4464-ba0a-b3850e98cb01","resolution":{"observed_at":"2026-07-31T11:04:04.552295Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2404.06654","last_updated":"2024-08-06T21:48:58Z","snapshot_observed_at":"2026-07-06T17:58:00.820879Z","submitted_at":"2024-04-09T23:41:27Z","title":"RULER: What's the Real Context Size of Your Long-Context Language Models?","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.06654","snapshot_observed_at":"2026-07-31T11:04:04.554289Z","title":"arXiv preprint arXiv:2404.06654 , year =","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.24593","last_updated":"2026-07-27T15:58:07Z","snapshot_observed_at":"2026-08-04T16:51:01.691381Z","submitted_at":"2026-07-27T15:58:07Z","title":"PIVOT: Efficient Query-Group Indexing for Token-Level Sparse Attention","version":1},"reference_index":18,"source":"arxiv_source","source_observed_at":"2026-07-31T11:04:04.554289Z"},"links":{"cited_paper":"/paper/2404.06654","citing_paper":"/paper/2607.24593"},"observation_digest":"sha256:180accb8f1b2703dc6f9d745611d72537078e9579b07238de6d96151d962679f","observation_id":"a33ac9d3-315e-48df-8ea6-70c3e551d9f0","resolution":{"observed_at":"2026-07-31T11:04:04.554289Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T11:04:04.556473Z","title":"Proceedings of the 29th Symposium on Operating Systems Principles (SOSP) , year =","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.24593","last_updated":"2026-07-27T15:58:07Z","snapshot_observed_at":"2026-08-04T16:51:01.691381Z","submitted_at":"2026-07-27T15:58:07Z","title":"PIVOT: Efficient Query-Group Indexing for Token-Level Sparse Attention","version":1},"reference_index":19,"source":"arxiv_source","source_observed_at":"2026-07-31T11:04:04.556473Z"},"links":{"citing_paper":"/paper/2607.24593"},"observation_digest":"sha256:702c9fb91e8b16aaa41419cc2f06854adba0e45b2a8980596994e2b0102ec01e","observation_id":"8722fe95-a255-41d6-9ba5-cc823a669c2d","resolution":{"observed_at":"2026-07-31T11:04:04.556473Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T11:04:04.558523Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.24593","last_updated":"2026-07-27T15:58:07Z","snapshot_observed_at":"2026-08-04T16:51:01.691381Z","submitted_at":"2026-07-27T15:58:07Z","title":"PIVOT: Efficient Query-Group Indexing for Token-Level Sparse Attention","version":1},"reference_index":20,"source":"arxiv_source","source_observed_at":"2026-07-31T11:04:04.558523Z"},"links":{"citing_paper":"/paper/2607.24593"},"observation_digest":"sha256:d2b5dc65cd6404e6923b29945a38d87ccbc603e435cfb4d49de0aa68d83d8f90","observation_id":"e17c3b08-9cde-4284-bb7a-3f2e8dd642e0","resolution":{"observed_at":"2026-07-31T11:04:04.558523Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T11:04:04.560504Z","title":"int8 (): 8-bit matrix multiplication for transformers at scale , author=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.24593","last_updated":"2026-07-27T15:58:07Z","snapshot_observed_at":"2026-08-04T16:51:01.691381Z","submitted_at":"2026-07-27T15:58:07Z","title":"PIVOT: Efficient Query-Group Indexing for Token-Level Sparse Attention","version":1},"reference_index":21,"source":"arxiv_source","source_observed_at":"2026-07-31T11:04:04.560504Z"},"links":{"citing_paper":"/paper/2607.24593"},"observation_digest":"sha256:fb2181da8cb2e629fef3200d844ace8de0660cd9e8f2ffff7b379e9a50879d77","observation_id":"ea4defef-a86c-432e-8ace-22542cefc9dc","resolution":{"observed_at":"2026-07-31T11:04:04.560504Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2307.09288","last_updated":"2023-07-19T17:08:59Z","snapshot_observed_at":"2026-08-02T11:57:18.735747Z","submitted_at":"2023-07-18T14:31:57Z","title":"Llama 2: Open Foundation and Fine-Tuned Chat Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2307.09288","snapshot_observed_at":"2026-07-31T11:04:04.562507Z","title":"arXiv preprint arXiv:2307.09288 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.24593","last_updated":"2026-07-27T15:58:07Z","snapshot_observed_at":"2026-08-04T16:51:01.691381Z","submitted_at":"2026-07-27T15:58:07Z","title":"PIVOT: Efficient Query-Group Indexing for Token-Level Sparse Attention","version":1},"reference_index":22,"source":"arxiv_source","source_observed_at":"2026-07-31T11:04:04.562507Z"},"links":{"cited_paper":"/paper/2307.09288","citing_paper":"/paper/2607.24593"},"observation_digest":"sha256:e7cd72fc26c6e983a53b1fbe98c3901dd302c571e47b67e4d34f688a38ec79da","observation_id":"62bead20-80b3-4e12-a418-4b7253b61019","resolution":{"observed_at":"2026-07-31T11:04:04.562507Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2505.09388","last_updated":"2025-05-14T13:41:34Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-05-14T13:41:34Z","title":"Qwen3 Technical Report","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2505.09388","snapshot_observed_at":"2026-07-31T11:04:04.564666Z","title":"arXiv preprint arXiv:2505.09388 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.24593","last_updated":"2026-07-27T15:58:07Z","snapshot_observed_at":"2026-08-04T16:51:01.691381Z","submitted_at":"2026-07-27T15:58:07Z","title":"PIVOT: Efficient Query-Group Indexing for Token-Level Sparse Attention","version":1},"reference_index":23,"source":"arxiv_source","source_observed_at":"2026-07-31T11:04:04.564666Z"},"links":{"cited_paper":"/paper/2505.09388","citing_paper":"/paper/2607.24593"},"observation_digest":"sha256:286fdba5a804929d5003f9462878768261c4ee4cf00a882cfa710fac1e957f4b","observation_id":"da9c64c8-a41a-4133-8b40-a97e5539efca","resolution":{"observed_at":"2026-07-31T11:04:04.564666Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T11:04:04.566891Z","title":"International conference on algorithmic learning theory , pages=","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2607.24593","last_updated":"2026-07-27T15:58:07Z","snapshot_observed_at":"2026-08-04T16:51:01.691381Z","submitted_at":"2026-07-27T15:58:07Z","title":"PIVOT: Efficient Query-Group Indexing for Token-Level Sparse Attention","version":1},"reference_index":24,"source":"arxiv_source","source_observed_at":"2026-07-31T11:04:04.566891Z"},"links":{"citing_paper":"/paper/2607.24593"},"observation_digest":"sha256:0afb8333ba61694e533c1523831a11e08595fe74314103947efc80c1439aa224","observation_id":"1faed358-89f8-4158-9fd6-86fc405e42b1","resolution":{"observed_at":"2026-07-31T11:04:04.566891Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2507.20534","last_updated":"2026-02-03T04:57:00Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-07-28T05:35:43Z","title":"Kimi K2: Open Agentic Intelligence","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2507.20534","snapshot_observed_at":"2026-07-31T11:04:04.569004Z","title":"arXiv preprint arXiv:2507.20534 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.24593","last_updated":"2026-07-27T15:58:07Z","snapshot_observed_at":"2026-08-04T16:51:01.691381Z","submitted_at":"2026-07-27T15:58:07Z","title":"PIVOT: Efficient Query-Group Indexing for Token-Level Sparse Attention","version":1},"reference_index":25,"source":"arxiv_source","source_observed_at":"2026-07-31T11:04:04.569004Z"},"links":{"cited_paper":"/paper/2507.20534","citing_paper":"/paper/2607.24593"},"observation_digest":"sha256:ae84058e0b8c5f8c70f8701080f0782b91ea23abdbf1f58a5e22eb2ccf692b0f","observation_id":"68b4bd60-4ad1-41d1-b6af-632a4bbdadf5","resolution":{"observed_at":"2026-07-31T11:04:04.569004Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T11:04:04.571340Z","title":"arXiv preprint arXiv:2509.24663 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.24593","last_updated":"2026-07-27T15:58:07Z","snapshot_observed_at":"2026-08-04T16:51:01.691381Z","submitted_at":"2026-07-27T15:58:07Z","title":"PIVOT: Efficient Query-Group Indexing for Token-Level Sparse Attention","version":1},"reference_index":26,"source":"arxiv_source","source_observed_at":"2026-07-31T11:04:04.571340Z"},"links":{"citing_paper":"/paper/2607.24593"},"observation_digest":"sha256:2241e2f9d9166084483bb2e823696b5329690ebb40dde95db7497e8b3744071d","observation_id":"d7cf0e14-d605-43d8-9695-59c2a54a749a","resolution":{"observed_at":"2026-07-31T11:04:04.571340Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T11:04:04.573511Z","title":"arXiv preprint arXiv:2602.03560 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.24593","last_updated":"2026-07-27T15:58:07Z","snapshot_observed_at":"2026-08-04T16:51:01.691381Z","submitted_at":"2026-07-27T15:58:07Z","title":"PIVOT: Efficient Query-Group Indexing for Token-Level Sparse Attention","version":1},"reference_index":27,"source":"arxiv_source","source_observed_at":"2026-07-31T11:04:04.573511Z"},"links":{"citing_paper":"/paper/2607.24593"},"observation_digest":"sha256:557ebe930fa2b65d524aac5e0fe1e27124b55ff5d993223204bdd1ba961c866a","observation_id":"47c3b097-9b55-4159-93c0-d87eeee41ec0","resolution":{"observed_at":"2026-07-31T11:04:04.573511Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T11:04:04.575512Z","title":"International Conference on Learning Representations , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.24593","last_updated":"2026-07-27T15:58:07Z","snapshot_observed_at":"2026-08-04T16:51:01.691381Z","submitted_at":"2026-07-27T15:58:07Z","title":"PIVOT: Efficient Query-Group Indexing for Token-Level Sparse Attention","version":1},"reference_index":28,"source":"arxiv_source","source_observed_at":"2026-07-31T11:04:04.575512Z"},"links":{"citing_paper":"/paper/2607.24593"},"observation_digest":"sha256:6c1cd927f4b01f7d5f40d08b9bb54b769a3ca71ad3c7c6a51cb5cc5a06e1e6c1","observation_id":"4fcbc682-8a47-4590-90bd-2401f0776dae","resolution":{"observed_at":"2026-07-31T11:04:04.575512Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T11:04:04.577595Z","title":"arXiv preprint arXiv:2509.24745 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.24593","last_updated":"2026-07-27T15:58:07Z","snapshot_observed_at":"2026-08-04T16:51:01.691381Z","submitted_at":"2026-07-27T15:58:07Z","title":"PIVOT: Efficient Query-Group Indexing for Token-Level Sparse Attention","version":1},"reference_index":29,"source":"arxiv_source","source_observed_at":"2026-07-31T11:04:04.577595Z"},"links":{"citing_paper":"/paper/2607.24593"},"observation_digest":"sha256:76dbe593bb5c49010406b0f9e3cc282e5b5466a15ed153b70f1e32c9199f55b8","observation_id":"94eebb90-568f-45ec-a046-864ec081f201","resolution":{"observed_at":"2026-07-31T11:04:04.577595Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.13276","last_updated":"2025-02-17T02:24:47Z","snapshot_observed_at":"2026-08-05T05:19:17.395325Z","submitted_at":"2024-10-17T07:07:09Z","title":"SeerAttention: Learning Intrinsic Sparse Attention in Your LLMs","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.13276","snapshot_observed_at":"2026-07-31T11:04:04.579723Z","title":"arXiv preprint arXiv:2410.13276 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.24593","last_updated":"2026-07-27T15:58:07Z","snapshot_observed_at":"2026-08-04T16:51:01.691381Z","submitted_at":"2026-07-27T15:58:07Z","title":"PIVOT: Efficient Query-Group Indexing for Token-Level Sparse Attention","version":1},"reference_index":30,"source":"arxiv_source","source_observed_at":"2026-07-31T11:04:04.579723Z"},"links":{"cited_paper":"/paper/2410.13276","citing_paper":"/paper/2607.24593"},"observation_digest":"sha256:c3bf6f394aa9e91c3636aeb00f5c347c46194c2481f9ec7aec618f7e52149ab6","observation_id":"fb9c76ae-c1f1-4464-b8aa-2b608a2e0421","resolution":{"observed_at":"2026-07-31T11:04:04.579723Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2606.04511","last_updated":"2026-06-03T06:42:05Z","snapshot_observed_at":"2026-07-06T23:44:37.797802Z","submitted_at":"2026-06-03T06:42:05Z","title":"SparDA: Sparse Decoupled Attention for Efficient Long-Context LLM Inference","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2606.04511","snapshot_observed_at":"2026-07-31T11:04:04.582053Z","title":"arXiv preprint arXiv:2606.04511 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.24593","last_updated":"2026-07-27T15:58:07Z","snapshot_observed_at":"2026-08-04T16:51:01.691381Z","submitted_at":"2026-07-27T15:58:07Z","title":"PIVOT: Efficient Query-Group Indexing for Token-Level Sparse Attention","version":1},"reference_index":31,"source":"arxiv_source","source_observed_at":"2026-07-31T11:04:04.582053Z"},"links":{"cited_paper":"/paper/2606.04511","citing_paper":"/paper/2607.24593"},"observation_digest":"sha256:7db50b2b47c24a457ed5aeb58c7bba6e078d60292b79f6677f7cbbe2629f7fdf","observation_id":"8bb9cb9a-5b29-467e-b518-3caf6038c74c","resolution":{"observed_at":"2026-07-31T11:04:04.582053Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2506.08889","last_updated":"2025-06-10T15:17:26Z","snapshot_observed_at":"2026-07-06T21:39:52.201136Z","submitted_at":"2025-06-10T15:17:26Z","title":"SeerAttention-R: Sparse Attention Adaptation for Long Reasoning","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2506.08889","snapshot_observed_at":"2026-07-31T11:04:04.584216Z","title":"arXiv preprint arXiv:2506.08889 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.24593","last_updated":"2026-07-27T15:58:07Z","snapshot_observed_at":"2026-08-04T16:51:01.691381Z","submitted_at":"2026-07-27T15:58:07Z","title":"PIVOT: Efficient Query-Group Indexing for Token-Level Sparse Attention","version":1},"reference_index":32,"source":"arxiv_source","source_observed_at":"2026-07-31T11:04:04.584216Z"},"links":{"cited_paper":"/paper/2506.08889","citing_paper":"/paper/2607.24593"},"observation_digest":"sha256:71c247b8e33ad44d5a7ee51e26567eb143ab460bf6fef63f40fd3f41c734f75e","observation_id":"5ee18d25-c704-4b04-9543-64c33b112150","resolution":{"observed_at":"2026-07-31T11:04:04.584216Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2605.16928","last_updated":"2026-06-08T02:38:05Z","snapshot_observed_at":"2026-08-02T02:41:46.596895Z","submitted_at":"2026-05-16T10:51:58Z","title":"Full Attention Strikes Back: Transferring Full Attention into Sparse within Hundred Training Steps","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2605.16928","snapshot_observed_at":"2026-07-31T11:04:04.586667Z","title":"arXiv preprint arXiv:2605.16928 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.24593","last_updated":"2026-07-27T15:58:07Z","snapshot_observed_at":"2026-08-04T16:51:01.691381Z","submitted_at":"2026-07-27T15:58:07Z","title":"PIVOT: Efficient Query-Group Indexing for Token-Level Sparse Attention","version":1},"reference_index":33,"source":"arxiv_source","source_observed_at":"2026-07-31T11:04:04.586667Z"},"links":{"cited_paper":"/paper/2605.16928","citing_paper":"/paper/2607.24593"},"observation_digest":"sha256:a9c34a933af3cb757c2501932ce2e8ec6c829689fbea510050d7b43359d7d159","observation_id":"ee1f1960-9c4a-4ef6-ba97-63858619af08","resolution":{"observed_at":"2026-07-31T11:04:04.586667Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2505.07608","last_updated":"2025-06-05T11:49:09Z","snapshot_observed_at":"2026-07-06T21:22:35.652600Z","submitted_at":"2025-05-12T14:30:11Z","title":"MiMo: Unlocking the Reasoning Potential of Language Model -- From Pretraining to Posttraining","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2505.07608","snapshot_observed_at":"2026-07-31T11:04:04.589430Z","title":"arXiv preprint arXiv:2505.07608 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.24593","last_updated":"2026-07-27T15:58:07Z","snapshot_observed_at":"2026-08-04T16:51:01.691381Z","submitted_at":"2026-07-27T15:58:07Z","title":"PIVOT: Efficient Query-Group Indexing for Token-Level Sparse Attention","version":1},"reference_index":34,"source":"arxiv_source","source_observed_at":"2026-07-31T11:04:04.589430Z"},"links":{"cited_paper":"/paper/2505.07608","citing_paper":"/paper/2607.24593"},"observation_digest":"sha256:8695245bf22498cf4b7da48e9f114dfcb01fb3d4720bfceb6e55043821731784","observation_id":"284e5753-3ba2-4d82-a1bb-310cb2ee18f0","resolution":{"observed_at":"2026-07-31T11:04:04.589430Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.10774","last_updated":"2024-08-26T21:01:02Z","snapshot_observed_at":"2026-07-06T18:31:36.504766Z","submitted_at":"2024-06-16T01:33:02Z","title":"Quest: Query-Aware Sparsity for Efficient Long-Context LLM Inference","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.10774","snapshot_observed_at":"2026-07-31T11:04:04.591687Z","title":"arXiv preprint arXiv:2406.10774 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.24593","last_updated":"2026-07-27T15:58:07Z","snapshot_observed_at":"2026-08-04T16:51:01.691381Z","submitted_at":"2026-07-27T15:58:07Z","title":"PIVOT: Efficient Query-Group Indexing for Token-Level Sparse Attention","version":1},"reference_index":35,"source":"arxiv_source","source_observed_at":"2026-07-31T11:04:04.591687Z"},"links":{"cited_paper":"/paper/2406.10774","citing_paper":"/paper/2607.24593"},"observation_digest":"sha256:090dc39bd8c821ea6a0f19932a005cf2d3898b1a81aab4897998fcd27ee218ad","observation_id":"eb91c79e-9676-4e67-b63f-7dfaed4ce0fd","resolution":{"observed_at":"2026-07-31T11:04:04.591687Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2607.24593","last_updated":"2026-07-27T15:58:07Z","latest_version":1,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-04T16:51:01.691381Z","submitted_at":"2026-07-27T15:58:07Z","title":"PIVOT: Efficient Query-Group Indexing for Token-Level Sparse Attention"},"reference_resolution":{"displayed":35,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":35,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":35},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"thesis":"As of 5 August 2026, this Paper Citation Record lists 35 of 35 outbound references and 0 inbound Pith citation observations for arXiv:2607.24593."}