{"as_of":"2026-08-22T10:44:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:d5e07b9c550148ec6edde0b64933b1adca4f1780d3f41021bc4c99ca75f43e84","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":33,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":33,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-22T06:32:14.747728+00:00","state":"measured"},{"denominator":33,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":33,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-16T10:28:52.508979Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"pith","source_observed_at":"2026-07-10T01:36:44.234407Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2410.21465","last_updated":"2025-04-25T19:40:54Z","snapshot_observed_at":"2026-08-19T20:01:29.403616Z","submitted_at":"2024-10-28T19:08:12Z","title":"ShadowKV: KV Cache in Shadows for High-Throughput Long-Context LLM Inference","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.21465","snapshot_observed_at":"2026-08-12T11:56:33.082127Z","title":"Shadowkv: Kv cache in shadows for high-throughput long-context llm inference, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2411.17685","last_updated":"2024-11-26T18:52:06Z","snapshot_observed_at":"2026-08-18T06:32:27.100646Z","submitted_at":"2024-11-26T18:52:06Z","title":"Attamba: Attending To Multi-Token States","version":1},"reference_index":17,"source":"arxiv_source","source_observed_at":"2026-08-12T11:56:33.082127Z"},"links":{"cited_paper":"/paper/2410.21465","citing_paper":"/paper/2411.17685"},"observation_digest":"sha256:a2c8f15c1b532f66f0bc6b434b050e47bc75a2037282fece8062d2d272a31df3","observation_id":"d00e5596-cb68-4190-a5e7-2648c9153629","resolution":{"observed_at":"2026-08-12T11:56:33.082127Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.21465","last_updated":"2025-04-25T19:40:54Z","snapshot_observed_at":"2026-08-19T20:01:29.403616Z","submitted_at":"2024-10-28T19:08:12Z","title":"ShadowKV: KV Cache in Shadows for High-Throughput Long-Context LLM Inference","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.21465","snapshot_observed_at":"2026-08-11T16:14:06.050044Z","title":"Shadowkv: Kv cache in shadows for high-throughput long-context llm inference","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2412.10319","last_updated":"2025-03-11T14:02:04Z","snapshot_observed_at":"2026-08-13T16:34:08.378180Z","submitted_at":"2024-12-13T17:59:52Z","title":"SCBench: A KV Cache-Centric Analysis of Long-Context Methods","version":2},"reference_index":77,"source":"arxiv_source","source_observed_at":"2026-08-11T16:14:06.050044Z"},"links":{"cited_paper":"/paper/2410.21465","citing_paper":"/paper/2412.10319"},"observation_digest":"sha256:e637221701396651c272c20a15c8966bbfb667ff1d8114516a021a64ca174ab6","observation_id":"26a87e2a-968c-4121-9028-452a777bf42d","resolution":{"observed_at":"2026-08-11T16:14:06.050044Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.21465","last_updated":"2025-04-25T19:40:54Z","snapshot_observed_at":"2026-08-19T20:01:29.403616Z","submitted_at":"2024-10-28T19:08:12Z","title":"ShadowKV: KV Cache in Shadows for High-Throughput Long-Context LLM Inference","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.21465","snapshot_observed_at":"2026-08-11T00:38:46.981217Z","title":"Shadowkv: Kv cache in shadows for high-throughput long-context llm inference,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.19442","last_updated":"2025-07-30T05:24:46Z","snapshot_observed_at":"2026-08-20T14:01:46.841630Z","submitted_at":"2024-12-27T04:17:57Z","title":"A Survey on Large Language Model Acceleration based on KV Cache Management","version":3},"reference_index":71,"source":"pdf_text","source_observed_at":"2026-08-11T00:38:46.981217Z"},"links":{"cited_paper":"/paper/2410.21465","citing_paper":"/paper/2412.19442"},"observation_digest":"sha256:20acbf334b61fe9b0cd16c40d0f8857e3f2f316547df60ca5f5cb4d051199681","observation_id":"8cb3a882-91f5-4131-810b-d02aac224e29","resolution":{"observed_at":"2026-08-11T00:38:46.981217Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.21465","last_updated":"2025-04-25T19:40:54Z","snapshot_observed_at":"2026-08-19T20:01:29.403616Z","submitted_at":"2024-10-28T19:08:12Z","title":"ShadowKV: KV Cache in Shadows for High-Throughput Long-Context LLM Inference","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.21465","snapshot_observed_at":"2026-08-09T13:26:55.444082Z","title":"Shadowkv: Kv cache in shadows for high-throughput long-context llm inference","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.02617","last_updated":"2025-02-04T08:52:13Z","snapshot_observed_at":"2026-08-15T21:58:34.589761Z","submitted_at":"2025-02-04T08:52:13Z","title":"PolarQuant: Quantizing KV Caches with Polar Transformation","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-09T13:26:55.444082Z"},"links":{"cited_paper":"/paper/2410.21465","citing_paper":"/paper/2502.02617"},"observation_digest":"sha256:fe7b914e80cc11a19c64af687643c99d4621e5cdca9b26f6f504953590ad2ddf","observation_id":"a12d4335-85ec-4560-8a79-12532bb34198","resolution":{"observed_at":"2026-08-09T13:26:55.444082Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.21465","last_updated":"2025-04-25T19:40:54Z","snapshot_observed_at":"2026-08-19T20:01:29.403616Z","submitted_at":"2024-10-28T19:08:12Z","title":"ShadowKV: KV Cache in Shadows for High-Throughput Long-Context LLM Inference","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.21465","snapshot_observed_at":"2026-08-09T11:11:17.779908Z","title":"Shadowkv: Kv cache in shadows for high-throughput long-context llm inference, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.02789","last_updated":"2025-05-19T18:24:50Z","snapshot_observed_at":"2026-08-17T01:39:21.989668Z","submitted_at":"2025-02-05T00:22:06Z","title":"Speculative Prefill: Turbocharging TTFT with Lightweight and Training-Free Token Importance Estimation","version":2},"reference_index":44,"source":"arxiv_source","source_observed_at":"2026-08-09T11:11:17.779908Z"},"links":{"cited_paper":"/paper/2410.21465","citing_paper":"/paper/2502.02789"},"observation_digest":"sha256:af1280f340f025781cad5303a70fae8fb10c1f2edbe46d5691051e1e74360c91","observation_id":"858ec0b8-1b25-42de-a93e-597ba1aa812a","resolution":{"observed_at":"2026-08-09T11:11:17.779908Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.21465","last_updated":"2025-04-25T19:40:54Z","snapshot_observed_at":"2026-08-19T20:01:29.403616Z","submitted_at":"2024-10-28T19:08:12Z","title":"ShadowKV: KV Cache in Shadows for High-Throughput Long-Context LLM Inference","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.21465","snapshot_observed_at":"2026-08-08T13:46:11.424138Z","title":"Shadowkv: Kv cache in shadows for high-throughput long-context llm inference.arXiv preprint arXiv:2410.21465 , 2024a","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2502.09647","last_updated":"2025-03-05T16:14:16Z","snapshot_observed_at":"2026-08-15T21:43:38.379452Z","submitted_at":"2025-02-11T00:04:32Z","title":"Unveiling Simplicities of Attention: Adaptive Long-Context Head Identification","version":2},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-08T13:46:11.424138Z"},"links":{"cited_paper":"/paper/2410.21465","citing_paper":"/paper/2502.09647"},"observation_digest":"sha256:e2776b60d0f614486a8d294bbe586ff5f5a222396e011648c8a15d5fffe209f2","observation_id":"e9cacb04-604f-45c6-97e3-80b3e285eaeb","resolution":{"observed_at":"2026-08-08T13:46:11.424138Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.21465","last_updated":"2025-04-25T19:40:54Z","snapshot_observed_at":"2026-08-19T20:01:29.403616Z","submitted_at":"2024-10-28T19:08:12Z","title":"ShadowKV: KV Cache in Shadows for High-Throughput Long-Context LLM Inference","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.21465","snapshot_observed_at":"2026-08-16T10:28:52.508979Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2504.18154","last_updated":"2025-04-25T08:06:22Z","snapshot_observed_at":"2026-08-18T18:55:10.997004Z","submitted_at":"2025-04-25T08:06:22Z","title":"EcoServe: Enabling Cost-effective LLM Serving with Proactive Intra- and Inter-Instance Orchestration","version":1},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-16T10:28:52.508979Z"},"links":{"cited_paper":"/paper/2410.21465","citing_paper":"/paper/2504.18154"},"observation_digest":"sha256:7b933ab683f50c61b2ca3c504bd372f34add5ba71f8baa0104749ada79d267f4","observation_id":"d19f214f-d719-492e-83ff-27baf0ba7ae9","resolution":{"observed_at":"2026-08-16T10:28:52.508979Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.21465","last_updated":"2025-04-25T19:40:54Z","snapshot_observed_at":"2026-08-19T20:01:29.403616Z","submitted_at":"2024-10-28T19:08:12Z","title":"ShadowKV: KV Cache in Shadows for High-Throughput Long-Context LLM Inference","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.21465","snapshot_observed_at":"2026-08-07T13:32:32.885907Z","title":"Shadowkv: Kv cache in shadows for high-throughput long-context llm inference, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.21487","last_updated":"2025-05-27T17:54:07Z","snapshot_observed_at":"2026-08-07T13:24:37.243429Z","submitted_at":"2025-05-27T17:54:07Z","title":"Hardware-Efficient Attention for Fast Decoding","version":1},"reference_index":64,"source":"arxiv_source","source_observed_at":"2026-08-07T13:32:32.885907Z"},"links":{"cited_paper":"/paper/2410.21465","citing_paper":"/paper/2505.21487"},"observation_digest":"sha256:9acb2fb3672e8b0c4d78bbf637ae046b2c8459a65ba48ef4ddd9746d7cbcb137","observation_id":"2393bfa6-f739-4a74-81c7-0b8bee2b5202","resolution":{"observed_at":"2026-08-07T13:32:32.885907Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.21465","last_updated":"2025-04-25T19:40:54Z","snapshot_observed_at":"2026-08-19T20:01:29.403616Z","submitted_at":"2024-10-28T19:08:12Z","title":"ShadowKV: KV Cache in Shadows for High-Throughput Long-Context LLM Inference","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.21465","snapshot_observed_at":"2026-08-07T11:26:00.895653Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.02572","last_updated":"2025-06-03T07:53:32Z","snapshot_observed_at":"2026-08-22T01:58:19.972521Z","submitted_at":"2025-06-03T07:53:32Z","title":"HATA: Trainable and Hardware-Efficient Hash-Aware Top-k Attention for Scalable Large Model Inference","version":1},"reference_index":28,"source":"arxiv_source","source_observed_at":"2026-08-07T11:26:00.895653Z"},"links":{"cited_paper":"/paper/2410.21465","citing_paper":"/paper/2506.02572"},"observation_digest":"sha256:258502a4e2e2239681870fb3fd7f0f8f7cd6907a5fb7066a9d12a998f743396b","observation_id":"2178ee9d-6c7e-4b91-bb98-be1e9a5b5e6b","resolution":{"observed_at":"2026-08-07T11:26:00.895653Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.21465","last_updated":"2025-04-25T19:40:54Z","snapshot_observed_at":"2026-08-19T20:01:29.403616Z","submitted_at":"2024-10-28T19:08:12Z","title":"ShadowKV: KV Cache in Shadows for High-Throughput Long-Context LLM Inference","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.21465","snapshot_observed_at":"2026-08-07T10:30:34.376950Z","title":"16 Hanshi Sun, Li-Wen Chang, Wenlei Bao, Size Zheng, Ningxin Zheng, Xin Liu, Harry Dong, Yuejie Chi, and Beidi Chen","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2506.05333","last_updated":"2025-06-20T01:25:25Z","snapshot_observed_at":"2026-08-07T10:18:59.977399Z","submitted_at":"2025-06-05T17:59:24Z","title":"Kinetics: Rethinking Test-Time Scaling Laws","version":3},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-08-07T10:30:34.376950Z"},"links":{"cited_paper":"/paper/2410.21465","citing_paper":"/paper/2506.05333"},"observation_digest":"sha256:4fa1992fec9374ac0be172d60ffc5acd4aa49b997776b901bd994db349be87f4","observation_id":"2122038e-f8b7-4b8f-8ae8-c0b6172ef580","resolution":{"observed_at":"2026-08-07T10:30:34.376950Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.21465","last_updated":"2025-04-25T19:40:54Z","snapshot_observed_at":"2026-08-19T20:01:29.403616Z","submitted_at":"2024-10-28T19:08:12Z","title":"ShadowKV: KV Cache in Shadows for High-Throughput Long-Context LLM Inference","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.21465","snapshot_observed_at":"2026-08-07T12:38:39.632271Z","title":"Shadowkv: Kv cache in shadows for high-throughput long-context llm inference","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2506.15704","last_updated":"2025-05-30T02:35:59Z","snapshot_observed_at":"2026-08-17T01:59:46.007512Z","submitted_at":"2025-05-30T02:35:59Z","title":"Learn from the Past: Fast Sparse Indexing for Large Language Model Decoding","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-07T12:38:39.632271Z"},"links":{"cited_paper":"/paper/2410.21465","citing_paper":"/paper/2506.15704"},"observation_digest":"sha256:67963f0cdaa20fccb8e773ac771a5c8a1f3867e9c30c6b91f31b6d9fa5d8a984","observation_id":"9545cbf4-34d8-46f4-9560-9cacafa53065","resolution":{"observed_at":"2026-08-07T12:38:39.632271Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.21465","last_updated":"2025-04-25T19:40:54Z","snapshot_observed_at":"2026-08-19T20:01:29.403616Z","submitted_at":"2024-10-28T19:08:12Z","title":"ShadowKV: KV Cache in Shadows for High-Throughput Long-Context LLM Inference","version":3},"cited_work":{"arxiv_id":"2410.21465","doi":null,"metadata_source":"pith","pith_arxiv_id":"2410.21465","snapshot_observed_at":"2026-07-10T01:36:44.234407Z","title":"Shadowkv: Kv cache in shadows for high-throughput long-context llm inference","venue":"cs.LG","work_id":"23d180f7-15ab-4477-b880-cd63ca2f6a95","year":2024},"citing_paper":{"arxiv_id":"2509.21623","last_updated":"2026-04-16T21:29:54Z","snapshot_observed_at":"2026-08-14T12:21:00.182138Z","submitted_at":"2025-09-25T21:42:27Z","title":"OjaKV: Context-Aware Online Low-Rank KV Cache Compression","version":2},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-05-18T13:26:02.980973Z"},"links":{"cited_paper":"/paper/2410.21465","citing_paper":"/paper/2509.21623"},"observation_digest":"sha256:cc0fc445058e8260874464a50b3435c02bece54f1b070460b19475eb932675e7","observation_id":"30d37a71-f428-4826-b899-af7e042db1c9","resolution":{"observed_at":"2026-05-18T13:26:24.732629Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.21465","last_updated":"2025-04-25T19:40:54Z","snapshot_observed_at":"2026-08-19T20:01:29.403616Z","submitted_at":"2024-10-28T19:08:12Z","title":"ShadowKV: KV Cache in Shadows for High-Throughput Long-Context LLM Inference","version":3},"cited_work":{"arxiv_id":"2410.21465","doi":null,"metadata_source":"pith","pith_arxiv_id":"2410.21465","snapshot_observed_at":"2026-07-10T01:36:44.234407Z","title":"Shadowkv: Kv cache in shadows for high-throughput long-context llm inference","venue":"cs.LG","work_id":"23d180f7-15ab-4477-b880-cd63ca2f6a95","year":2024},"citing_paper":{"arxiv_id":"2601.13684","last_updated":"2026-04-18T08:34:09Z","snapshot_observed_at":"2026-08-17T05:04:07.975490Z","submitted_at":"2026-01-20T07:35:06Z","title":"HeteroCache: A Dynamic Retrieval Approach to Heterogeneous KV Cache Compression for Long-Context LLM Inference","version":2},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-05-16T13:16:31.568604Z"},"links":{"cited_paper":"/paper/2410.21465","citing_paper":"/paper/2601.13684"},"observation_digest":"sha256:1bb1b33a530bfd594bc05380d16685d2683db466a6283fdc582edccdde8e72eb","observation_id":"5ed6dcbf-1129-423e-a47c-1dd72e511631","resolution":{"observed_at":"2026-05-16T13:17:54.885057Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.21465","last_updated":"2025-04-25T19:40:54Z","snapshot_observed_at":"2026-08-19T20:01:29.403616Z","submitted_at":"2024-10-28T19:08:12Z","title":"ShadowKV: KV Cache in Shadows for High-Throughput Long-Context LLM Inference","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.21465","snapshot_observed_at":"2026-08-03T03:37:50.132428Z","title":"Shadowkv: Kv cache in shadows for high-throughput long-context llm inference.arXiv preprint arXiv:2410.21465,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2602.07721","last_updated":"2026-05-28T23:34:23Z","snapshot_observed_at":"2026-08-13T15:58:59.754405Z","submitted_at":"2026-02-07T22:26:45Z","title":"ParisKV: Fast and Drift-Robust KV-Cache Retrieval for Long-Context LLMs","version":3},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-03T03:37:50.132428Z"},"links":{"cited_paper":"/paper/2410.21465","citing_paper":"/paper/2602.07721"},"observation_digest":"sha256:a1ab9198f5700f8ee35e59017a16876e764ad5f755f9165662d162b93de4a943","observation_id":"2c23b3d6-85bb-4a6a-a169-544ac13a77e2","resolution":{"observed_at":"2026-08-03T03:37:50.132428Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.21465","last_updated":"2025-04-25T19:40:54Z","snapshot_observed_at":"2026-08-19T20:01:29.403616Z","submitted_at":"2024-10-28T19:08:12Z","title":"ShadowKV: KV Cache in Shadows for High-Throughput Long-Context LLM Inference","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.21465","snapshot_observed_at":"2026-08-02T23:32:13.355912Z","title":"Shadowkv: Kv cache in shadows for high-throughput long-context llm inference, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2602.13692","last_updated":"2026-06-30T17:19:18Z","snapshot_observed_at":"2026-08-13T15:15:57.031830Z","submitted_at":"2026-02-14T09:26:41Z","title":"ThunderAgent: A Simple, Fast and Program-Aware Agentic Inference System","version":3},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-02T23:32:13.355912Z"},"links":{"cited_paper":"/paper/2410.21465","citing_paper":"/paper/2602.13692"},"observation_digest":"sha256:e55a844f54d1f828dc3b4025a9d203f6e565eec6938eb19c9011969d9a2181f7","observation_id":"7e99a5d4-efe5-4e04-8563-65f6fea7ffa4","resolution":{"observed_at":"2026-08-02T23:32:13.355912Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.21465","last_updated":"2025-04-25T19:40:54Z","snapshot_observed_at":"2026-08-19T20:01:29.403616Z","submitted_at":"2024-10-28T19:08:12Z","title":"ShadowKV: KV Cache in Shadows for High-Throughput Long-Context LLM Inference","version":3},"cited_work":{"arxiv_id":"2410.21465","doi":null,"metadata_source":"pith","pith_arxiv_id":"2410.21465","snapshot_observed_at":"2026-07-10T01:36:44.234407Z","title":"Shadowkv: Kv cache in shadows for high-throughput long-context llm inference","venue":"cs.LG","work_id":"23d180f7-15ab-4477-b880-cd63ca2f6a95","year":2024},"citing_paper":{"arxiv_id":"2603.10726","last_updated":"2026-05-20T10:27:28Z","snapshot_observed_at":"2026-08-16T00:06:10.598472Z","submitted_at":"2026-03-11T12:59:12Z","title":"PrefixWall: Mitigating Prefix Caching Side Channels in Shared LLM Systems","version":2},"reference_index":61,"source":"pdf_text","source_observed_at":"2026-05-21T12:14:09.509302Z"},"links":{"cited_paper":"/paper/2410.21465","citing_paper":"/paper/2603.10726"},"observation_digest":"sha256:0cb5f418161ad5600d7bd922ffeee7b6e51071603b00d4e6b93b6477fe0fba3c","observation_id":"6c2560c7-d8c1-4487-9f62-eb0da98abf70","resolution":{"observed_at":"2026-05-21T12:15:06.756004Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.21465","last_updated":"2025-04-25T19:40:54Z","snapshot_observed_at":"2026-08-19T20:01:29.403616Z","submitted_at":"2024-10-28T19:08:12Z","title":"ShadowKV: KV Cache in Shadows for High-Throughput Long-Context LLM Inference","version":3},"cited_work":{"arxiv_id":"2410.21465","doi":null,"metadata_source":"pith","pith_arxiv_id":"2410.21465","snapshot_observed_at":"2026-07-10T01:36:44.234407Z","title":"Shadowkv: Kv cache in shadows for high-throughput long-context llm inference","venue":"cs.LG","work_id":"23d180f7-15ab-4477-b880-cd63ca2f6a95","year":2024},"citing_paper":{"arxiv_id":"2604.11627","last_updated":"2026-04-13T15:38:22Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2026-04-13T15:38:22Z","title":"POINTS-Long: Adaptive Dual-Mode Visual Reasoning in MLLMs","version":1},"reference_index":76,"source":"pdf_text","source_observed_at":"2026-05-10T15:23:08.671342Z"},"links":{"cited_paper":"/paper/2410.21465","citing_paper":"/paper/2604.11627"},"observation_digest":"sha256:19b02e9962654360cbfd41719c59f43313aeb611bb33de30d535f37e98ebb8c5","observation_id":"b68da320-0dbd-4fb0-a708-e93f8985d0d7","resolution":{"observed_at":"2026-05-11T10:41:04.055209Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.21465","last_updated":"2025-04-25T19:40:54Z","snapshot_observed_at":"2026-08-19T20:01:29.403616Z","submitted_at":"2024-10-28T19:08:12Z","title":"ShadowKV: KV Cache in Shadows for High-Throughput Long-Context LLM Inference","version":3},"cited_work":{"arxiv_id":"2410.21465","doi":null,"metadata_source":"pith","pith_arxiv_id":"2410.21465","snapshot_observed_at":"2026-07-10T01:36:44.234407Z","title":"Shadowkv: Kv cache in shadows for high-throughput long-context llm inference","venue":"cs.LG","work_id":"23d180f7-15ab-4477-b880-cd63ca2f6a95","year":2024},"citing_paper":{"arxiv_id":"2604.17708","last_updated":"2026-06-02T08:20:34Z","snapshot_observed_at":"2026-08-12T12:58:25.847937Z","submitted_at":"2026-04-20T01:44:18Z","title":"Co-evolving Agent Architectures and Interpretable Reasoning for Automated Optimization","version":2},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-07-05T16:27:03.569327Z"},"links":{"cited_paper":"/paper/2410.21465","citing_paper":"/paper/2604.17708"},"observation_digest":"sha256:3804d87c3db4b0e56de7ba100900b8fab5463aa9de5ea5fc422b24e92364248c","observation_id":"264a3734-c713-461d-9c20-76f23d6b4347","resolution":{"observed_at":"2026-07-05T16:31:16.090941Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.21465","last_updated":"2025-04-25T19:40:54Z","snapshot_observed_at":"2026-08-19T20:01:29.403616Z","submitted_at":"2024-10-28T19:08:12Z","title":"ShadowKV: KV Cache in Shadows for High-Throughput Long-Context LLM Inference","version":3},"cited_work":{"arxiv_id":"2410.21465","doi":null,"metadata_source":"pith","pith_arxiv_id":"2410.21465","snapshot_observed_at":"2026-07-10T01:36:44.234407Z","title":"Shadowkv: Kv cache in shadows for high-throughput long-context llm inference","venue":"cs.LG","work_id":"23d180f7-15ab-4477-b880-cd63ca2f6a95","year":2024},"citing_paper":{"arxiv_id":"2604.17709","last_updated":"2026-06-03T15:24:49Z","snapshot_observed_at":"2026-08-11T01:56:08.814671Z","submitted_at":"2026-04-20T01:47:48Z","title":"DeInfer: Efficient Parallel Inferencing for Decomposed Large Language Models","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-05-10T05:38:17.542951Z"},"links":{"cited_paper":"/paper/2410.21465","citing_paper":"/paper/2604.17709"},"observation_digest":"sha256:684351996e428a287593fffb6033e4356c149e2d3ab0e7e35b341457a70e5e3c","observation_id":"ef1bd31d-b45d-4b87-82fd-7d12adf05460","resolution":{"observed_at":"2026-05-10T05:41:02.125976Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.21465","last_updated":"2025-04-25T19:40:54Z","snapshot_observed_at":"2026-08-19T20:01:29.403616Z","submitted_at":"2024-10-28T19:08:12Z","title":"ShadowKV: KV Cache in Shadows for High-Throughput Long-Context LLM Inference","version":3},"cited_work":{"arxiv_id":"2410.21465","doi":null,"metadata_source":"pith","pith_arxiv_id":"2410.21465","snapshot_observed_at":"2026-07-10T01:36:44.234407Z","title":"Shadowkv: Kv cache in shadows for high-throughput long-context llm inference","venue":"cs.LG","work_id":"23d180f7-15ab-4477-b880-cd63ca2f6a95","year":2024},"citing_paper":{"arxiv_id":"2605.07719","last_updated":"2026-05-08T13:24:39Z","snapshot_observed_at":"2026-08-15T00:27:43.410921Z","submitted_at":"2026-05-08T13:24:39Z","title":"An Efficient Hybrid Sparse Attention with CPU-GPU Parallelism for Long-Context Inference","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-05-11T02:56:28.828593Z"},"links":{"cited_paper":"/paper/2410.21465","citing_paper":"/paper/2605.07719"},"observation_digest":"sha256:0271ac891dd09797744779a55f80995aa40ed82eeb44c596762d12995910dab9","observation_id":"e8b5f251-9dad-492a-9869-7db3bf39207a","resolution":{"observed_at":"2026-05-11T03:05:54.076381Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.21465","last_updated":"2025-04-25T19:40:54Z","snapshot_observed_at":"2026-08-19T20:01:29.403616Z","submitted_at":"2024-10-28T19:08:12Z","title":"ShadowKV: KV Cache in Shadows for High-Throughput Long-Context LLM Inference","version":3},"cited_work":{"arxiv_id":"2410.21465","doi":null,"metadata_source":"pith","pith_arxiv_id":"2410.21465","snapshot_observed_at":"2026-07-10T01:36:44.234407Z","title":"Shadowkv: Kv cache in shadows for high-throughput long-context llm inference","venue":"cs.LG","work_id":"23d180f7-15ab-4477-b880-cd63ca2f6a95","year":2024},"citing_paper":{"arxiv_id":"2605.09649","last_updated":"2026-05-10T16:47:50Z","snapshot_observed_at":"2026-08-12T14:17:43.612550Z","submitted_at":"2026-05-10T16:47:50Z","title":"Make Each Token Count: Towards Improving Long-Context Performance with KV Cache Eviction","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-05-12T05:02:25.513351Z"},"links":{"cited_paper":"/paper/2410.21465","citing_paper":"/paper/2605.09649"},"observation_digest":"sha256:6d45a1390ab663f3635929280ac9b00cf979fc49bc7a2982a4c172dcd6489603","observation_id":"56215eab-bc41-47e8-9332-38d08d6b5f28","resolution":{"observed_at":"2026-05-12T05:41:26.478146Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.21465","last_updated":"2025-04-25T19:40:54Z","snapshot_observed_at":"2026-08-19T20:01:29.403616Z","submitted_at":"2024-10-28T19:08:12Z","title":"ShadowKV: KV Cache in Shadows for High-Throughput Long-Context LLM Inference","version":3},"cited_work":{"arxiv_id":"2410.21465","doi":null,"metadata_source":"pith","pith_arxiv_id":"2410.21465","snapshot_observed_at":"2026-07-10T01:36:44.234407Z","title":"Shadowkv: Kv cache in shadows for high-throughput long-context llm inference","venue":"cs.LG","work_id":"23d180f7-15ab-4477-b880-cd63ca2f6a95","year":2024},"citing_paper":{"arxiv_id":"2605.18753","last_updated":"2026-05-18T17:59:52Z","snapshot_observed_at":"2026-08-17T02:28:01.764345Z","submitted_at":"2026-05-18T17:59:52Z","title":"DashAttention: Differentiable and Adaptive Sparse Hierarchical Attention","version":1},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-05-20T10:50:12.926232Z"},"links":{"cited_paper":"/paper/2410.21465","citing_paper":"/paper/2605.18753"},"observation_digest":"sha256:073f47f31f87a6ec0bd7096fb011a4b891cac1c2d656733b0cd1f63d1b8c68dc","observation_id":"2de173d6-f8da-40a1-b937-7a2692db3acb","resolution":{"observed_at":"2026-05-20T10:53:13.593270Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.21465","last_updated":"2025-04-25T19:40:54Z","snapshot_observed_at":"2026-08-19T20:01:29.403616Z","submitted_at":"2024-10-28T19:08:12Z","title":"ShadowKV: KV Cache in Shadows for High-Throughput Long-Context LLM Inference","version":3},"cited_work":{"arxiv_id":"2410.21465","doi":null,"metadata_source":"pith","pith_arxiv_id":"2410.21465","snapshot_observed_at":"2026-07-10T01:36:44.234407Z","title":"Shadowkv: Kv cache in shadows for high-throughput long-context llm inference","venue":"cs.LG","work_id":"23d180f7-15ab-4477-b880-cd63ca2f6a95","year":2024},"citing_paper":{"arxiv_id":"2606.00866","last_updated":"2026-05-30T19:44:25Z","snapshot_observed_at":"2026-08-12T17:45:10.515441Z","submitted_at":"2026-05-30T19:44:25Z","title":"Idleness is Relative: Exploiting Tool-Call Idle Windows for Offloading in Agentic Systems with MORI","version":1},"reference_index":54,"source":"pdf_text","source_observed_at":"2026-06-28T17:30:56.324289Z"},"links":{"cited_paper":"/paper/2410.21465","citing_paper":"/paper/2606.00866"},"observation_digest":"sha256:b1cc13481f4502dd46b1fedf5a7c376d57f94c25ec9b6185725307893a174a8a","observation_id":"063962f1-9fef-4255-8ef2-5789fd2de9ef","resolution":{"observed_at":"2026-07-01T21:06:13.602579Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.21465","last_updated":"2025-04-25T19:40:54Z","snapshot_observed_at":"2026-08-19T20:01:29.403616Z","submitted_at":"2024-10-28T19:08:12Z","title":"ShadowKV: KV Cache in Shadows for High-Throughput Long-Context LLM Inference","version":3},"cited_work":{"arxiv_id":"2410.21465","doi":null,"metadata_source":"pith","pith_arxiv_id":"2410.21465","snapshot_observed_at":"2026-07-10T01:36:44.234407Z","title":"Shadowkv: Kv cache in shadows for high-throughput long-context llm inference","venue":"cs.LG","work_id":"23d180f7-15ab-4477-b880-cd63ca2f6a95","year":2024},"citing_paper":{"arxiv_id":"2606.03825","last_updated":"2026-06-02T16:07:55Z","snapshot_observed_at":"2026-08-21T01:31:18.679421Z","submitted_at":"2026-06-02T16:07:55Z","title":"Dynamic Short Convolutions Improve Transformers","version":1},"reference_index":103,"source":"arxiv_source","source_observed_at":"2026-06-28T10:48:50.103004Z"},"links":{"cited_paper":"/paper/2410.21465","citing_paper":"/paper/2606.03825"},"observation_digest":"sha256:4492cdb38e209dc9c8ebdc0be9c518ea90cbd5b5b022e161d1037c26fe6f51e7","observation_id":"fc0b3de4-46fa-42f1-a6d9-c72042f05787","resolution":{"observed_at":"2026-07-02T02:36:26.995018Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.21465","last_updated":"2025-04-25T19:40:54Z","snapshot_observed_at":"2026-08-19T20:01:29.403616Z","submitted_at":"2024-10-28T19:08:12Z","title":"ShadowKV: KV Cache in Shadows for High-Throughput Long-Context LLM Inference","version":3},"cited_work":{"arxiv_id":"2410.21465","doi":null,"metadata_source":"pith","pith_arxiv_id":"2410.21465","snapshot_observed_at":"2026-07-10T01:36:44.234407Z","title":"Shadowkv: Kv cache in shadows for high-throughput long-context llm inference","venue":"cs.LG","work_id":"23d180f7-15ab-4477-b880-cd63ca2f6a95","year":2024},"citing_paper":{"arxiv_id":"2606.06453","last_updated":"2026-06-04T17:48:17Z","snapshot_observed_at":"2026-08-08T14:28:58.546665Z","submitted_at":"2026-06-04T17:48:17Z","title":"Vortex: Efficient and Programmable Sparse Attention Serving for AI Agents","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-06-28T01:07:14.691347Z"},"links":{"cited_paper":"/paper/2410.21465","citing_paper":"/paper/2606.06453"},"observation_digest":"sha256:13a1fe917ca9861d0529a52d74ba09dace569aad633ad233eafb0e8fabd3fc96","observation_id":"8373b765-eb4e-468d-851a-30e80082cc0c","resolution":{"observed_at":"2026-07-02T13:36:59.437987Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.21465","last_updated":"2025-04-25T19:40:54Z","snapshot_observed_at":"2026-08-19T20:01:29.403616Z","submitted_at":"2024-10-28T19:08:12Z","title":"ShadowKV: KV Cache in Shadows for High-Throughput Long-Context LLM Inference","version":3},"cited_work":{"arxiv_id":"2410.21465","doi":null,"metadata_source":"pith","pith_arxiv_id":"2410.21465","snapshot_observed_at":"2026-07-10T01:36:44.234407Z","title":"Shadowkv: Kv cache in shadows for high-throughput long-context llm inference","venue":"cs.LG","work_id":"23d180f7-15ab-4477-b880-cd63ca2f6a95","year":2024},"citing_paper":{"arxiv_id":"2606.30389","last_updated":"2026-06-29T14:43:25Z","snapshot_observed_at":"2026-08-13T11:37:17.528311Z","submitted_at":"2026-06-29T14:43:25Z","title":"Predict, Reuse, and Repair: Accelerating Dynamic Sparse Attention for Long-Context LLM Decoding","version":1},"reference_index":36,"source":"arxiv_source","source_observed_at":"2026-06-30T07:03:08.617257Z"},"links":{"cited_paper":"/paper/2410.21465","citing_paper":"/paper/2606.30389"},"observation_digest":"sha256:5c0d6e2cc8c8f380ab80c072af9cbefcfa229b3a1f532d611ed1d2b2585c894e","observation_id":"fa206f97-18aa-4903-b6a7-7d047684e8c8","resolution":{"observed_at":"2026-06-30T07:04:20.907471Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.21465","last_updated":"2025-04-25T19:40:54Z","snapshot_observed_at":"2026-08-19T20:01:29.403616Z","submitted_at":"2024-10-28T19:08:12Z","title":"ShadowKV: KV Cache in Shadows for High-Throughput Long-Context LLM Inference","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.21465","snapshot_observed_at":"2026-07-12T09:50:23.266920Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.02574","last_updated":"2026-06-30T16:12:40Z","snapshot_observed_at":"2026-08-13T16:31:58.188253Z","submitted_at":"2026-06-30T16:12:40Z","title":"From Tensor Buffer to Distributed Memory Hierarchy: A Survey of KV Cache Management for LLM Serving","version":1},"reference_index":111,"source":"pdf_text","source_observed_at":"2026-07-12T09:50:23.266920Z"},"links":{"cited_paper":"/paper/2410.21465","citing_paper":"/paper/2607.02574"},"observation_digest":"sha256:d7a354857919ad3d6a2f9a6451bb3e991f17ddb545a0c43341842fa441b28e5a","observation_id":"25da11a0-f559-4f99-b474-27a1f651eaf0","resolution":{"observed_at":"2026-07-12T09:50:23.266920Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.21465","last_updated":"2025-04-25T19:40:54Z","snapshot_observed_at":"2026-08-19T20:01:29.403616Z","submitted_at":"2024-10-28T19:08:12Z","title":"ShadowKV: KV Cache in Shadows for High-Throughput Long-Context LLM Inference","version":3},"cited_work":{"arxiv_id":"2410.21465","doi":null,"metadata_source":"pith","pith_arxiv_id":"2410.21465","snapshot_observed_at":"2026-07-10T01:36:44.234407Z","title":"Shadowkv: Kv cache in shadows for high-throughput long-context llm inference","venue":"cs.LG","work_id":"23d180f7-15ab-4477-b880-cd63ca2f6a95","year":2024},"citing_paper":{"arxiv_id":"2607.08032","last_updated":"2026-07-09T01:15:03Z","snapshot_observed_at":"2026-08-18T23:35:57.173534Z","submitted_at":"2026-07-09T01:15:03Z","title":"What to Keep, What to Forget: A Rate--Distortion View of Memory Compaction in LLMs and Agents","version":1},"reference_index":109,"source":"pdf_text","source_observed_at":"2026-07-10T01:26:59.421158Z"},"links":{"cited_paper":"/paper/2410.21465","citing_paper":"/paper/2607.08032"},"observation_digest":"sha256:2d3f1d2148154b281802327b96fd3f3cd7b1cd3c65617eac6edee15a6b60e992","observation_id":"d9e434ec-3b15-4c4e-9c3a-aad7ad305b00","resolution":{"observed_at":"2026-07-10T01:36:44.235599Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-22T06:32:14.747728+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.21465","last_updated":"2025-04-25T19:40:54Z","snapshot_observed_at":"2026-08-19T20:01:29.403616Z","submitted_at":"2024-10-28T19:08:12Z","title":"ShadowKV: KV Cache in Shadows for High-Throughput Long-Context LLM Inference","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.21465","snapshot_observed_at":"2026-08-04T13:43:55.710075Z","title":"arXiv preprint arXiv:2410.21465 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.02150","last_updated":"2026-08-05T06:57:37Z","snapshot_observed_at":"2026-08-15T14:31:12.914781Z","submitted_at":"2026-08-03T12:34:31Z","title":"PhyCheck: Fine-Grained Evidence-Grounded Dataset for Physical Law Understanding in Video-LLMs","version":1},"reference_index":75,"source":"arxiv_source","source_observed_at":"2026-08-04T13:43:55.710075Z"},"links":{"cited_paper":"/paper/2410.21465","citing_paper":"/paper/2608.02150"},"observation_digest":"sha256:626010651e27cc59455e7c6844fc5bb8207d4afa95a1080588fb4cca7ba9070e","observation_id":"27517b0b-e45e-47a2-b685-9bab0a290d64","resolution":{"observed_at":"2026-08-04T13:43:55.710075Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.21465","last_updated":"2025-04-25T19:40:54Z","snapshot_observed_at":"2026-08-19T20:01:29.403616Z","submitted_at":"2024-10-28T19:08:12Z","title":"ShadowKV: KV Cache in Shadows for High-Throughput Long-Context LLM Inference","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.21465","snapshot_observed_at":"2026-08-07T00:11:49.064199Z","title":"arXiv preprint arXiv:2410.21465 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.02150","last_updated":"2026-08-05T06:57:37Z","snapshot_observed_at":"2026-08-15T14:31:12.914781Z","submitted_at":"2026-08-03T12:34:31Z","title":"PhyCheck: Fine-Grained Evidence-Grounded Dataset for Physical Law Understanding in Video-LLMs","version":3},"reference_index":75,"source":"arxiv_source","source_observed_at":"2026-08-07T00:11:49.064199Z"},"links":{"cited_paper":"/paper/2410.21465","citing_paper":"/paper/2608.02150"},"observation_digest":"sha256:83b7383d6717bcc823648e164dc4e69b502812fff6dd9c3ebb605fad97a46986","observation_id":"8a0700cc-3573-43c3-8a64-b4cf4e03cc54","resolution":{"observed_at":"2026-08-07T00:11:49.064199Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.21465","last_updated":"2025-04-25T19:40:54Z","snapshot_observed_at":"2026-08-19T20:01:29.403616Z","submitted_at":"2024-10-28T19:08:12Z","title":"ShadowKV: KV Cache in Shadows for High-Throughput Long-Context LLM Inference","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.21465","snapshot_observed_at":"2026-08-05T23:10:35.446082Z","title":"ShadowKV: KV Cache in Shadows for High-Throughput Long-Context LLM Inference","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2608.03228","last_updated":"2026-08-05T20:50:42Z","snapshot_observed_at":"2026-08-14T09:50:26.014755Z","submitted_at":"2026-08-04T06:59:29Z","title":"SAKI: Score-Aware Low-Rank Key Indexing with Random-Matrix Noise Correction for KV Retrieval","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-05T23:10:35.446082Z"},"links":{"cited_paper":"/paper/2410.21465","citing_paper":"/paper/2608.03228"},"observation_digest":"sha256:3acc2368b59e2d46a0dad265cd437d13f7f4d342f4d0c4a1811cb94b22bec8e3","observation_id":"a205d2e9-b4d4-41cd-b526-f9ceeb3fcd79","resolution":{"observed_at":"2026-08-05T23:10:35.446082Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.21465","last_updated":"2025-04-25T19:40:54Z","snapshot_observed_at":"2026-08-19T20:01:29.403616Z","submitted_at":"2024-10-28T19:08:12Z","title":"ShadowKV: KV Cache in Shadows for High-Throughput Long-Context LLM Inference","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.21465","snapshot_observed_at":"2026-08-08T00:54:42.703743Z","title":"ShadowKV: KV Cache in Shadows for High-Throughput Long-Context LLM Inference","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2608.03228","last_updated":"2026-08-05T20:50:42Z","snapshot_observed_at":"2026-08-14T09:50:26.014755Z","submitted_at":"2026-08-04T06:59:29Z","title":"SAKI: Score-Aware Low-Rank Key Indexing with Random-Matrix Noise Correction for KV Retrieval","version":2},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-08T00:54:42.703743Z"},"links":{"cited_paper":"/paper/2410.21465","citing_paper":"/paper/2608.03228"},"observation_digest":"sha256:b4e676872c476e9a9f1fadf5aab10b3ddea549cc90bd9210b6ca7836e8a612d6","observation_id":"9437c568-adbd-4aae-a0e7-bfa78d3487cb","resolution":{"observed_at":"2026-08-08T00:54:42.703743Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.21465","last_updated":"2025-04-25T19:40:54Z","snapshot_observed_at":"2026-08-19T20:01:29.403616Z","submitted_at":"2024-10-28T19:08:12Z","title":"ShadowKV: KV Cache in Shadows for High-Throughput Long-Context LLM Inference","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.21465","snapshot_observed_at":"2026-08-12T00:32:37.689283Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2608.08097","last_updated":"2026-08-08T12:27:03Z","snapshot_observed_at":"2026-08-18T04:59:06.437774Z","submitted_at":"2026-08-08T12:27:03Z","title":"OasisKV: Scaling In-Decode KV Cache Beyond HBM with Lookahead Sparse Prefetching","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-12T00:32:37.689283Z"},"links":{"cited_paper":"/paper/2410.21465","citing_paper":"/paper/2608.08097"},"observation_digest":"sha256:aebc6bdaf1c3ee3c137f2fff9496d2b443da4868f2b05a510838ed244c348f5f","observation_id":"3a033333-3dc6-44ff-94f6-e1e6a600920c","resolution":{"observed_at":"2026-08-12T00:32:37.689283Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2410.21465/citation-record","integrity":"/paper/2410.21465/integrity","json":"/paper/2410.21465/citation-record.json","paper":"/paper/2410.21465"},"outbound":[],"paper":{"arxiv_id":"2410.21465","last_updated":"2025-04-25T19:40:54Z","latest_version":3,"primary_category":"cs.LG","snapshot_observed_at":"2026-08-19T20:01:29.403616Z","submitted_at":"2024-10-28T19:08:12Z","title":"ShadowKV: KV Cache in Shadows for High-Throughput Long-Context LLM Inference"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-22T06:32:14.747728+00:00","source":"crossref"},{"observed_at":"2026-08-22T06:32:06.552537+00:00","source":"retraction_watch"}],"thesis":"As of 22 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 33 inbound Pith citation observations for arXiv:2410.21465."}