{"as_of":"2026-08-08T20:01:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:894fb41d4e3f40b5a33e955cd2f768394aafe2d3c59d9414643af350f279ef4c","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":10,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":10,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-08T06:32:00.761636+00:00","state":"measured"},{"denominator":10,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":10,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-08T14:29:13.944332Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-05-18T08:12:01.964814Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2407.12820","last_updated":"2025-03-30T08:13:50Z","snapshot_observed_at":"2026-07-06T18:48:00.070861Z","submitted_at":"2024-07-01T13:05:42Z","title":"PQCache: Product Quantization-based KVCache for Long Context LLM Inference","version":2},"cited_work":{"arxiv_id":"2407.12820","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2407.12820","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Pqcache: Product quantization-based kvcache for long context llm inference","venue":null,"work_id":"4369ba22-96b5-4e5b-ab6d-72b22ce8949a","year":2024},"citing_paper":{"arxiv_id":"2409.10516","last_updated":"2024-12-31T07:11:00Z","snapshot_observed_at":"2026-08-07T09:13:28.333702Z","submitted_at":"2024-09-16T17:59:52Z","title":"RetrievalAttention: Accelerating Long-Context LLM Inference via Vector Retrieval","version":3},"reference_index":118,"source":"arxiv_source","source_observed_at":"2026-05-18T08:12:01.798459Z"},"links":{"cited_paper":"/paper/2407.12820","citing_paper":"/paper/2409.10516"},"observation_digest":"sha256:7f1a13eff0595e85aa136199907994b0199ba214932adab0e309c98ce7813f11","observation_id":"cc2cedf8-7723-493a-9b8e-dc54aff2add3","resolution":{"observed_at":"2026-05-18T08:12:01.969232Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2407.12820","last_updated":"2025-03-30T08:13:50Z","snapshot_observed_at":"2026-07-06T18:48:00.070861Z","submitted_at":"2024-07-01T13:05:42Z","title":"PQCache: Product Quantization-based KVCache for Long Context LLM Inference","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.12820","snapshot_observed_at":"2026-08-08T14:29:13.944332Z","title":"Pqcache: Product quantization-based kvcache for long context llm inference, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.06766","last_updated":"2025-02-12T15:55:37Z","snapshot_observed_at":"2026-08-08T14:22:26.963746Z","submitted_at":"2025-02-10T18:47:04Z","title":"Exploiting Sparsity for Long Context Inference: Million Token Contexts on Commodity GPUs","version":2},"reference_index":44,"source":"arxiv_source","source_observed_at":"2026-08-08T14:29:13.944332Z"},"links":{"cited_paper":"/paper/2407.12820","citing_paper":"/paper/2502.06766"},"observation_digest":"sha256:82511003e844a1afa19e1b68c38e4a09515b7efb9b67037225b477f499c519dc","observation_id":"a0e7f2bf-fb87-479a-8ca3-3ca35c60e3a2","resolution":{"observed_at":"2026-08-08T14:29:13.944332Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.12820","last_updated":"2025-03-30T08:13:50Z","snapshot_observed_at":"2026-07-06T18:48:00.070861Z","submitted_at":"2024-07-01T13:05:42Z","title":"PQCache: Product Quantization-based KVCache for Long Context LLM Inference","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.12820","snapshot_observed_at":"2026-08-07T14:16:42.391215Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.19586","last_updated":"2025-05-27T03:16:32Z","snapshot_observed_at":"2026-08-07T14:09:10.072359Z","submitted_at":"2025-05-26T07:00:04Z","title":"TailorKV: A Hybrid Framework for Long-Context Inference via Tailored KV Cache Optimization","version":2},"reference_index":36,"source":"arxiv_source","source_observed_at":"2026-08-07T14:16:42.391215Z"},"links":{"cited_paper":"/paper/2407.12820","citing_paper":"/paper/2505.19586"},"observation_digest":"sha256:4b90006cad3b69a8cfaef53e29433cb02b15ffcec0de110870220822ccc93f7d","observation_id":"d06da1e1-b521-4fa2-8675-91a428ba9431","resolution":{"observed_at":"2026-08-07T14:16:42.391215Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.12820","last_updated":"2025-03-30T08:13:50Z","snapshot_observed_at":"2026-07-06T18:48:00.070861Z","submitted_at":"2024-07-01T13:05:42Z","title":"PQCache: Product Quantization-based KVCache for Long Context LLM Inference","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.12820","snapshot_observed_at":"2026-08-07T12:39:05.152541Z","title":"Pqcache: Product quantization- based kvcache for long context llm inference,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.24179","last_updated":"2025-05-30T03:40:24Z","snapshot_observed_at":"2026-08-07T12:28:26.705118Z","submitted_at":"2025-05-30T03:40:24Z","title":"SALE : Low-bit Estimation for Efficient Sparse Attention in Long-context LLM Prefilling","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-07T12:39:05.152541Z"},"links":{"cited_paper":"/paper/2407.12820","citing_paper":"/paper/2505.24179"},"observation_digest":"sha256:13f84c0d400ebbf8e920e739b738797db90aca26a1f9098f6a7bcc57bac59be9","observation_id":"f0d3eeb6-e310-4965-99c8-b85aa2d41096","resolution":{"observed_at":"2026-08-07T12:39:05.152541Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.12820","last_updated":"2025-03-30T08:13:50Z","snapshot_observed_at":"2026-07-06T18:48:00.070861Z","submitted_at":"2024-07-01T13:05:42Z","title":"PQCache: Product Quantization-based KVCache for Long Context LLM Inference","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.12820","snapshot_observed_at":"2026-08-07T05:06:32.745585Z","title":"Pqcache: Product quantization-based kvcache for long context llm inference","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.08889","last_updated":"2025-06-10T15:17:26Z","snapshot_observed_at":"2026-08-07T11:00:03.513055Z","submitted_at":"2025-06-10T15:17:26Z","title":"SeerAttention-R: Sparse Attention Adaptation for Long Reasoning","version":1},"reference_index":75,"source":"pdf_text","source_observed_at":"2026-08-07T05:06:32.745585Z"},"links":{"cited_paper":"/paper/2407.12820","citing_paper":"/paper/2506.08889"},"observation_digest":"sha256:95d3ffcc1f1b01398dbe15c884c26d632e3ed5b08cf089d2cd0994acb5609e91","observation_id":"506d2d3c-2bcb-4aea-9efb-b6071374512a","resolution":{"observed_at":"2026-08-07T05:06:32.745585Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.12820","last_updated":"2025-03-30T08:13:50Z","snapshot_observed_at":"2026-07-06T18:48:00.070861Z","submitted_at":"2024-07-01T13:05:42Z","title":"PQCache: Product Quantization-based KVCache for Long Context LLM Inference","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.12820","snapshot_observed_at":"2026-08-07T12:38:40.900312Z","title":"Pqcache: Product quantization-based kvcache for long context llm inference","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.15704","last_updated":"2025-05-30T02:35:59Z","snapshot_observed_at":"2026-08-07T12:30:17.322980Z","submitted_at":"2025-05-30T02:35:59Z","title":"Learn from the Past: Fast Sparse Indexing for Large Language Model Decoding","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-07T12:38:40.900312Z"},"links":{"cited_paper":"/paper/2407.12820","citing_paper":"/paper/2506.15704"},"observation_digest":"sha256:ced4e6335a07bc56030c39d3542338253ccf0db3519d3d4f9f6ba244f2700ba3","observation_id":"0a5774c3-cbdc-4a02-aab5-6c6eaea9326c","resolution":{"observed_at":"2026-08-07T12:38:40.900312Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.12820","last_updated":"2025-03-30T08:13:50Z","snapshot_observed_at":"2026-07-06T18:48:00.070861Z","submitted_at":"2024-07-01T13:05:42Z","title":"PQCache: Product Quantization-based KVCache for Long Context LLM Inference","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.12820","snapshot_observed_at":"2026-08-06T23:00:11.220828Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.20187","last_updated":"2025-07-02T05:12:29Z","snapshot_observed_at":"2026-08-07T20:42:38.846597Z","submitted_at":"2025-06-25T07:26:42Z","title":"Breaking the Boundaries of Long-Context LLM Inference: Adaptive KV Management on a Single Commodity GPU","version":2},"reference_index":67,"source":"pdf_text","source_observed_at":"2026-08-06T23:00:11.220828Z"},"links":{"cited_paper":"/paper/2407.12820","citing_paper":"/paper/2506.20187"},"observation_digest":"sha256:614c6e1e6b16ca9fa3e6ae6bf929e47686497058fc1d930f83586b05c6e092c6","observation_id":"2f441ba8-0dff-49a9-a7c4-17d4fbf2545c","resolution":{"observed_at":"2026-08-06T23:00:11.220828Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.12820","last_updated":"2025-03-30T08:13:50Z","snapshot_observed_at":"2026-07-06T18:48:00.070861Z","submitted_at":"2024-07-01T13:05:42Z","title":"PQCache: Product Quantization-based KVCache for Long Context LLM Inference","version":2},"cited_work":{"arxiv_id":"2407.12820","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2407.12820","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Pqcache: Product quantization-based kvcache for long context llm inference","venue":null,"work_id":"4369ba22-96b5-4e5b-ab6d-72b22ce8949a","year":2024},"citing_paper":{"arxiv_id":"2604.10539","last_updated":"2026-04-12T09:02:20Z","snapshot_observed_at":"2026-07-06T22:59:09.778842Z","submitted_at":"2026-04-12T09:02:20Z","title":"IceCache: Memory-efficient KV-cache Management for Long-Sequence LLMs","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-05-10T15:46:43.682254Z"},"links":{"cited_paper":"/paper/2407.12820","citing_paper":"/paper/2604.10539"},"observation_digest":"sha256:c880bb06560f937bc50ee4bc8d19f9cdb1e5550224722eb60a31cd71824b89ec","observation_id":"c88da4b3-afe2-46ae-b2bc-57b68f33f65d","resolution":{"observed_at":"2026-05-11T09:51:01.188085Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2407.12820","last_updated":"2025-03-30T08:13:50Z","snapshot_observed_at":"2026-07-06T18:48:00.070861Z","submitted_at":"2024-07-01T13:05:42Z","title":"PQCache: Product Quantization-based KVCache for Long Context LLM Inference","version":2},"cited_work":{"arxiv_id":"2407.12820","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2407.12820","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Pqcache: Product quantization-based kvcache for long context llm inference","venue":null,"work_id":"4369ba22-96b5-4e5b-ab6d-72b22ce8949a","year":2024},"citing_paper":{"arxiv_id":"2604.18137","last_updated":"2026-04-20T12:04:51Z","snapshot_observed_at":"2026-07-06T23:05:04.279268Z","submitted_at":"2026-04-20T12:04:51Z","title":"AQPIM: Breaking the PIM Capacity Wall for LLMs with In-Memory Activation Quantization","version":1},"reference_index":77,"source":"pdf_text","source_observed_at":"2026-05-10T04:09:11.684432Z"},"links":{"cited_paper":"/paper/2407.12820","citing_paper":"/paper/2604.18137"},"observation_digest":"sha256:ee114f112c6083e47eed7a7c512d7f4337e4738aaf338559c4627eb24050ca28","observation_id":"58f9a7d8-cea5-49c0-9c7d-b171f4d489bf","resolution":{"observed_at":"2026-05-11T12:06:04.343444Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2407.12820","last_updated":"2025-03-30T08:13:50Z","snapshot_observed_at":"2026-07-06T18:48:00.070861Z","submitted_at":"2024-07-01T13:05:42Z","title":"PQCache: Product Quantization-based KVCache for Long Context LLM Inference","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.12820","snapshot_observed_at":"2026-07-31T11:49:11.799232Z","title":"PQCache: Product quantization-based KVCache for long context LLM inference.Proceedings of the ACM on Management of Data, 3(3):201:1–201:30, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.24555","last_updated":"2026-07-27T15:28:52Z","snapshot_observed_at":"2026-08-02T06:50:39.173138Z","submitted_at":"2026-07-27T15:28:52Z","title":"LOCKS: Page-Local Compact Key Summaries for Efficient Long-Context Decoding","version":1},"reference_index":72,"source":"pdf_text","source_observed_at":"2026-07-31T11:49:11.799232Z"},"links":{"cited_paper":"/paper/2407.12820","citing_paper":"/paper/2607.24555"},"observation_digest":"sha256:472abb664119f9749edefba2d434755b9454ef1889e401bc8a3716d22bb44eb0","observation_id":"e1473abf-d04b-484b-a419-926e0cb80548","resolution":{"observed_at":"2026-07-31T11:49:11.799232Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2407.12820/citation-record","integrity":"/paper/2407.12820/integrity","json":"/paper/2407.12820/citation-record.json","paper":"/paper/2407.12820"},"outbound":[],"paper":{"arxiv_id":"2407.12820","last_updated":"2025-03-30T08:13:50Z","latest_version":2,"primary_category":"cs.CL","snapshot_observed_at":"2026-07-06T18:48:00.070861Z","submitted_at":"2024-07-01T13:05:42Z","title":"PQCache: Product Quantization-based KVCache for Long Context LLM Inference"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"thesis":"As of 8 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 10 inbound Pith citation observations for arXiv:2407.12820."}