{"as_of":"2026-08-05T08:52:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:6c23918c309aa231da3b6f549682c701ba892a3768ee9dd2793c325b91650bc2","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":14,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":14,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-05T06:32:48.257954+00:00","state":"measured"},{"denominator":14,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":14,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-04T02:48:35.073858Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"pith","source_observed_at":"2026-07-10T01:36:44.326797Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2405.04437","last_updated":"2025-01-29T04:10:41Z","snapshot_observed_at":"2026-08-05T02:13:00.249260Z","submitted_at":"2024-05-07T16:00:32Z","title":"vAttention: Dynamic Memory Management for Serving LLMs without PagedAttention","version":3},"cited_work":{"arxiv_id":"2405.04437","doi":null,"metadata_source":"pith","pith_arxiv_id":"2405.04437","snapshot_observed_at":"2026-07-10T01:36:44.326797Z","title":"vAttention: Dynamic Memory Management for Serving LLMs without PagedAttention","venue":"cs.LG","work_id":"3610d9ca-4370-41c0-9fcb-cc624f5b4d77","year":2024},"citing_paper":{"arxiv_id":"2506.15155","last_updated":"2026-05-06T17:46:08Z","snapshot_observed_at":"2026-07-06T21:44:04.066471Z","submitted_at":"2025-06-18T05:56:01Z","title":"eLLM: Elastic Memory Management Framework for Efficient LLM Serving","version":2},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-05-19T09:41:45.544399Z"},"links":{"cited_paper":"/paper/2405.04437","citing_paper":"/paper/2506.15155"},"observation_digest":"sha256:0f83961ffb5085088acbf219cf686cb06d0024f167b6ed1f63531267a2b52cb3","observation_id":"4cbac68f-45a7-42af-88ab-5209a864596e","resolution":{"observed_at":"2026-05-19T09:42:13.910745Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2405.04437","last_updated":"2025-01-29T04:10:41Z","snapshot_observed_at":"2026-08-05T02:13:00.249260Z","submitted_at":"2024-05-07T16:00:32Z","title":"vAttention: Dynamic Memory Management for Serving LLMs without PagedAttention","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.04437","snapshot_observed_at":"2026-07-13T22:20:51.838309Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2603.18897","last_updated":"2026-06-16T09:34:16Z","snapshot_observed_at":"2026-07-13T22:20:51.137198Z","submitted_at":"2026-03-19T13:36:50Z","title":"Parallelizing Tool Execution and LLM Generation for Low-Latency Agent Serving","version":3},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-07-13T22:20:51.838309Z"},"links":{"cited_paper":"/paper/2405.04437","citing_paper":"/paper/2603.18897"},"observation_digest":"sha256:179f3fddfda2c8e2e9fb87b9ed40a0f7fcc8eb3394fc46b66f835903c7c34ff2","observation_id":"0f85f268-14ee-492e-93a6-fd93f1c5731b","resolution":{"observed_at":"2026-07-13T22:20:51.838309Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.04437","last_updated":"2025-01-29T04:10:41Z","snapshot_observed_at":"2026-08-05T02:13:00.249260Z","submitted_at":"2024-05-07T16:00:32Z","title":"vAttention: Dynamic Memory Management for Serving LLMs without PagedAttention","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.04437","snapshot_observed_at":"2026-07-13T09:09:32.932051Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2604.06035","last_updated":"2026-06-02T01:46:36Z","snapshot_observed_at":"2026-08-01T13:27:50.254470Z","submitted_at":"2026-04-07T16:30:54Z","title":"cuRAMSES: Scalable AMR Optimizations for Large-Scale Cosmological Simulations","version":2},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-07-13T09:09:32.932051Z"},"links":{"cited_paper":"/paper/2405.04437","citing_paper":"/paper/2604.06035"},"observation_digest":"sha256:598f39e9ff533915e92b4d8bffcee29290a2f4de77a9320c7b68d11027c0fa33","observation_id":"deea181f-38ac-4696-886d-b92bb444cd29","resolution":{"observed_at":"2026-07-13T09:09:32.932051Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.04437","last_updated":"2025-01-29T04:10:41Z","snapshot_observed_at":"2026-08-05T02:13:00.249260Z","submitted_at":"2024-05-07T16:00:32Z","title":"vAttention: Dynamic Memory Management for Serving LLMs without PagedAttention","version":3},"cited_work":{"arxiv_id":"2405.04437","doi":null,"metadata_source":"pith","pith_arxiv_id":"2405.04437","snapshot_observed_at":"2026-07-10T01:36:44.326797Z","title":"vAttention: Dynamic Memory Management for Serving LLMs without PagedAttention","venue":"cs.LG","work_id":"3610d9ca-4370-41c0-9fcb-cc624f5b4d77","year":2024},"citing_paper":{"arxiv_id":"2604.06036","last_updated":"2026-04-09T09:40:36Z","snapshot_observed_at":"2026-07-06T22:54:34.877944Z","submitted_at":"2026-04-07T16:31:45Z","title":"CodecSight: Leveraging Video Codec Signals for Efficient Streaming VLM Inference","version":3},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-05-10T18:45:26.555395Z"},"links":{"cited_paper":"/paper/2405.04437","citing_paper":"/paper/2604.06036"},"observation_digest":"sha256:9f15ca583c5b9eec7be2b5165bbb414472bbc82f92da32c7dfbe44a131210be2","observation_id":"9d0d39ea-3a51-4936-870d-f1591f511c78","resolution":{"observed_at":"2026-05-11T00:00:52.993784Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2405.04437","last_updated":"2025-01-29T04:10:41Z","snapshot_observed_at":"2026-08-05T02:13:00.249260Z","submitted_at":"2024-05-07T16:00:32Z","title":"vAttention: Dynamic Memory Management for Serving LLMs without PagedAttention","version":3},"cited_work":{"arxiv_id":"2405.04437","doi":null,"metadata_source":"pith","pith_arxiv_id":"2405.04437","snapshot_observed_at":"2026-07-10T01:36:44.326797Z","title":"vAttention: Dynamic Memory Management for Serving LLMs without PagedAttention","venue":"cs.LG","work_id":"3610d9ca-4370-41c0-9fcb-cc624f5b4d77","year":2024},"citing_paper":{"arxiv_id":"2605.09735","last_updated":"2026-06-30T03:39:11Z","snapshot_observed_at":"2026-07-06T23:21:49.499951Z","submitted_at":"2026-05-10T20:10:26Z","title":"KV-RM: Regularizing KV-Cache Movement for Static-Graph LLM Serving","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-05-12T03:57:30.104609Z"},"links":{"cited_paper":"/paper/2405.04437","citing_paper":"/paper/2605.09735"},"observation_digest":"sha256:0c8df428d75d168e4960afeec7cfc5a412f62c01c44430849c077d98e7f060f8","observation_id":"1421fa3b-f095-4da7-88e5-d953020acd85","resolution":{"observed_at":"2026-05-12T06:51:27.063199Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2405.04437","last_updated":"2025-01-29T04:10:41Z","snapshot_observed_at":"2026-08-05T02:13:00.249260Z","submitted_at":"2024-05-07T16:00:32Z","title":"vAttention: Dynamic Memory Management for Serving LLMs without PagedAttention","version":3},"cited_work":{"arxiv_id":"2405.04437","doi":null,"metadata_source":"pith","pith_arxiv_id":"2405.04437","snapshot_observed_at":"2026-07-10T01:36:44.326797Z","title":"vAttention: Dynamic Memory Management for Serving LLMs without PagedAttention","venue":"cs.LG","work_id":"3610d9ca-4370-41c0-9fcb-cc624f5b4d77","year":2024},"citing_paper":{"arxiv_id":"2605.09735","last_updated":"2026-06-30T03:39:11Z","snapshot_observed_at":"2026-07-06T23:21:49.499951Z","submitted_at":"2026-05-10T20:10:26Z","title":"KV-RM: Regularizing KV-Cache Movement for Static-Graph LLM Serving","version":2},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-07-01T08:05:44.256565Z"},"links":{"cited_paper":"/paper/2405.04437","citing_paper":"/paper/2605.09735"},"observation_digest":"sha256:0010fbb515073ed5955ea3a7c133d9b7dae061a9bc6e87eee9e44772ebec91a6","observation_id":"85c56d4b-216e-499d-8bbe-7f370039632a","resolution":{"observed_at":"2026-07-01T08:15:32.450518Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2405.04437","last_updated":"2025-01-29T04:10:41Z","snapshot_observed_at":"2026-08-05T02:13:00.249260Z","submitted_at":"2024-05-07T16:00:32Z","title":"vAttention: Dynamic Memory Management for Serving LLMs without PagedAttention","version":3},"cited_work":{"arxiv_id":"2405.04437","doi":null,"metadata_source":"pith","pith_arxiv_id":"2405.04437","snapshot_observed_at":"2026-07-10T01:36:44.326797Z","title":"vAttention: Dynamic Memory Management for Serving LLMs without PagedAttention","venue":"cs.LG","work_id":"3610d9ca-4370-41c0-9fcb-cc624f5b4d77","year":2024},"citing_paper":{"arxiv_id":"2606.20295","last_updated":"2026-07-24T07:54:17Z","snapshot_observed_at":"2026-08-02T10:48:55.986116Z","submitted_at":"2026-06-18T14:33:09Z","title":"Token-Operations-Oriented Inference Optimization Techniques for Large Models","version":1},"reference_index":200,"source":"pdf_text","source_observed_at":"2026-06-26T16:15:22.543601Z"},"links":{"cited_paper":"/paper/2405.04437","citing_paper":"/paper/2606.20295"},"observation_digest":"sha256:859a4cf63a8e4ed07284dfd285751947c78a7982607ef37dfeab1b185b1761a3","observation_id":"861b87d2-2475-4f59-9a8c-87be60491b3a","resolution":{"observed_at":"2026-07-04T05:09:36.541010Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2405.04437","last_updated":"2025-01-29T04:10:41Z","snapshot_observed_at":"2026-08-05T02:13:00.249260Z","submitted_at":"2024-05-07T16:00:32Z","title":"vAttention: Dynamic Memory Management for Serving LLMs without PagedAttention","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.04437","snapshot_observed_at":"2026-08-02T10:49:22.213295Z","title":"vAttention: Dynamic Memory Management for Serving LLMs without PagedAttention, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2606.20295","last_updated":"2026-07-24T07:54:17Z","snapshot_observed_at":"2026-08-02T10:48:55.986116Z","submitted_at":"2026-06-18T14:33:09Z","title":"Token-Operations-Oriented Inference Optimization Techniques for Large Models","version":2},"reference_index":187,"source":"pdf_text","source_observed_at":"2026-08-02T10:49:22.213295Z"},"links":{"cited_paper":"/paper/2405.04437","citing_paper":"/paper/2606.20295"},"observation_digest":"sha256:02746f52f830c86dac7d1b83147f318dcf1f237ab3f1b53064ac82b2b99ff947","observation_id":"420d3c7e-fe8b-4d58-91c7-782b6db6bb72","resolution":{"observed_at":"2026-08-02T10:49:22.213295Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.04437","last_updated":"2025-01-29T04:10:41Z","snapshot_observed_at":"2026-08-05T02:13:00.249260Z","submitted_at":"2024-05-07T16:00:32Z","title":"vAttention: Dynamic Memory Management for Serving LLMs without PagedAttention","version":3},"cited_work":{"arxiv_id":"2405.04437","doi":null,"metadata_source":"pith","pith_arxiv_id":"2405.04437","snapshot_observed_at":"2026-07-10T01:36:44.326797Z","title":"vAttention: Dynamic Memory Management for Serving LLMs without PagedAttention","venue":"cs.LG","work_id":"3610d9ca-4370-41c0-9fcb-cc624f5b4d77","year":2024},"citing_paper":{"arxiv_id":"2606.20537","last_updated":"2026-06-18T17:49:36Z","snapshot_observed_at":"2026-07-06T23:55:38.979306Z","submitted_at":"2026-06-18T17:49:36Z","title":"Execution-State Capsules: Graph-Bound Execution-State Checkpoint and Restore for Low-Latency, Small-Batch, On-Device Physical-AI Serving","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-06-26T17:39:47.477338Z"},"links":{"cited_paper":"/paper/2405.04437","citing_paper":"/paper/2606.20537"},"observation_digest":"sha256:a3162b793ccd719a573d96c31672374d41b2a95235c1da8252d352fed34d04bb","observation_id":"2ead7c60-3ad8-49b8-a333-4a79cadf0d32","resolution":{"observed_at":"2026-07-04T03:49:29.923241Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2405.04437","last_updated":"2025-01-29T04:10:41Z","snapshot_observed_at":"2026-08-05T02:13:00.249260Z","submitted_at":"2024-05-07T16:00:32Z","title":"vAttention: Dynamic Memory Management for Serving LLMs without PagedAttention","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.04437","snapshot_observed_at":"2026-08-02T10:40:08.018921Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2606.21848","last_updated":"2026-07-31T19:30:57Z","snapshot_observed_at":"2026-08-05T08:15:51.286218Z","submitted_at":"2026-06-20T03:12:30Z","title":"Keyless Attention: Value-Space Routing and Value-Only Caching for Efficient Transformers","version":2},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-02T10:40:08.018921Z"},"links":{"cited_paper":"/paper/2405.04437","citing_paper":"/paper/2606.21848"},"observation_digest":"sha256:b4a685d96edea45e6305e734b4848609891fc613ea9d74d745b1a216927776d4","observation_id":"c6687955-f488-49d7-b25e-5c4ee1f5523a","resolution":{"observed_at":"2026-08-02T10:40:08.018921Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.04437","last_updated":"2025-01-29T04:10:41Z","snapshot_observed_at":"2026-08-05T02:13:00.249260Z","submitted_at":"2024-05-07T16:00:32Z","title":"vAttention: Dynamic Memory Management for Serving LLMs without PagedAttention","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.04437","snapshot_observed_at":"2026-08-04T02:48:35.073858Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2606.21848","last_updated":"2026-07-31T19:30:57Z","snapshot_observed_at":"2026-08-05T08:15:51.286218Z","submitted_at":"2026-06-20T03:12:30Z","title":"Keyless Attention: Value-Space Routing and Value-Only Caching for Efficient Transformers","version":3},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-04T02:48:35.073858Z"},"links":{"cited_paper":"/paper/2405.04437","citing_paper":"/paper/2606.21848"},"observation_digest":"sha256:09228632d9573d6598176f90a5fca826f2bed72d92c43ae1b75631d3553fa560","observation_id":"0778cff1-9044-4a61-a263-8fbf6d1c9011","resolution":{"observed_at":"2026-08-04T02:48:35.073858Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.04437","last_updated":"2025-01-29T04:10:41Z","snapshot_observed_at":"2026-08-05T02:13:00.249260Z","submitted_at":"2024-05-07T16:00:32Z","title":"vAttention: Dynamic Memory Management for Serving LLMs without PagedAttention","version":3},"cited_work":{"arxiv_id":"2405.04437","doi":null,"metadata_source":"pith","pith_arxiv_id":"2405.04437","snapshot_observed_at":"2026-07-10T01:36:44.326797Z","title":"vAttention: Dynamic Memory Management for Serving LLMs without PagedAttention","venue":"cs.LG","work_id":"3610d9ca-4370-41c0-9fcb-cc624f5b4d77","year":2024},"citing_paper":{"arxiv_id":"2606.26666","last_updated":"2026-07-01T08:54:32Z","snapshot_observed_at":"2026-08-03T23:08:24.179797Z","submitted_at":"2026-06-25T06:56:43Z","title":"PersistentKV: Page-Aware Decode Scheduling for Long-Context LLM Serving on Commodity GPUs","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-06-26T05:23:08.673299Z"},"links":{"cited_paper":"/paper/2405.04437","citing_paper":"/paper/2606.26666"},"observation_digest":"sha256:af900a656181318288be5322c796ba4e6d9c6a175b35bacc8865c2e47fd5b80e","observation_id":"d400d8ad-2f2c-46c2-9112-85c7cd931a85","resolution":{"observed_at":"2026-07-04T13:19:50.308008Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2405.04437","last_updated":"2025-01-29T04:10:41Z","snapshot_observed_at":"2026-08-05T02:13:00.249260Z","submitted_at":"2024-05-07T16:00:32Z","title":"vAttention: Dynamic Memory Management for Serving LLMs without PagedAttention","version":3},"cited_work":{"arxiv_id":"2405.04437","doi":null,"metadata_source":"pith","pith_arxiv_id":"2405.04437","snapshot_observed_at":"2026-07-10T01:36:44.326797Z","title":"vAttention: Dynamic Memory Management for Serving LLMs without PagedAttention","venue":"cs.LG","work_id":"3610d9ca-4370-41c0-9fcb-cc624f5b4d77","year":2024},"citing_paper":{"arxiv_id":"2606.26666","last_updated":"2026-07-01T08:54:32Z","snapshot_observed_at":"2026-08-03T23:08:24.179797Z","submitted_at":"2026-06-25T06:56:43Z","title":"PersistentKV: Page-Aware Decode Scheduling for Long-Context LLM Serving on Commodity GPUs","version":2},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-07-02T21:04:43.874528Z"},"links":{"cited_paper":"/paper/2405.04437","citing_paper":"/paper/2606.26666"},"observation_digest":"sha256:f4019e95d8ee061589954daa77b5ca74fe7869d59c32a96b97caeb86f6869586","observation_id":"1b961289-45f5-427e-93d7-a5c03cd83641","resolution":{"observed_at":"2026-07-02T21:07:23.419736Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2405.04437","last_updated":"2025-01-29T04:10:41Z","snapshot_observed_at":"2026-08-05T02:13:00.249260Z","submitted_at":"2024-05-07T16:00:32Z","title":"vAttention: Dynamic Memory Management for Serving LLMs without PagedAttention","version":3},"cited_work":{"arxiv_id":"2405.04437","doi":null,"metadata_source":"pith","pith_arxiv_id":"2405.04437","snapshot_observed_at":"2026-07-10T01:36:44.326797Z","title":"vAttention: Dynamic Memory Management for Serving LLMs without PagedAttention","venue":"cs.LG","work_id":"3610d9ca-4370-41c0-9fcb-cc624f5b4d77","year":2024},"citing_paper":{"arxiv_id":"2607.08032","last_updated":"2026-07-09T01:15:03Z","snapshot_observed_at":"2026-08-02T16:38:21.894642Z","submitted_at":"2026-07-09T01:15:03Z","title":"What to Keep, What to Forget: A Rate--Distortion View of Memory Compaction in LLMs and Agents","version":1},"reference_index":95,"source":"pdf_text","source_observed_at":"2026-07-10T01:26:59.421158Z"},"links":{"cited_paper":"/paper/2405.04437","citing_paper":"/paper/2607.08032"},"observation_digest":"sha256:7a005fac69d8c8f5d440af509741047205e0b3bbd09c9bfc026adf9b7f6a0108","observation_id":"a56edc14-4e35-4800-815f-b329b3a2a82f","resolution":{"observed_at":"2026-07-10T01:36:44.327959Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-05T06:32:48.257954+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2405.04437/citation-record","integrity":"/paper/2405.04437/integrity","json":"/paper/2405.04437/citation-record.json","paper":"/paper/2405.04437"},"outbound":[],"paper":{"arxiv_id":"2405.04437","last_updated":"2025-01-29T04:10:41Z","latest_version":3,"primary_category":"cs.LG","snapshot_observed_at":"2026-08-05T02:13:00.249260Z","submitted_at":"2024-05-07T16:00:32Z","title":"vAttention: Dynamic Memory Management for Serving LLMs without PagedAttention"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-05T06:32:48.257954+00:00","source":"crossref"},{"observed_at":"2026-08-05T06:32:44.755628+00:00","source":"retraction_watch"}],"thesis":"As of 5 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 14 inbound Pith citation observations for arXiv:2405.04437."}