{"as_of":"2026-08-04T18:46:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:364489beb2f8e6c986e56fffe1a8f373ad8d16980d4bb47de6883cf50c72146d","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":16,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":16,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-04T06:34:03.388597+00:00","state":"measured"},{"denominator":16,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":16,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-03T04:01:06.957435Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"pith","source_observed_at":"2026-07-10T06:15:00.866473Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2305.17118","last_updated":"2023-08-28T22:48:46Z","snapshot_observed_at":"2026-08-01T09:57:16.282788Z","submitted_at":"2023-05-26T17:39:58Z","title":"Scissorhands: Exploiting the Persistence of Importance Hypothesis for LLM KV Cache Compression at Test Time","version":2},"cited_work":{"arxiv_id":"2305.17118","doi":"10.48550/arxiv.2305.17118","metadata_source":"pith","pith_arxiv_id":"2305.17118","snapshot_observed_at":"2026-07-10T06:15:00.866473Z","title":"Scissorhands: Exploiting the persistence of importance hypothesis for llm kv cache compression at test time.arXiv preprint arXiv:2305.17118","venue":"cs.LG","work_id":"ff3209ef-b63a-4975-9d1d-ccab19932963","year":2023},"citing_paper":{"arxiv_id":"2310.01801","last_updated":"2024-10-29T18:26:09Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-03T05:17:08Z","title":"Model Tells You What to Discard: Adaptive KV Cache Compression for LLMs","version":4},"reference_index":88,"source":"arxiv_source","source_observed_at":"2026-05-17T11:11:21.460613Z"},"links":{"cited_paper":"/paper/2305.17118","citing_paper":"/paper/2310.01801"},"observation_digest":"sha256:16ad6f6e588741e094519a7a3af1e290facbb0afb6779853ed5f2a5f4368dd33","observation_id":"c178cfdb-18b7-4e47-8f63-e34dd649b11a","resolution":{"observed_at":"2026-05-17T11:11:21.567723Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.17118","last_updated":"2023-08-28T22:48:46Z","snapshot_observed_at":"2026-08-01T09:57:16.282788Z","submitted_at":"2023-05-26T17:39:58Z","title":"Scissorhands: Exploiting the Persistence of Importance Hypothesis for LLM KV Cache Compression at Test Time","version":2},"cited_work":{"arxiv_id":"2305.17118","doi":"10.48550/arxiv.2305.17118","metadata_source":"pith","pith_arxiv_id":"2305.17118","snapshot_observed_at":"2026-07-10T06:15:00.866473Z","title":"Scissorhands: Exploiting the persistence of importance hypothesis for llm kv cache compression at test time.arXiv preprint arXiv:2305.17118","venue":"cs.LG","work_id":"ff3209ef-b63a-4975-9d1d-ccab19932963","year":2023},"citing_paper":{"arxiv_id":"2604.17935","last_updated":"2026-04-20T08:15:17Z","snapshot_observed_at":"2026-07-06T23:04:55.190916Z","submitted_at":"2026-04-20T08:15:17Z","title":"How Much Cache Does Reasoning Need? Depth-Cache Tradeoffs in KV-Compressed Transformers","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-05-10T05:17:52.344313Z"},"links":{"cited_paper":"/paper/2305.17118","citing_paper":"/paper/2604.17935"},"observation_digest":"sha256:9cbe962bfdde91e60f8a9998ec749535ebe155863d77856958d3d2e2d2ff4bc3","observation_id":"cdcd808f-6c0c-4030-ad27-3f72bae25eff","resolution":{"observed_at":"2026-05-10T09:28:39.165223Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.17118","last_updated":"2023-08-28T22:48:46Z","snapshot_observed_at":"2026-08-01T09:57:16.282788Z","submitted_at":"2023-05-26T17:39:58Z","title":"Scissorhands: Exploiting the Persistence of Importance Hypothesis for LLM KV Cache Compression at Test Time","version":2},"cited_work":{"arxiv_id":"2305.17118","doi":"10.48550/arxiv.2305.17118","metadata_source":"pith","pith_arxiv_id":"2305.17118","snapshot_observed_at":"2026-07-10T06:15:00.866473Z","title":"Scissorhands: Exploiting the persistence of importance hypothesis for llm kv cache compression at test time.arXiv preprint arXiv:2305.17118","venue":"cs.LG","work_id":"ff3209ef-b63a-4975-9d1d-ccab19932963","year":2023},"citing_paper":{"arxiv_id":"2605.05219","last_updated":"2026-04-17T09:24:58Z","snapshot_observed_at":"2026-07-06T23:17:52.987860Z","submitted_at":"2026-04-17T09:24:58Z","title":"Sparse Prefix Caching for Hybrid and Recurrent LLM Serving","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-05-10T08:49:24.880528Z"},"links":{"cited_paper":"/paper/2305.17118","citing_paper":"/paper/2605.05219"},"observation_digest":"sha256:ac0b1adcd268f69307e12bcfc4fb4a53c3e1cbd1c493364e3c3ab6a1a402a536","observation_id":"484e545c-25c6-40b5-81fb-c203a0fa326a","resolution":{"observed_at":"2026-05-10T08:53:04.451224Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.17118","last_updated":"2023-08-28T22:48:46Z","snapshot_observed_at":"2026-08-01T09:57:16.282788Z","submitted_at":"2023-05-26T17:39:58Z","title":"Scissorhands: Exploiting the Persistence of Importance Hypothesis for LLM KV Cache Compression at Test Time","version":2},"cited_work":{"arxiv_id":"2305.17118","doi":"10.48550/arxiv.2305.17118","metadata_source":"pith","pith_arxiv_id":"2305.17118","snapshot_observed_at":"2026-07-10T06:15:00.866473Z","title":"Scissorhands: Exploiting the persistence of importance hypothesis for llm kv cache compression at test time.arXiv preprint arXiv:2305.17118","venue":"cs.LG","work_id":"ff3209ef-b63a-4975-9d1d-ccab19932963","year":2023},"citing_paper":{"arxiv_id":"2605.06763","last_updated":"2026-05-11T02:30:41Z","snapshot_observed_at":"2026-07-06T23:19:10.574595Z","submitted_at":"2026-05-07T17:37:56Z","title":"Sparse Attention as a Range Searching Problem: Towards an Inference-Efficient Index for KV Cache","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-05-11T01:15:35.871863Z"},"links":{"cited_paper":"/paper/2305.17118","citing_paper":"/paper/2605.06763"},"observation_digest":"sha256:2d37d1ba9e2f62fdb89ec6b2327efb7db2777b4a809e08b193cca12c505f6ac7","observation_id":"6dd109ad-40e5-482b-9813-25a5f50f0733","resolution":{"observed_at":"2026-05-11T01:15:50.543513Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.17118","last_updated":"2023-08-28T22:48:46Z","snapshot_observed_at":"2026-08-01T09:57:16.282788Z","submitted_at":"2023-05-26T17:39:58Z","title":"Scissorhands: Exploiting the Persistence of Importance Hypothesis for LLM KV Cache Compression at Test Time","version":2},"cited_work":{"arxiv_id":"2305.17118","doi":"10.48550/arxiv.2305.17118","metadata_source":"pith","pith_arxiv_id":"2305.17118","snapshot_observed_at":"2026-07-10T06:15:00.866473Z","title":"Scissorhands: Exploiting the persistence of importance hypothesis for llm kv cache compression at test time.arXiv preprint arXiv:2305.17118","venue":"cs.LG","work_id":"ff3209ef-b63a-4975-9d1d-ccab19932963","year":2023},"citing_paper":{"arxiv_id":"2605.09735","last_updated":"2026-06-30T03:39:11Z","snapshot_observed_at":"2026-07-06T23:21:49.499951Z","submitted_at":"2026-05-10T20:10:26Z","title":"KV-RM: Regularizing KV-Cache Movement for Static-Graph LLM Serving","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-05-12T03:57:30.104609Z"},"links":{"cited_paper":"/paper/2305.17118","citing_paper":"/paper/2605.09735"},"observation_digest":"sha256:7c8a1a87483eb2985dec3a2b07b5246c4aa3f066732f5354706a06fc41acb4b7","observation_id":"e22ec146-2c0f-4458-9d05-7af6d5906756","resolution":{"observed_at":"2026-05-12T06:51:27.079903Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.17118","last_updated":"2023-08-28T22:48:46Z","snapshot_observed_at":"2026-08-01T09:57:16.282788Z","submitted_at":"2023-05-26T17:39:58Z","title":"Scissorhands: Exploiting the Persistence of Importance Hypothesis for LLM KV Cache Compression at Test Time","version":2},"cited_work":{"arxiv_id":"2305.17118","doi":"10.48550/arxiv.2305.17118","metadata_source":"pith","pith_arxiv_id":"2305.17118","snapshot_observed_at":"2026-07-10T06:15:00.866473Z","title":"Scissorhands: Exploiting the persistence of importance hypothesis for llm kv cache compression at test time.arXiv preprint arXiv:2305.17118","venue":"cs.LG","work_id":"ff3209ef-b63a-4975-9d1d-ccab19932963","year":2023},"citing_paper":{"arxiv_id":"2605.09735","last_updated":"2026-06-30T03:39:11Z","snapshot_observed_at":"2026-07-06T23:21:49.499951Z","submitted_at":"2026-05-10T20:10:26Z","title":"KV-RM: Regularizing KV-Cache Movement for Static-Graph LLM Serving","version":2},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-07-01T08:05:44.256565Z"},"links":{"cited_paper":"/paper/2305.17118","citing_paper":"/paper/2605.09735"},"observation_digest":"sha256:26714e9863f4938afe4092e4605d8beed30beb24fffdf72c9cbad2cdd174b558","observation_id":"d3a1e9e6-3c28-48b2-917b-636e4341ee28","resolution":{"observed_at":"2026-07-01T08:15:32.435930Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.17118","last_updated":"2023-08-28T22:48:46Z","snapshot_observed_at":"2026-08-01T09:57:16.282788Z","submitted_at":"2023-05-26T17:39:58Z","title":"Scissorhands: Exploiting the Persistence of Importance Hypothesis for LLM KV Cache Compression at Test Time","version":2},"cited_work":{"arxiv_id":"2305.17118","doi":"10.48550/arxiv.2305.17118","metadata_source":"pith","pith_arxiv_id":"2305.17118","snapshot_observed_at":"2026-07-10T06:15:00.866473Z","title":"Scissorhands: Exploiting the persistence of importance hypothesis for llm kv cache compression at test time.arXiv preprint arXiv:2305.17118","venue":"cs.LG","work_id":"ff3209ef-b63a-4975-9d1d-ccab19932963","year":2023},"citing_paper":{"arxiv_id":"2605.14037","last_updated":"2026-05-13T18:58:16Z","snapshot_observed_at":"2026-08-04T00:53:16.855757Z","submitted_at":"2026-05-13T18:58:16Z","title":"Self-Pruned Key-Value Attention: Learning When to Write by Predicting Future Utility","version":1},"reference_index":70,"source":"arxiv_source","source_observed_at":"2026-05-15T05:35:09.705532Z"},"links":{"cited_paper":"/paper/2305.17118","citing_paper":"/paper/2605.14037"},"observation_digest":"sha256:875242b85fbaf7d7fabdbe408c941037f39e4018066a3189f2bb505d45ee6377","observation_id":"4c96ca8d-3847-4452-a6ed-7b4aea34c55d","resolution":{"observed_at":"2026-05-15T05:39:47.960219Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.17118","last_updated":"2023-08-28T22:48:46Z","snapshot_observed_at":"2026-08-01T09:57:16.282788Z","submitted_at":"2023-05-26T17:39:58Z","title":"Scissorhands: Exploiting the Persistence of Importance Hypothesis for LLM KV Cache Compression at Test Time","version":2},"cited_work":{"arxiv_id":"2305.17118","doi":"10.48550/arxiv.2305.17118","metadata_source":"pith","pith_arxiv_id":"2305.17118","snapshot_observed_at":"2026-07-10T06:15:00.866473Z","title":"Scissorhands: Exploiting the persistence of importance hypothesis for llm kv cache compression at test time.arXiv preprint arXiv:2305.17118","venue":"cs.LG","work_id":"ff3209ef-b63a-4975-9d1d-ccab19932963","year":2023},"citing_paper":{"arxiv_id":"2605.18053","last_updated":"2026-05-18T08:41:34Z","snapshot_observed_at":"2026-07-06T23:28:59.503300Z","submitted_at":"2026-05-18T08:41:34Z","title":"Protection Is (Nearly) All You Need: Structural Protection Dominates Scoring in Globally Capped KV Eviction","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-05-20T12:13:14.295659Z"},"links":{"cited_paper":"/paper/2305.17118","citing_paper":"/paper/2605.18053"},"observation_digest":"sha256:f4a0f71b46f4d667203d811a71bd8d124bacf37be5d4c59a3f9399f04e8c4216","observation_id":"3df71234-c3e9-459c-9142-3e2715476967","resolution":{"observed_at":"2026-05-20T12:13:15.990718Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.17118","last_updated":"2023-08-28T22:48:46Z","snapshot_observed_at":"2026-08-01T09:57:16.282788Z","submitted_at":"2023-05-26T17:39:58Z","title":"Scissorhands: Exploiting the Persistence of Importance Hypothesis for LLM KV Cache Compression at Test Time","version":2},"cited_work":{"arxiv_id":"2305.17118","doi":"10.48550/arxiv.2305.17118","metadata_source":"pith","pith_arxiv_id":"2305.17118","snapshot_observed_at":"2026-07-10T06:15:00.866473Z","title":"Scissorhands: Exploiting the persistence of importance hypothesis for llm kv cache compression at test time.arXiv preprint arXiv:2305.17118","venue":"cs.LG","work_id":"ff3209ef-b63a-4975-9d1d-ccab19932963","year":2023},"citing_paper":{"arxiv_id":"2605.22337","last_updated":"2026-05-23T16:54:41Z","snapshot_observed_at":"2026-07-06T23:32:39.761431Z","submitted_at":"2026-05-21T11:24:46Z","title":"Meta-Soft: Leveraging Composable Meta-Tokens for Context-Preserving KV Cache Compression","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-05-22T05:10:53.133346Z"},"links":{"cited_paper":"/paper/2305.17118","citing_paper":"/paper/2605.22337"},"observation_digest":"sha256:b660a8886f26426697f3df9e205a568cd26e0bf0a4d21d7061a52bcd73044e02","observation_id":"6e200117-5c5c-4bf6-928f-bb44c7bd0075","resolution":{"observed_at":"2026-05-22T05:11:06.173077Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.17118","last_updated":"2023-08-28T22:48:46Z","snapshot_observed_at":"2026-08-01T09:57:16.282788Z","submitted_at":"2023-05-26T17:39:58Z","title":"Scissorhands: Exploiting the Persistence of Importance Hypothesis for LLM KV Cache Compression at Test Time","version":2},"cited_work":{"arxiv_id":"2305.17118","doi":"10.48550/arxiv.2305.17118","metadata_source":"pith","pith_arxiv_id":"2305.17118","snapshot_observed_at":"2026-07-10T06:15:00.866473Z","title":"Scissorhands: Exploiting the persistence of importance hypothesis for llm kv cache compression at test time.arXiv preprint arXiv:2305.17118","venue":"cs.LG","work_id":"ff3209ef-b63a-4975-9d1d-ccab19932963","year":2023},"citing_paper":{"arxiv_id":"2605.22337","last_updated":"2026-05-23T16:54:41Z","snapshot_observed_at":"2026-07-06T23:32:39.761431Z","submitted_at":"2026-05-21T11:24:46Z","title":"Meta-Soft: Leveraging Composable Meta-Tokens for Context-Preserving KV Cache Compression","version":2},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-06-30T17:28:08.750049Z"},"links":{"cited_paper":"/paper/2305.17118","citing_paper":"/paper/2605.22337"},"observation_digest":"sha256:f8f77401870d439ce846430b44b648a3a1a69f972581fb468855476a8ab955b6","observation_id":"2e03c416-d9b1-44d5-a6ae-60e3d9206e92","resolution":{"observed_at":"2026-06-30T17:34:57.846036Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.17118","last_updated":"2023-08-28T22:48:46Z","snapshot_observed_at":"2026-08-01T09:57:16.282788Z","submitted_at":"2023-05-26T17:39:58Z","title":"Scissorhands: Exploiting the Persistence of Importance Hypothesis for LLM KV Cache Compression at Test Time","version":2},"cited_work":{"arxiv_id":"2305.17118","doi":"10.48550/arxiv.2305.17118","metadata_source":"pith","pith_arxiv_id":"2305.17118","snapshot_observed_at":"2026-07-10T06:15:00.866473Z","title":"Scissorhands: Exploiting the persistence of importance hypothesis for llm kv cache compression at test time.arXiv preprint arXiv:2305.17118","venue":"cs.LG","work_id":"ff3209ef-b63a-4975-9d1d-ccab19932963","year":2023},"citing_paper":{"arxiv_id":"2606.20295","last_updated":"2026-07-24T07:54:17Z","snapshot_observed_at":"2026-08-02T10:48:55.986116Z","submitted_at":"2026-06-18T14:33:09Z","title":"Token-Operations-Oriented Inference Optimization Techniques for Large Models","version":1},"reference_index":118,"source":"pdf_text","source_observed_at":"2026-06-26T16:15:22.543601Z"},"links":{"cited_paper":"/paper/2305.17118","citing_paper":"/paper/2606.20295"},"observation_digest":"sha256:1c41ebb04611aaea827253f2e7c2e809b56cca6a92e05ddecdccd867afe64be1","observation_id":"3543cda0-796d-4076-90b9-766b7a4b438b","resolution":{"observed_at":"2026-07-04T05:09:36.729583Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.17118","last_updated":"2023-08-28T22:48:46Z","snapshot_observed_at":"2026-08-01T09:57:16.282788Z","submitted_at":"2023-05-26T17:39:58Z","title":"Scissorhands: Exploiting the Persistence of Importance Hypothesis for LLM KV Cache Compression at Test Time","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.17118","snapshot_observed_at":"2026-08-02T10:49:11.815517Z","title":"Scissorhands: Exploiting the Persistence of Importance Hypothesis for LLM KV Cache Compression at Test Time.Advances in Neural Information Processing Systems, 36:52342–52364, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2606.20295","last_updated":"2026-07-24T07:54:17Z","snapshot_observed_at":"2026-08-02T10:48:55.986116Z","submitted_at":"2026-06-18T14:33:09Z","title":"Token-Operations-Oriented Inference Optimization Techniques for Large Models","version":2},"reference_index":109,"source":"pdf_text","source_observed_at":"2026-08-02T10:49:11.815517Z"},"links":{"cited_paper":"/paper/2305.17118","citing_paper":"/paper/2606.20295"},"observation_digest":"sha256:3d46b7903bc42e84b2c4d50776d56bd54409d476bd1b3fe657b6dc05fbdbd6e3","observation_id":"1bc2f9f6-7465-45bc-acf9-f4e20ef1245c","resolution":{"observed_at":"2026-08-02T10:49:11.815517Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.17118","last_updated":"2023-08-28T22:48:46Z","snapshot_observed_at":"2026-08-01T09:57:16.282788Z","submitted_at":"2023-05-26T17:39:58Z","title":"Scissorhands: Exploiting the Persistence of Importance Hypothesis for LLM KV Cache Compression at Test Time","version":2},"cited_work":{"arxiv_id":"2305.17118","doi":"10.48550/arxiv.2305.17118","metadata_source":"pith","pith_arxiv_id":"2305.17118","snapshot_observed_at":"2026-07-10T06:15:00.866473Z","title":"Scissorhands: Exploiting the persistence of importance hypothesis for llm kv cache compression at test time.arXiv preprint arXiv:2305.17118","venue":"cs.LG","work_id":"ff3209ef-b63a-4975-9d1d-ccab19932963","year":2023},"citing_paper":{"arxiv_id":"2606.24033","last_updated":"2026-06-23T00:17:48Z","snapshot_observed_at":"2026-08-02T14:58:46.342813Z","submitted_at":"2026-06-23T00:17:48Z","title":"RoPE-Aware Bit Allocation for KV-Cache Quantization","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-06-26T01:05:13.790381Z"},"links":{"cited_paper":"/paper/2305.17118","citing_paper":"/paper/2606.24033"},"observation_digest":"sha256:3436c356f929f56cc18fee454511e23e334760be41f1b41e9809707020eb3553","observation_id":"f76728fe-3fbc-42b7-aa8c-8533340d74d4","resolution":{"observed_at":"2026-07-04T15:59:57.363534Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.17118","last_updated":"2023-08-28T22:48:46Z","snapshot_observed_at":"2026-08-01T09:57:16.282788Z","submitted_at":"2023-05-26T17:39:58Z","title":"Scissorhands: Exploiting the Persistence of Importance Hypothesis for LLM KV Cache Compression at Test Time","version":2},"cited_work":{"arxiv_id":"2305.17118","doi":"10.48550/arxiv.2305.17118","metadata_source":"pith","pith_arxiv_id":"2305.17118","snapshot_observed_at":"2026-07-10T06:15:00.866473Z","title":"Scissorhands: Exploiting the persistence of importance hypothesis for llm kv cache compression at test time.arXiv preprint arXiv:2305.17118","venue":"cs.LG","work_id":"ff3209ef-b63a-4975-9d1d-ccab19932963","year":2023},"citing_paper":{"arxiv_id":"2607.01065","last_updated":"2026-07-01T15:25:21Z","snapshot_observed_at":"2026-07-07T00:06:39.908636Z","submitted_at":"2026-07-01T15:25:21Z","title":"GSRQ: Gain-Shape Residual Quantization for Sub-1-bit KV Cache","version":1},"reference_index":2,"source":"arxiv_source","source_observed_at":"2026-07-02T15:55:40.177742Z"},"links":{"cited_paper":"/paper/2305.17118","citing_paper":"/paper/2607.01065"},"observation_digest":"sha256:5aeae94142947a9a3d5a288bad1672434be0b671ed971fcfc9ccf455e1090332","observation_id":"1d340047-4201-478b-8df8-4f35f725eb73","resolution":{"observed_at":"2026-07-02T15:57:06.386693Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.17118","last_updated":"2023-08-28T22:48:46Z","snapshot_observed_at":"2026-08-01T09:57:16.282788Z","submitted_at":"2023-05-26T17:39:58Z","title":"Scissorhands: Exploiting the Persistence of Importance Hypothesis for LLM KV Cache Compression at Test Time","version":2},"cited_work":{"arxiv_id":"2305.17118","doi":"10.48550/arxiv.2305.17118","metadata_source":"pith","pith_arxiv_id":"2305.17118","snapshot_observed_at":"2026-07-10T06:15:00.866473Z","title":"Scissorhands: Exploiting the persistence of importance hypothesis for llm kv cache compression at test time.arXiv preprint arXiv:2305.17118","venue":"cs.LG","work_id":"ff3209ef-b63a-4975-9d1d-ccab19932963","year":2023},"citing_paper":{"arxiv_id":"2607.08032","last_updated":"2026-07-09T01:15:03Z","snapshot_observed_at":"2026-08-02T16:38:21.894642Z","submitted_at":"2026-07-09T01:15:03Z","title":"What to Keep, What to Forget: A Rate--Distortion View of Memory Compaction in LLMs and Agents","version":1},"reference_index":75,"source":"pdf_text","source_observed_at":"2026-07-10T01:26:59.421158Z"},"links":{"cited_paper":"/paper/2305.17118","citing_paper":"/paper/2607.08032"},"observation_digest":"sha256:aeb507b24bc41db1390d9ca1132cb83c873a1f046002460001bb00233b9e0cc4","observation_id":"0cd2b41a-325f-473d-9382-7064ffa5a2ad","resolution":{"observed_at":"2026-07-10T01:36:44.244964Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.17118","last_updated":"2023-08-28T22:48:46Z","snapshot_observed_at":"2026-08-01T09:57:16.282788Z","submitted_at":"2023-05-26T17:39:58Z","title":"Scissorhands: Exploiting the Persistence of Importance Hypothesis for LLM KV Cache Compression at Test Time","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.17118","snapshot_observed_at":"2026-08-03T04:01:06.957435Z","title":"Zirui Liu, Jiayi Yuan, Hongye Jin, Shaochen Zhong, Zhaozhuo Xu, Vladimir Braverman, Beidi Chen, and Xia Hu","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.29591","last_updated":"2026-07-31T16:16:45Z","snapshot_observed_at":"2026-08-04T18:23:28.095912Z","submitted_at":"2026-07-31T16:16:45Z","title":"ResKV: Reconstructing Omitted Attention Contributions for Fixed-Budget KV Cache Compression","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-03T04:01:06.957435Z"},"links":{"cited_paper":"/paper/2305.17118","citing_paper":"/paper/2607.29591"},"observation_digest":"sha256:14e64783d5a431fe197bf2a8ec0cf6d7849563a6b315edcb9ac6eb58528c895c","observation_id":"17a8bfe0-40ed-42a5-ac19-8c2c133a793b","resolution":{"observed_at":"2026-08-03T04:01:06.957435Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2305.17118/citation-record","integrity":"/paper/2305.17118/integrity","json":"/paper/2305.17118/citation-record.json","paper":"/paper/2305.17118"},"outbound":[],"paper":{"arxiv_id":"2305.17118","last_updated":"2023-08-28T22:48:46Z","latest_version":2,"primary_category":"cs.LG","snapshot_observed_at":"2026-08-01T09:57:16.282788Z","submitted_at":"2023-05-26T17:39:58Z","title":"Scissorhands: Exploiting the Persistence of Importance Hypothesis for LLM KV Cache Compression at Test Time"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"thesis":"As of 4 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 16 inbound Pith citation observations for arXiv:2305.17118."}