{"bundle_type":"pith_open_graph_bundle","bundle_version":"1.0","pith_number":"pith:2026:7EWU5DCQS5YPYOCE7L36SQEXQI","short_pith_number":"pith:7EWU5DCQ","canonical_record":{"source":{"id":"2605.18825","kind":"arxiv","version":1},"metadata":{"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.LG","submitted_at":"2026-05-12T18:38:24Z","cross_cats_sorted":["cs.DC"],"title_canon_sha256":"59d82b2a6bf1a27059242039985132d8b35b7c124db736a816f9308478de5b40","abstract_canon_sha256":"277766eb94fa49476b74fc7108c1fb2215bfef9328a1c423c77ae9b0ef59a0ef"},"schema_version":"1.0"},"canonical_sha256":"f92d4e8c509770fc3844faf7e94097822940271b55ad9237c2ca9e76677e2fdb","source":{"kind":"arxiv","id":"2605.18825","version":1},"source_aliases":[{"alias_kind":"arxiv","alias_value":"2605.18825","created_at":"2026-05-20T00:06:24Z"},{"alias_kind":"arxiv_version","alias_value":"2605.18825v1","created_at":"2026-05-20T00:06:24Z"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2605.18825","created_at":"2026-05-20T00:06:24Z"},{"alias_kind":"pith_short_12","alias_value":"7EWU5DCQS5YP","created_at":"2026-05-20T00:06:24Z"},{"alias_kind":"pith_short_16","alias_value":"7EWU5DCQS5YPYOCE","created_at":"2026-05-20T00:06:24Z"},{"alias_kind":"pith_short_8","alias_value":"7EWU5DCQ","created_at":"2026-05-20T00:06:24Z"}],"events":[{"event_type":"record_created","subject_pith_number":"pith:2026:7EWU5DCQS5YPYOCE7L36SQEXQI","target":"record","payload":{"canonical_record":{"source":{"id":"2605.18825","kind":"arxiv","version":1},"metadata":{"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.LG","submitted_at":"2026-05-12T18:38:24Z","cross_cats_sorted":["cs.DC"],"title_canon_sha256":"59d82b2a6bf1a27059242039985132d8b35b7c124db736a816f9308478de5b40","abstract_canon_sha256":"277766eb94fa49476b74fc7108c1fb2215bfef9328a1c423c77ae9b0ef59a0ef"},"schema_version":"1.0"},"canonical_sha256":"f92d4e8c509770fc3844faf7e94097822940271b55ad9237c2ca9e76677e2fdb","receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-05-20T00:06:24.594131Z","signature_b64":"9bzAv2yaggato1ZBEnK1HT9J5e4MwoERNMXgXdK38pbZ0XehOw8YDzfjBq+M2YaOPIsi8jHNSycYP75mzroCAg==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"f92d4e8c509770fc3844faf7e94097822940271b55ad9237c2ca9e76677e2fdb","last_reissued_at":"2026-05-20T00:06:24.593473Z","signature_status":"signed_v1","first_computed_at":"2026-05-20T00:06:24.593473Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"source_kind":"arxiv","source_id":"2605.18825","source_version":1,"attestation_state":"computed"},"signer":{"signer_id":"pith.science","signer_type":"pith_registry","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"created_at":"2026-05-20T00:06:24Z","supersedes":[],"prev_event":null,"signature":{"signature_status":"signed_v1","algorithm":"ed25519","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signature_b64":"OerBf1wFfFm7p101Lk7C6nOfUTXCWKeLsR2rGys+CbyhVUnl/2QHHjI5xskuY/5lTW1QHh2DTJfmgbuoodjmBw==","signed_message":"open_graph_event_sha256_bytes","signed_at":"2026-05-26T01:08:38.370427Z"},"content_sha256":"b1d8456f3008e24cc8a11583df5b36eeb795e8c8c6a627bd95fd2706c22a0dad","schema_version":"1.0","event_id":"sha256:b1d8456f3008e24cc8a11583df5b36eeb795e8c8c6a627bd95fd2706c22a0dad"},{"event_type":"graph_snapshot","subject_pith_number":"pith:2026:7EWU5DCQS5YPYOCE7L36SQEXQI","target":"graph","payload":{"graph_snapshot":{"paper":{"title":"Not All Tokens Are Worth Caching: Learning Semantic-Aware Eviction for LLM Prefix Caches","license":"http://creativecommons.org/licenses/by/4.0/","headline":"","cross_cats":["cs.DC"],"primary_cat":"cs.LG","authors_text":"Jiatong Ji, Qingsong Liu, Ruizhi Pu, Shaoke Fang, Wenfei Wu, Ziang Li","submitted_at":"2026-05-12T18:38:24Z","abstract_excerpt":"Prefix caching is a key optimization in Large Language Model (LLM) serving, reusing attention Key-Value (KV) states across requests with shared prompt prefixes to reduce expensive prefill computation. However, its benefit depends critically on the eviction policy as GPU memory is scarce, and existing policies such as LRU largely treat cached blocks uniformly. This view ignores a fundamental property of LLM prompts: not all tokens are equally worth caching. We show that different token types within a prompt, including system prompts, user queries, tool outputs, model responses, and chain-of-tho"},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2605.18825","kind":"arxiv","version":1},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2605.18825/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"verdict_id":null},"signer":{"signer_id":"pith.science","signer_type":"pith_registry","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"created_at":"2026-05-20T00:06:24Z","supersedes":[],"prev_event":null,"signature":{"signature_status":"signed_v1","algorithm":"ed25519","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signature_b64":"2tBsOLYWdk8J2EbivruitqujzL2WjutX8PhqB9OytiNICrF5P7MOc6iR83Ccr4PkYFztwjTvQPRBq4i9pqS2DQ==","signed_message":"open_graph_event_sha256_bytes","signed_at":"2026-05-26T01:08:38.371086Z"},"content_sha256":"a8f30b28410f5c499037595e59aa6b82df36dd54fa671ce18a3e3d833afc7dfb","schema_version":"1.0","event_id":"sha256:a8f30b28410f5c499037595e59aa6b82df36dd54fa671ce18a3e3d833afc7dfb"}],"timestamp_proofs":[],"mirror_hints":[{"mirror_type":"https","name":"Pith Resolver","base_url":"https://pith.science","bundle_url":"https://pith.science/pith/7EWU5DCQS5YPYOCE7L36SQEXQI/bundle.json","state_url":"https://pith.science/pith/7EWU5DCQS5YPYOCE7L36SQEXQI/state.json","well_known_bundle_url":"https://pith.science/.well-known/pith/7EWU5DCQS5YPYOCE7L36SQEXQI/bundle.json","status":"primary"}],"public_keys":[{"key_id":"pith-v1-2026-05","algorithm":"ed25519","format":"raw","public_key_b64":"stVStoiQhXFxp4s2pdzPNoqVNBMojDU/fJ2db5S3CbM=","public_key_hex":"b2d552b68890857171a78b36a5dccf368a953413288c353f7c9d9d6f94b709b3","fingerprint_sha256_b32_first128bits":"RVFV5Z2OI2J3ZUO7ERDEBCYNKS","fingerprint_sha256_hex":"8d4b5ee74e4693bcd1df2446408b0d54","rotates_at":null,"url":"https://pith.science/pith-signing-key.json","notes":"Pith uses this Ed25519 key to sign canonical record SHA-256 digests. Verify with: ed25519_verify(public_key, message=canonical_sha256_bytes, signature=base64decode(signature_b64))."}],"merge_version":"pith-open-graph-merge-v1","built_at":"2026-05-26T01:08:38Z","links":{"resolver":"https://pith.science/pith/7EWU5DCQS5YPYOCE7L36SQEXQI","bundle":"https://pith.science/pith/7EWU5DCQS5YPYOCE7L36SQEXQI/bundle.json","state":"https://pith.science/pith/7EWU5DCQS5YPYOCE7L36SQEXQI/state.json","well_known_bundle":"https://pith.science/.well-known/pith/7EWU5DCQS5YPYOCE7L36SQEXQI/bundle.json"},"state":{"state_type":"pith_open_graph_state","state_version":"1.0","pith_number":"pith:2026:7EWU5DCQS5YPYOCE7L36SQEXQI","merge_version":"pith-open-graph-merge-v1","event_count":2,"valid_event_count":2,"invalid_event_count":0,"equivocation_count":0,"current":{"canonical_record":{"metadata":{"abstract_canon_sha256":"277766eb94fa49476b74fc7108c1fb2215bfef9328a1c423c77ae9b0ef59a0ef","cross_cats_sorted":["cs.DC"],"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.LG","submitted_at":"2026-05-12T18:38:24Z","title_canon_sha256":"59d82b2a6bf1a27059242039985132d8b35b7c124db736a816f9308478de5b40"},"schema_version":"1.0","source":{"id":"2605.18825","kind":"arxiv","version":1}},"source_aliases":[{"alias_kind":"arxiv","alias_value":"2605.18825","created_at":"2026-05-20T00:06:24Z"},{"alias_kind":"arxiv_version","alias_value":"2605.18825v1","created_at":"2026-05-20T00:06:24Z"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2605.18825","created_at":"2026-05-20T00:06:24Z"},{"alias_kind":"pith_short_12","alias_value":"7EWU5DCQS5YP","created_at":"2026-05-20T00:06:24Z"},{"alias_kind":"pith_short_16","alias_value":"7EWU5DCQS5YPYOCE","created_at":"2026-05-20T00:06:24Z"},{"alias_kind":"pith_short_8","alias_value":"7EWU5DCQ","created_at":"2026-05-20T00:06:24Z"}],"graph_snapshots":[{"event_id":"sha256:a8f30b28410f5c499037595e59aa6b82df36dd54fa671ce18a3e3d833afc7dfb","target":"graph","created_at":"2026-05-20T00:06:24Z","signer":{"key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signer_id":"pith.science","signer_type":"pith_registry"},"payload":{"graph_snapshot":{"author_claims":{"count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","strong_count":0},"builder_version":"pith-number-builder-2026-05-17-v1","claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"integrity":{"available":true,"clean":true,"detectors_run":[],"endpoint":"/pith/2605.18825/integrity.json","findings":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938","summary":{"advisory":0,"by_detector":{},"critical":0,"informational":0}},"paper":{"abstract_excerpt":"Prefix caching is a key optimization in Large Language Model (LLM) serving, reusing attention Key-Value (KV) states across requests with shared prompt prefixes to reduce expensive prefill computation. However, its benefit depends critically on the eviction policy as GPU memory is scarce, and existing policies such as LRU largely treat cached blocks uniformly. This view ignores a fundamental property of LLM prompts: not all tokens are equally worth caching. We show that different token types within a prompt, including system prompts, user queries, tool outputs, model responses, and chain-of-tho","authors_text":"Jiatong Ji, Qingsong Liu, Ruizhi Pu, Shaoke Fang, Wenfei Wu, Ziang Li","cross_cats":["cs.DC"],"headline":"","license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.LG","submitted_at":"2026-05-12T18:38:24Z","title":"Not All Tokens Are Worth Caching: Learning Semantic-Aware Eviction for LLM Prefix Caches"},"references":{"count":0,"internal_anchors":0,"resolved_work":0,"sample":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2605.18825","kind":"arxiv","version":1},"verdict":{"created_at":null,"id":null,"model_set":{},"one_line_summary":"","pipeline_version":null,"pith_extraction_headline":"","strongest_claim":"","weakest_assumption":""}},"verdict_id":null}}],"author_attestations":[],"timestamp_anchors":[],"storage_attestations":[],"citation_signatures":[],"replication_records":[],"corrections":[],"mirror_hints":[],"record_created":{"event_id":"sha256:b1d8456f3008e24cc8a11583df5b36eeb795e8c8c6a627bd95fd2706c22a0dad","target":"record","created_at":"2026-05-20T00:06:24Z","signer":{"key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signer_id":"pith.science","signer_type":"pith_registry"},"payload":{"attestation_state":"computed","canonical_record":{"metadata":{"abstract_canon_sha256":"277766eb94fa49476b74fc7108c1fb2215bfef9328a1c423c77ae9b0ef59a0ef","cross_cats_sorted":["cs.DC"],"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.LG","submitted_at":"2026-05-12T18:38:24Z","title_canon_sha256":"59d82b2a6bf1a27059242039985132d8b35b7c124db736a816f9308478de5b40"},"schema_version":"1.0","source":{"id":"2605.18825","kind":"arxiv","version":1}},"canonical_sha256":"f92d4e8c509770fc3844faf7e94097822940271b55ad9237c2ca9e76677e2fdb","receipt":{"algorithm":"ed25519","builder_version":"pith-number-builder-2026-05-17-v1","canonical_sha256":"f92d4e8c509770fc3844faf7e94097822940271b55ad9237c2ca9e76677e2fdb","first_computed_at":"2026-05-20T00:06:24.593473Z","key_id":"pith-v1-2026-05","kind":"pith_receipt","last_reissued_at":"2026-05-20T00:06:24.593473Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","receipt_version":"0.3","signature_b64":"9bzAv2yaggato1ZBEnK1HT9J5e4MwoERNMXgXdK38pbZ0XehOw8YDzfjBq+M2YaOPIsi8jHNSycYP75mzroCAg==","signature_status":"signed_v1","signed_at":"2026-05-20T00:06:24.594131Z","signed_message":"canonical_sha256_bytes"},"source_id":"2605.18825","source_kind":"arxiv","source_version":1}}},"equivocations":[],"invalid_events":[],"applied_event_ids":["sha256:b1d8456f3008e24cc8a11583df5b36eeb795e8c8c6a627bd95fd2706c22a0dad","sha256:a8f30b28410f5c499037595e59aa6b82df36dd54fa671ce18a3e3d833afc7dfb"],"state_sha256":"765858bc4aa88c8a8235d76a62da871054e1d9cdf52f4ac4e6a7aa480ec7cbd0"},"bundle_signature":{"signature_status":"signed_v1","algorithm":"ed25519","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signature_b64":"dWDYwzIDD5ach9iD+OhPm+B9QCvFotbFh04mNjrgwTeF5CbuXvSH96PISYbiINYYi5CEbxfc3/q9z0CyeUmWCw==","signed_message":"bundle_sha256_bytes","signed_at":"2026-05-26T01:08:38.374800Z","bundle_sha256":"cb5f0856a966bd88f3ca5dccb4083f4687e908a5b90b248f923ce7a248825297"}}