{"as_of":"2026-08-13T13:29:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:476d646aaa91c87804b3f8ec72f7dd43bdc4771caef1fd76cea1f0bb5cdb94f8","coverage":[{"denominator":51,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":51,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-12T15:58:48.247129Z","state":"measured"},{"denominator":51,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":51,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-13T06:32:02.005865+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2411.13820/citation-record","integrity":"/paper/2411.13820/integrity","json":"/paper/2411.13820/citation-record.json","paper":"/paper/2411.13820"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2112.00861","last_updated":"2021-12-09T21:40:22Z","snapshot_observed_at":"2026-08-10T12:03:10.373653Z","submitted_at":"2021-12-01T22:24:34Z","title":"A General Language Assistant as a Laboratory for Alignment","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2112.00861","snapshot_observed_at":"2026-08-12T15:58:48.050699Z","title":"Brown, Jack Clark, Sam McCandlish, Chris Olah, and Jared Kaplan","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-12T15:48:15.716052Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.050699Z"},"links":{"cited_paper":"/paper/2112.00861","citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:723885fe118987183993350da5103c36127bca75c1b7bd18e95c23f5b9248dfb","observation_id":"2abab399-c656-4168-8eb6-e8a656bf5c3b","resolution":{"observed_at":"2026-08-12T15:58:48.050699Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.871744Z","title":"Baeza-Yates and Felipe Saint-Jean","venue":null,"work_id":"04f26d6f-eb19-4676-8bfc-54e6fe9cdeac","year":2003},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-12T15:48:15.716052Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.055211Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:948658c35abd320b5f4c8bd028636c839dfe6d880f19349e5fdcb5e04405de99","observation_id":"344203b4-e725-4fb2-90e0-810cc47893bc","resolution":{"observed_at":"2026-08-12T15:58:48.876060Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.855128Z","title":"GPTCache: An open-source semantic cache for LLM applications enabling faster answers and cost savings","venue":null,"work_id":"57d79c06-5c3a-4c24-9f1c-3fa4d7a00a2e","year":2023},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-12T15:48:15.716052Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.059033Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:6c6f13ba8c1269078f404e3c9c8bd50898674abe9321467779ab50fe1a6f8b55","observation_id":"22190558-f55d-49da-9dde-c0889ff741ab","resolution":{"observed_at":"2026-08-12T15:58:48.861489Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.833310Z","title":"Jordan, Joseph E","venue":null,"work_id":"50a783dd-ab77-44cc-85f3-c734a9a84e10","year":2024},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-12T15:48:15.716052Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.067059Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:1c489f577e8a6b729596ee645aa705eec9485f376e1d5c0856a5d0e6ba9d0b09","observation_id":"27d6ce61-9612-489f-a787-f7c86fa94047","resolution":{"observed_at":"2026-08-12T15:58:48.838176Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2405.04434","last_updated":"2024-06-19T06:04:17Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-05-07T15:56:43Z","title":"DeepSeek-V2: A Strong, Economical, and Efficient Mixture-of-Experts Language Model","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.04434","snapshot_observed_at":"2026-08-12T15:58:48.071378Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-12T15:48:15.716052Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.071378Z"},"links":{"cited_paper":"/paper/2405.04434","citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:a71316deec01bcccb3ffbb7d237fcd07fcd9a89e2517331033f17ac3d52b6003","observation_id":"c1e4f6b5-10d6-4f12-8071-c0bb3b1d3877","resolution":{"observed_at":"2026-08-12T15:58:48.071378Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2412.19437","last_updated":"2025-02-18T17:26:38Z","snapshot_observed_at":"2026-08-11T01:48:59.557045Z","submitted_at":"2024-12-27T04:03:16Z","title":"DeepSeek-V3 Technical Report","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.19437","snapshot_observed_at":"2026-08-12T15:58:48.075969Z","title":"Zhang, Han Bao, Hanwei Xu, Haocheng Wang, Haowei Zhang, Honghui Ding, Huajian Xin, Huazuo Gao, Hui Li, Hui Qu, J","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-12T15:48:15.716052Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.075969Z"},"links":{"cited_paper":"/paper/2412.19437","citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:4df270082f4884a204111b5b85b1f2e41b57267fd9444d3e1db7a654edd4e5ce","observation_id":"5328728b-8302-47b7-8681-5b26f7cb8a89","resolution":{"observed_at":"2026-08-12T15:58:48.075969Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.21783","last_updated":"2024-11-23T23:27:33Z","snapshot_observed_at":"2026-08-10T16:40:37.411115Z","submitted_at":"2024-07-31T17:54:27Z","title":"The Llama 3 Herd of Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.21783","snapshot_observed_at":"2026-08-12T15:58:48.080782Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-12T15:48:15.716052Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.080782Z"},"links":{"cited_paper":"/paper/2407.21783","citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:44842d0e0ace44b7a8fe099616817effa8014648042aaf4e2d2f2ef0315729ba","observation_id":"6d5c70ba-d067-41ca-92d7-61f3448ccb4b","resolution":{"observed_at":"2026-08-12T15:58:48.080782Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.822038Z","title":"Boosting the perfor- mance of web search engines: Caching and prefetching query results by exploiting historical usage data","venue":null,"work_id":"d36a7c50-f895-4003-88a0-967fb87df439","year":2006},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-12T15:48:15.716052Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.084979Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:cdf0a5a9a5c8e1e6fbffdbadae239c532b7acf84f8b75516bf31a9245b340cc7","observation_id":"b1ff2bd5-0485-4dbd-9abb-efdf031817d6","resolution":{"observed_at":"2026-08-12T15:58:48.825813Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2312.10997","last_updated":"2024-03-27T09:16:57Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-12-18T07:47:33Z","title":"Retrieval-Augmented Generation for Large Language Models: A Survey","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.10997","snapshot_observed_at":"2026-08-12T15:58:48.092990Z","title":"Retrieval-augmented generation for large language models: A survey","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-12T15:48:15.716052Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.092990Z"},"links":{"cited_paper":"/paper/2312.10997","citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:af1117eb82c12f40db519c1584b281375c1bb67997e62390d0628397b8227971","observation_id":"5ea15936-7068-4813-a226-d02a1d076ed5","resolution":{"observed_at":"2026-08-12T15:58:48.092990Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2403.19708","last_updated":"2024-06-30T23:50:38Z","snapshot_observed_at":"2026-08-13T00:47:23.092746Z","submitted_at":"2024-03-23T10:42:49Z","title":"Cost-Efficient Large Language Model Serving for Multi-turn Conversations with CachedAttention","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.19708","snapshot_observed_at":"2026-08-12T15:58:48.097190Z","title":"Privacy-aware semantic cache for large language models","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-12T15:48:15.716052Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.097190Z"},"links":{"cited_paper":"/paper/2403.19708","citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:8ef77401349b08c4eb9587293e3cb498f86abbacb90b63ff35a2bd2de959b554","observation_id":"4cbadb51-9cac-4a66-a3be-15320991c338","resolution":{"observed_at":"2026-08-12T15:58:48.097190Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.101011Z","title":"Efficient memory management for large language model serving with pagedattention","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-12T15:48:15.716052Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.101011Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:69405deb4a6d32a2161aa7df67f6a483578fcd4b344342018af7c3df847b385c","observation_id":"a0f343f8-9ec2-40bc-9b43-b15ea0c5aeb5","resolution":{"observed_at":"2026-08-12T15:58:48.101011Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.802857Z","title":"Predictive caching and prefetching of query results in search engines","venue":null,"work_id":"01ed570a-397d-4989-849e-f4dce2b4ce55","year":2003},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-12T15:48:15.716052Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.104821Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:926be1da2a12531386c77e879d9c0ea42c36baca0a97535534f770d91e77fb15","observation_id":"545658c7-f2e4-4a1f-9c57-469b91ab1e67","resolution":{"observed_at":"2026-08-12T15:58:48.807271Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.790891Z","title":"von Riedemann, Cong Zhang, and Jiangchuan Liu","venue":null,"work_id":"f1777ceb-71d6-43c1-aca9-441611b4dead","year":2024},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-12T15:48:15.716052Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.108549Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:37a52feea8af34ba2b8f00d08a89bcdbe5b4468d57321a0e8973a873fd1aa3ae","observation_id":"e45bdd50-287f-41dd-97dd-ee74c5ef7a9e","resolution":{"observed_at":"2026-08-12T15:58:48.794791Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.778957Z","title":"Three-level caching for efficient query processing in large web search engines","venue":null,"work_id":"ac19456f-3840-4415-bf16-69b6ee514673","year":2005},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-12T15:48:15.716052Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.112557Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:d9a4258912c6539da711071d2f3ab58380d41be99b7f26cf7d306a9a8d4412ad","observation_id":"a52ced6d-fa0a-4d03-a700-4390a44dc713","resolution":{"observed_at":"2026-08-12T15:58:48.782959Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.767230Z","title":"New caching techniques for web search engines","venue":null,"work_id":"9ced39de-3789-4aeb-91c6-316aee3a934f","year":2010},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-12T15:48:15.716052Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.116208Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:690bdf432608a3c00fdda00a438d0040b2b81d100a5da7fc3e28a8b9565b2a69","observation_id":"5bfe0ea8-678b-4b98-bbea-1935b14ce1a2","resolution":{"observed_at":"2026-08-12T15:58:48.770863Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.756960Z","title":"Markatos","venue":null,"work_id":"b0ce5fb4-f531-4c17-a900-e24f85be6e28","year":2001},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-12T15:48:15.716052Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.119909Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:a14d3554ec0e20f7ff0882c26766d5a66b39de0237ed7b5c11f8d14d4af839db","observation_id":"928c0530-b343-4312-887c-68dbea1d6cdf","resolution":{"observed_at":"2026-08-12T15:58:48.760554Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.746338Z","title":"Context-based semantic caching for llm applications","venue":null,"work_id":"f14f00bc-4f28-4c81-be99-a967e4afe93b","year":2024},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-12T15:48:15.716052Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.123660Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:bb6e5e3b8c4758f1cdbdf6e9f3542d793cd2b9767f44e70370b82b5a8aecdaab","observation_id":"93294676-7c33-42d1-a25a-d542b78a2ac4","resolution":{"observed_at":"2026-08-12T15:58:48.750157Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.735318Z","title":"Openai chatgpt, 2022","venue":null,"work_id":"b0cf169f-2616-47b8-a148-fdd22fd9d1bd","year":2022},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-12T15:48:15.716052Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.127561Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:1f23271fb1a0f279b2f3dfcef4545905684d1535d375a02e85ae417771b85397","observation_id":"67f5381b-d07a-4041-903e-aac89dae261a","resolution":{"observed_at":"2026-08-12T15:58:48.739502Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2401.08138","last_updated":"2024-01-16T06:16:33Z","snapshot_observed_at":"2026-08-13T04:42:17.285249Z","submitted_at":"2024-01-16T06:16:33Z","title":"LLMs for Test Input Generation for Semantic Caches","version":1},"cited_work":{"arxiv_id":"2401.08138","doi":null,"metadata_source":"pith","pith_arxiv_id":"2401.08138","snapshot_observed_at":"2026-08-12T15:58:48.354508Z","title":"LLMs for Test Input Generation for Semantic Caches","venue":"cs.SE","work_id":"076de712-c5f2-47e3-bffd-b3df0d95efb2","year":2024},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-12T15:48:15.716052Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.131050Z"},"links":{"cited_paper":"/paper/2401.08138","citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:6811c80b067668b36f1284f77c72683471a0a013ec2991f45823731c4db1d53e","observation_id":"1d4b706a-8d53-46d4-9537-d42729467a35","resolution":{"observed_at":"2026-08-12T15:58:48.361102Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.725732Z","title":"Sentence-bert: Sentence embeddings using siamese bert- networks","venue":null,"work_id":"311864a1-8542-447f-a60e-8710b04696a3","year":2019},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-12T15:48:15.716052Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.135036Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:888ffb150ecd7956304e9de3e8e5a221d89491cf6a5c878560b607a425fecac6","observation_id":"55c39471-6c74-4266-b09f-e877b2d90aba","resolution":{"observed_at":"2026-08-12T15:58:48.728852Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.716062Z","title":"Fonseca, Wagner Meira Jr., Berthier A","venue":null,"work_id":"3dae246b-8ffa-4a29-989d-68c60f1dfec0","year":2001},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-12T15:48:15.716052Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.138752Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:44dc3756b3b90fc59bcdf5df60ca2f37311e8f3a282db7e433a4f92a18604d7b","observation_id":"fd9dafe8-7aa6-4e33-bbdf-d6f2d4f9579c","resolution":{"observed_at":"2026-08-12T15:58:48.719812Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.142480Z","title":"The early bird catches the leak: Unveiling timing side channels in LLM serving systems","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-12T15:48:15.716052Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.142480Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:07c82d4faf78b151997b15eaddd719759259bb7bb1f5a7f4d93eb7acc18493eb","observation_id":"70151361-b4af-41ad-8e63-99047fe0a28c","resolution":{"observed_at":"2026-08-12T15:58:48.142480Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.704727Z","title":"MOSS: an open conversational large language model","venue":null,"work_id":"42c72140-db31-43e2-868a-be5ab28564ba","year":2024},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-12T15:48:15.716052Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.145517Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:e62ff986048ed100b0c4869563bde3fe9e9dd208e7561b668f1549b4522f5e43","observation_id":"5daaef7f-06e2-4b93-be1b-cdc8956929af","resolution":{"observed_at":"2026-08-12T15:58:48.708827Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.693135Z","title":"Sharegpt, 2023","venue":null,"work_id":"72f6a429-0f2d-4bf4-8eac-b711a50c4118","year":2023},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-12T15:48:15.716052Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.148683Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:518b09bfec75ab471806ad2477f7f039f97f3f9226f4baf5d5b65705f8968275","observation_id":"46740945-70fb-4e0a-b719-3a7ca5d55525","resolution":{"observed_at":"2026-08-12T15:58:48.696983Z","resolver_source":"raw_fallback","status":"parse_uncertain"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.681317Z","title":"Efficient streaming language models with attention sinks","venue":null,"work_id":"f5439a9e-b553-45eb-a57c-2adf2e3babff","year":2024},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-12T15:48:15.716052Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.151726Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:226b9943a2d6242135a9d140317ebc42715067b71c9eb9cb3932ac19384050e1","observation_id":"acafe0b1-14a7-4724-bcf4-0f7c94be0a1f","resolution":{"observed_at":"2026-08-12T15:58:48.684968Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.670205Z","title":"O’Hallaron","venue":null,"work_id":"5bb19433-eafd-4cca-85d4-2c0b098eb596","year":2002},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-12T15:48:15.716052Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.155258Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:38213ecc0ee5801d01deeac004844321d58c38a5e617185d8d3ff8f4df08c504","observation_id":"46c2cb6a-1aaf-4612-a74c-4819f842e9f7","resolution":{"observed_at":"2026-08-12T15:58:48.673960Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2412.15115","last_updated":"2025-01-03T02:18:21Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-12-19T17:56:09Z","title":"Qwen2.5 Technical Report","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.15115","snapshot_observed_at":"2026-08-12T15:58:48.158767Z","title":"Qwen2.5 technical report","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-12T15:48:15.716052Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.158767Z"},"links":{"cited_paper":"/paper/2412.15115","citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:160e277775b4eb7d51b3ae6670b4ecfe15eac772c2058fd4b85d1e65d2697631","observation_id":"db3fb3c9-0a18-4ddc-82b5-c3847940bc78","resolution":{"observed_at":"2026-08-12T15:58:48.158767Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.659452Z","title":"Chunkattention: Efficient self-attention with prefix-aware KV cache and two-phase partition","venue":null,"work_id":"7379c5a2-6d9c-42e9-8f0c-1f10ba33cecf","year":2024},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-12T15:48:15.716052Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.162353Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:aa853d82f3f5256a9efc384a3b1c5cb7b257361572a4813c121f6b991e5b8c35","observation_id":"9134765f-e8b9-4578-b3ff-df789ddf3010","resolution":{"observed_at":"2026-08-12T15:58:48.663350Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.648325Z","title":"Performance of compressed inverted list caching in search engines","venue":null,"work_id":"b0399d78-3918-4d08-bbd8-0971090253b1","year":2008},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-12T15:48:15.716052Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.165516Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:400fff47665af8701f5c44b7f69c60a364b46323ee82adfcbdbdd34270315351","observation_id":"8c58630c-38c1-4d0f-a4c1-00c5b9f5bc6b","resolution":{"observed_at":"2026-08-12T15:58:48.651757Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.168601Z","title":"Barrett, Zhangyang Wang, and Beidi Chen","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-12T15:48:15.716052Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.168601Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:97465930ecf1771c0ee173bab62f8c7580620a70dfaaf395bc09f0af482fe55e","observation_id":"eb972039-3c82-483d-9b62-d63b8f6d2df9","resolution":{"observed_at":"2026-08-12T15:58:48.168601Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.631254Z","title":"Wildchat: 1m chatgpt interaction logs in the wild","venue":null,"work_id":"57ca10cc-f6d5-4cd3-a43c-2d1c56c433d5","year":2024},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-12T15:48:15.716052Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.172494Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:d17f24d969d5979fe66f83a1e721893a81615aa226f3c22599866a31dd86cfeb","observation_id":"8d2ec94d-a311-48f4-87d4-cb468bb46d7e","resolution":{"observed_at":"2026-08-12T15:58:48.635067Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.620600Z","title":"Xing, Joseph E","venue":null,"work_id":"905dc580-7b3f-43da-9f82-bef0df03ccfd","year":2024},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-12T15:48:15.716052Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.176135Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:29bdd5c34730579ababde9b511703ea3ad31395681479f9b02524e778a40c02c","observation_id":"9f6f15a9-680b-43cb-ae4d-97bf965e4e33","resolution":{"observed_at":"2026-08-12T15:58:48.624149Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2312.07104","last_updated":"2024-06-06T00:10:06Z","snapshot_observed_at":"2026-08-11T10:42:54.633244Z","submitted_at":"2023-12-12T09:34:27Z","title":"SGLang: Efficient Execution of Structured Language Model Programs","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.07104","snapshot_observed_at":"2026-08-12T15:58:48.183725Z","title":"kinky date","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-12T15:48:15.716052Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.183725Z"},"links":{"cited_paper":"/paper/2312.07104","citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:fb58a2dcee879d4027c8309e3c1e4f8e8a4bdd8f17047f7c3cca4fefc0b80228","observation_id":"eb374e5a-b77a-4d85-84b3-6bd15b4e574a","resolution":{"observed_at":"2026-08-12T15:58:48.183725Z","resolver_source":null,"status":"malformed_identifier"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.188245Z","title":"Guidelines: • The answer NA means that the abstract and introduction do not include the claims made in the paper","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-12T15:48:15.716052Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.188245Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:e518a4e9c43041de80cde1d2f030873d3575248ffdf90d004b34e8bfb70a6166","observation_id":"5931fcfd-dfff-4c8f-be91-7307a1186dda","resolution":{"observed_at":"2026-08-12T15:58:48.188245Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.598565Z","title":"Limitations","venue":null,"work_id":"a3b4c566-4a34-4475-a661-79478e5138e1","year":null},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-12T15:48:15.716052Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.192197Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:1e7dba62f006197b534a0d06979622e6a0fc42183706640ee20942a69f6252ab","observation_id":"ebdb5508-c5ca-4c93-bfd8-f71db2fd09ee","resolution":{"observed_at":"2026-08-12T15:58:48.602113Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.588273Z","title":"Guidelines: • The answer NA means that the paper does not include theoretical results","venue":null,"work_id":"efef348b-f9b4-43cc-bd96-23b68f0bc0d8","year":null},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-12T15:48:15.716052Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.196316Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:e4e7180c26c70e6f90cbbdecbbc3adc38bcc65292c8218a46c40ebd51b156109","observation_id":"4cc2ff57-2cde-4661-969f-49d7ea375e7a","resolution":{"observed_at":"2026-08-12T15:58:48.592001Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.575588Z","title":"Guidelines: • The answer NA means that the paper does not include experiments","venue":null,"work_id":"c5408018-7e2f-4e24-b623-fca79ec0c0c2","year":null},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-12T15:48:15.716052Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.200359Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:2c8cf033531561c5e5db3d0b16d50430107d4ed92cc8ac08cf48ddd68f8a4a79","observation_id":"b8396268-01a7-4634-a7ea-ec54dfac6152","resolution":{"observed_at":"2026-08-12T15:58:48.579256Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.564948Z","title":"Guidelines: • The answer NA means that paper does not include experiments requiring code","venue":null,"work_id":"c709714d-fb41-43c7-bba3-63126bc539dc","year":null},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-12T15:48:15.716052Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.204341Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:8d3acdceb953256cfc0e4faaa4c9c41362544d46f38732b33475154d66871866","observation_id":"2a19b655-f77a-44ff-a18f-f21f7894b4eb","resolution":{"observed_at":"2026-08-12T15:58:48.568744Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.554452Z","title":"• The experimental setting should be presented in the core of the paper to a level of detail that is necessary to appreciate the results and make sense of them","venue":null,"work_id":"07215d91-0336-4f4e-afe1-91ce149d49ae","year":null},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-12T15:48:15.716052Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.208712Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:6cfdea57e4167e86cb796db8df4de194b4d84d6e6ad074e93d73d9c506ac0427","observation_id":"83d019a6-5743-4c91-bbd7-d997dc31b925","resolution":{"observed_at":"2026-08-12T15:58:48.558482Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.544288Z","title":"This experimental setup is also consistent with existing studies[12, 34]","venue":null,"work_id":"97962488-9555-4804-bf07-ee947ccd27e6","year":null},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-12T15:48:15.716052Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.212336Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:d22bb501c6cedbaa10d36efe5203ab92ed68c7a730728a4d4b576c13f7712342","observation_id":"be2ff848-1b0f-4946-8ae5-4490b02099a2","resolution":{"observed_at":"2026-08-12T15:58:48.547812Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.534267Z","title":"Guidelines: • The answer NA means that the paper does not include experiments","venue":null,"work_id":"83c5480f-7de0-4658-99d1-da611f255214","year":null},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-12T15:48:15.716052Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.216428Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:72b06acf977f9ee831cad19473f01c9615c2774e0ddeb91f753f68903ec5f529","observation_id":"800ef9d8-b73f-4f5b-acc0-669ff943bbb2","resolution":{"observed_at":"2026-08-12T15:58:48.537471Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.524317Z","title":"Guidelines: • The answer NA means that the authors have not reviewed the NeurIPS Code of Ethics","venue":null,"work_id":"1fc85633-97ef-431f-a3af-aae69912202b","year":null},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-12T15:48:15.716052Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.220373Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:77193e04d3ef7ea9a0cc9dcda9d79e118e3a7ba06eb688ed61e49cb8fb40e181","observation_id":"1d185f71-5ee1-4cbc-9f7a-2a0b69d7905f","resolution":{"observed_at":"2026-08-12T15:58:48.527890Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.512537Z","title":"Guidelines: • The answer NA means that there is no societal impact of the work performed","venue":null,"work_id":"6c12dcb9-aef9-4882-9edb-3662e0f56125","year":null},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-12T15:48:15.716052Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.223953Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:81b98fa59ec0bc5446ab65b70ea1a816f8474fd76f23b5bfd308e2dde1502dff","observation_id":"8451f0d9-f0b5-4737-af24-dca0c05e1a75","resolution":{"observed_at":"2026-08-12T15:58:48.516492Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.500580Z","title":"Guidelines: • The answer NA means that the paper poses no such risks","venue":null,"work_id":"7e6176c6-d58a-4266-bbd5-871793528aac","year":null},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-12T15:48:15.716052Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.227844Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:83997456aef51ec05c99fdf9c7555e0f5fa028e76d7ec11561d340494445dac5","observation_id":"9061bc30-2137-4789-b29a-cf01f4c521ac","resolution":{"observed_at":"2026-08-12T15:58:48.504972Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.488042Z","title":"Guidelines: • The answer NA means that the paper does not use existing assets","venue":null,"work_id":"f121fbc3-9650-43d6-8c73-87ae5e7a7e7f","year":null},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-12T15:48:15.716052Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.231431Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:e4d1be6f6bdb79fcdb440c322bfd0557cab8e973eb9f4ac0e4cb019f8f789c13","observation_id":"03375e8b-a654-4846-b014-a03e5f2d0c8f","resolution":{"observed_at":"2026-08-12T15:58:48.492294Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.235405Z","title":"Guidelines: • The answer NA means that the paper does not release new assets","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-12T15:48:15.716052Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.235405Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:b5363e7c6d08427f2a9e7c1308c4d31dde8e8f3a25ab23f0b3d3a7d48c28790f","observation_id":"c7a31830-0eae-4c82-986e-d214d17cedc7","resolution":{"observed_at":"2026-08-12T15:58:48.235405Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.239285Z","title":"Guidelines: • The answer NA means that the paper does not involve crowdsourcing nor research with human subjects","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-12T15:48:15.716052Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.239285Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:0719e19d59652ec86e87a0e9526043d164ea7ac84cdbcb70b5de098f35e93e3a","observation_id":"d157f936-cc07-43b8-aca2-8388b67ab52e","resolution":{"observed_at":"2026-08-12T15:58:48.239285Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.457648Z","title":"Guidelines: • The answer NA means that the paper does not involve crowdsourcing nor research with human subjects","venue":null,"work_id":"d59d90a2-3efe-4378-b781-507d725370f2","year":null},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-12T15:48:15.716052Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.243029Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:2b961d070cafaf8b8c6a32a537c3094a1d949999685eb23a48a12b56d0bbf7bb","observation_id":"446d0345-9ff9-4eb6-9f91-209a7b8a567c","resolution":{"observed_at":"2026-08-12T15:58:48.462095Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.444465Z","title":"Answer: [NA] Justification: The core method development in this research does not involve LLMs as any important, original, or non-standard components","venue":null,"work_id":"ec04fe49-8440-4de1-8062-9931d9fcb86f","year":2025},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-12T15:48:15.716052Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.247129Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:63f389bf31f9c7fd00cf3dd9c461ba246b3c54ced0a8fa615884de08b2e4ec14","observation_id":"5804080a-9472-46f2-a961-36321cc23137","resolution":{"observed_at":"2026-08-12T15:58:48.448600Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-13T06:32:02.005865+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.063168Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-12T15:48:15.716052Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":2023,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.063168Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:e8fd04c0a245de4386a9a31f1b7b8c16eb1ac67519ac641b01521b53f85e8612","observation_id":"f3280a63-f62c-4be1-a0d4-6a44e68241ed","resolution":{"observed_at":"2026-08-12T15:58:48.063168Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-12T15:58:48.179849Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","snapshot_observed_at":"2026-08-12T15:48:15.716052Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving","version":2},"reference_index":2024,"source":"pdf_text","source_observed_at":"2026-08-12T15:58:48.179849Z"},"links":{"citing_paper":"/paper/2411.13820"},"observation_digest":"sha256:6084a536f264045d128a32c555cd4d570c1a1e633c90c51c85b47b72953a8259","observation_id":"873708e1-60b9-40f8-83f5-05df7e040935","resolution":{"observed_at":"2026-08-12T15:58:48.179849Z","resolver_source":null,"status":"parse_uncertain"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2411.13820","last_updated":"2025-07-14T02:22:43Z","latest_version":2,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-12T15:48:15.716052Z","submitted_at":"2024-11-21T03:52:41Z","title":"InstCache: A Predictive Cache for LLM Serving"},"reference_resolution":{"displayed":51,"state_counts":{"malformed_identifier":1,"metadata_mismatch":0,"parse_uncertain":2,"unresolved":14,"verified_exact":1,"verified_fuzzy":33},"total_outbound_references":51},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-13T06:32:02.005865+00:00","source":"crossref"},{"observed_at":"2026-08-13T06:31:53.387327+00:00","source":"retraction_watch"}],"thesis":"As of 13 August 2026, this Paper Citation Record lists 51 of 51 outbound references and 0 inbound Pith citation observations for arXiv:2411.13820."}