{"as_of":"2026-08-07T18:00:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:7941f799ba6b85d6b2f727ccd6a7c5212de9204faa3e0edf01e5b7296ab2367c","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":11,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":11,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-07T06:34:17.273281+00:00","state":"measured"},{"denominator":11,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":11,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T00:14:54.845630Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-06-30T00:24:04.305830Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2501.12162","last_updated":"2025-05-17T07:09:10Z","snapshot_observed_at":"2026-08-05T21:38:55.998982Z","submitted_at":"2025-01-21T14:15:01Z","title":"AdaServe: Accelerating Multi-SLO LLM Serving with SLO-Customized Speculative Decoding","version":2},"cited_work":{"arxiv_id":"2501.12162","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2501.12162","snapshot_observed_at":"2026-06-30T00:24:04.305830Z","title":"Adaserve: Accelerating multi-slo llm serving with slo-customized speculative decoding.arXiv preprint arXiv:2501.12162","venue":null,"work_id":"29bbd187-e529-40c4-bb30-608f5bee82b7","year":2025},"citing_paper":{"arxiv_id":"2504.15965","last_updated":"2025-04-23T13:47:27Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-04-22T15:05:04Z","title":"From Human Memory to AI Memory: A Survey on Memory Mechanisms in the Era of LLMs","version":2},"reference_index":125,"source":"pdf_text","source_observed_at":"2026-05-17T11:05:09.588491Z"},"links":{"cited_paper":"/paper/2501.12162","citing_paper":"/paper/2504.15965"},"observation_digest":"sha256:7fe49754c02197f0183516a4bd97f2c49e589c64104899bd4f14fde7a771489e","observation_id":"4aace16e-ebb0-400d-9d31-96c44e5694fe","resolution":{"observed_at":"2026-05-17T11:05:09.853249Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.12162","last_updated":"2025-05-17T07:09:10Z","snapshot_observed_at":"2026-08-05T21:38:55.998982Z","submitted_at":"2025-01-21T14:15:01Z","title":"AdaServe: Accelerating Multi-SLO LLM Serving with SLO-Customized Speculative Decoding","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.12162","snapshot_observed_at":"2026-08-07T00:14:54.845630Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.20675","last_updated":"2025-06-17T20:06:08Z","snapshot_observed_at":"2026-08-07T00:07:24.139747Z","submitted_at":"2025-06-17T20:06:08Z","title":"Utility-Driven Speculative Decoding for Mixture-of-Experts","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-07T00:14:54.845630Z"},"links":{"cited_paper":"/paper/2501.12162","citing_paper":"/paper/2506.20675"},"observation_digest":"sha256:c7e740aae3301fa9abdff859c94fea3be44382be1b8254550946371c725e7b63","observation_id":"aae5d117-b459-44fc-a0df-b8277c46f359","resolution":{"observed_at":"2026-08-07T00:14:54.845630Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.12162","last_updated":"2025-05-17T07:09:10Z","snapshot_observed_at":"2026-08-05T21:38:55.998982Z","submitted_at":"2025-01-21T14:15:01Z","title":"AdaServe: Accelerating Multi-SLO LLM Serving with SLO-Customized Speculative Decoding","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.12162","snapshot_observed_at":"2026-08-06T15:06:47.960865Z","title":"Adaserve: Slo-customized llm serving with fine-grained speculative decoding","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.16731","last_updated":"2025-07-22T16:13:43Z","snapshot_observed_at":"2026-08-06T15:00:33.300567Z","submitted_at":"2025-07-22T16:13:43Z","title":"Collaborative Inference and Learning between Edge SLMs and Cloud LLMs: A Survey of Algorithms, Execution, and Open Challenges","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-06T15:06:47.960865Z"},"links":{"cited_paper":"/paper/2501.12162","citing_paper":"/paper/2507.16731"},"observation_digest":"sha256:63672922ae2ee614aa30fde815659a26cb82ca24cdda7349898287ae7bb09790","observation_id":"c5ae9cde-a63b-4669-ae61-f28a8dab7742","resolution":{"observed_at":"2026-08-06T15:06:47.960865Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.12162","last_updated":"2025-05-17T07:09:10Z","snapshot_observed_at":"2026-08-05T21:38:55.998982Z","submitted_at":"2025-01-21T14:15:01Z","title":"AdaServe: Accelerating Multi-SLO LLM Serving with SLO-Customized Speculative Decoding","version":2},"cited_work":{"arxiv_id":"2501.12162","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2501.12162","snapshot_observed_at":"2026-06-30T00:24:04.305830Z","title":"Adaserve: Accelerating multi-slo llm serving with slo-customized speculative decoding.arXiv preprint arXiv:2501.12162","venue":null,"work_id":"29bbd187-e529-40c4-bb30-608f5bee82b7","year":2025},"citing_paper":{"arxiv_id":"2601.20309","last_updated":"2026-05-18T19:51:16Z","snapshot_observed_at":"2026-07-06T22:43:17.026472Z","submitted_at":"2026-01-28T07:01:46Z","title":"SuperInfer: SLO-Aware Rotary Scheduling and Memory Management for LLM Inference on Superchips","version":2},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-05-21T15:26:01.283448Z"},"links":{"cited_paper":"/paper/2501.12162","citing_paper":"/paper/2601.20309"},"observation_digest":"sha256:a49cd88e292551c3b13b79e20d402e221570d951d32185c34bac143e07b88675","observation_id":"9560cedb-568d-41d3-afb4-b30209e55e71","resolution":{"observed_at":"2026-05-21T15:30:17.977958Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.12162","last_updated":"2025-05-17T07:09:10Z","snapshot_observed_at":"2026-08-05T21:38:55.998982Z","submitted_at":"2025-01-21T14:15:01Z","title":"AdaServe: Accelerating Multi-SLO LLM Serving with SLO-Customized Speculative Decoding","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.12162","snapshot_observed_at":"2026-08-02T21:14:48.629742Z","title":"Adaserve: Accelerating multi-slo llm serving with slo-customized speculative decoding","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2603.18016","last_updated":"2026-06-01T17:56:47Z","snapshot_observed_at":"2026-08-05T22:27:54.412362Z","submitted_at":"2026-02-24T17:24:50Z","title":"MineDraft: A Framework for Batch Parallel Speculative Decoding","version":2},"reference_index":17,"source":"arxiv_source","source_observed_at":"2026-08-02T21:14:48.629742Z"},"links":{"cited_paper":"/paper/2501.12162","citing_paper":"/paper/2603.18016"},"observation_digest":"sha256:726d0980e62c2a44bdac54e2c6ed1882e00c2899526ee0e2ce47d96fe32ae1ad","observation_id":"6722a6bf-1307-4619-a1e6-304ebba110c7","resolution":{"observed_at":"2026-08-02T21:14:48.629742Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.12162","last_updated":"2025-05-17T07:09:10Z","snapshot_observed_at":"2026-08-05T21:38:55.998982Z","submitted_at":"2025-01-21T14:15:01Z","title":"AdaServe: Accelerating Multi-SLO LLM Serving with SLO-Customized Speculative Decoding","version":2},"cited_work":{"arxiv_id":"2501.12162","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2501.12162","snapshot_observed_at":"2026-06-30T00:24:04.305830Z","title":"Adaserve: Accelerating multi-slo llm serving with slo-customized speculative decoding.arXiv preprint arXiv:2501.12162","venue":null,"work_id":"29bbd187-e529-40c4-bb30-608f5bee82b7","year":2025},"citing_paper":{"arxiv_id":"2604.17227","last_updated":"2026-04-19T03:20:30Z","snapshot_observed_at":"2026-07-06T23:04:23.730478Z","submitted_at":"2026-04-19T03:20:30Z","title":"Cloud-native and Distributed Systems for Efficient and Scalable Large Language Models -- A Research Agenda","version":1},"reference_index":79,"source":"pdf_text","source_observed_at":"2026-05-10T06:27:23.580445Z"},"links":{"cited_paper":"/paper/2501.12162","citing_paper":"/paper/2604.17227"},"observation_digest":"sha256:84f9958a80111377cea78a87b7939ee4cb58ad0961caae927d8da263863829cc","observation_id":"7e021acf-ec94-43f0-ad98-bddc56040aaa","resolution":{"observed_at":"2026-05-10T06:31:30.784210Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.12162","last_updated":"2025-05-17T07:09:10Z","snapshot_observed_at":"2026-08-05T21:38:55.998982Z","submitted_at":"2025-01-21T14:15:01Z","title":"AdaServe: Accelerating Multi-SLO LLM Serving with SLO-Customized Speculative Decoding","version":2},"cited_work":{"arxiv_id":"2501.12162","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2501.12162","snapshot_observed_at":"2026-06-30T00:24:04.305830Z","title":"Adaserve: Accelerating multi-slo llm serving with slo-customized speculative decoding.arXiv preprint arXiv:2501.12162","venue":null,"work_id":"29bbd187-e529-40c4-bb30-608f5bee82b7","year":2025},"citing_paper":{"arxiv_id":"2604.20503","last_updated":"2026-04-22T12:44:39Z","snapshot_observed_at":"2026-07-06T23:06:55.880484Z","submitted_at":"2026-04-22T12:44:39Z","title":"FASER: Fine-Grained Phase Management for Speculative Decoding in Dynamic LLM Serving","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-05-09T22:56:19.734262Z"},"links":{"cited_paper":"/paper/2501.12162","citing_paper":"/paper/2604.20503"},"observation_digest":"sha256:6aea7ceccf7ec7ba67fdf329719823136acabb6b6767a946d11b6772246c7c3d","observation_id":"dabdead1-cca5-4f88-b5b4-9b8fcd46bc27","resolution":{"observed_at":"2026-05-09T22:59:17.704940Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.12162","last_updated":"2025-05-17T07:09:10Z","snapshot_observed_at":"2026-08-05T21:38:55.998982Z","submitted_at":"2025-01-21T14:15:01Z","title":"AdaServe: Accelerating Multi-SLO LLM Serving with SLO-Customized Speculative Decoding","version":2},"cited_work":{"arxiv_id":"2501.12162","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2501.12162","snapshot_observed_at":"2026-06-30T00:24:04.305830Z","title":"Adaserve: Accelerating multi-slo llm serving with slo-customized speculative decoding.arXiv preprint arXiv:2501.12162","venue":null,"work_id":"29bbd187-e529-40c4-bb30-608f5bee82b7","year":2025},"citing_paper":{"arxiv_id":"2605.04357","last_updated":"2026-05-05T23:25:29Z","snapshot_observed_at":"2026-07-06T23:17:09.246351Z","submitted_at":"2026-05-05T23:25:29Z","title":"Coral: Cost-Efficient Multi-LLM Serving over Heterogeneous Cloud GPUs","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-05-08T16:45:39.853764Z"},"links":{"cited_paper":"/paper/2501.12162","citing_paper":"/paper/2605.04357"},"observation_digest":"sha256:46e86f59b294454731416ce878adc3429379533cf3339506fcfc4f55b8f35317","observation_id":"5648bcaf-9149-4720-a19c-7141d5db253f","resolution":{"observed_at":"2026-05-11T18:01:07.087778Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.12162","last_updated":"2025-05-17T07:09:10Z","snapshot_observed_at":"2026-08-05T21:38:55.998982Z","submitted_at":"2025-01-21T14:15:01Z","title":"AdaServe: Accelerating Multi-SLO LLM Serving with SLO-Customized Speculative Decoding","version":2},"cited_work":{"arxiv_id":"2501.12162","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2501.12162","snapshot_observed_at":"2026-06-30T00:24:04.305830Z","title":"Adaserve: Accelerating multi-slo llm serving with slo-customized speculative decoding.arXiv preprint arXiv:2501.12162","venue":null,"work_id":"29bbd187-e529-40c4-bb30-608f5bee82b7","year":2025},"citing_paper":{"arxiv_id":"2605.06914","last_updated":"2026-05-07T20:23:32Z","snapshot_observed_at":"2026-08-02T05:38:27.537862Z","submitted_at":"2026-05-07T20:23:32Z","title":"Regulating Branch Parallelism in LLM Serving","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-05-11T01:00:53.308946Z"},"links":{"cited_paper":"/paper/2501.12162","citing_paper":"/paper/2605.06914"},"observation_digest":"sha256:f20eac3aa2831e17d45eb49217fb603b204332ea051953ed742402b4c997673b","observation_id":"c8f10b5b-eab2-45a4-8e9f-af97cd7c8db6","resolution":{"observed_at":"2026-05-11T04:50:58.328653Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.12162","last_updated":"2025-05-17T07:09:10Z","snapshot_observed_at":"2026-08-05T21:38:55.998982Z","submitted_at":"2025-01-21T14:15:01Z","title":"AdaServe: Accelerating Multi-SLO LLM Serving with SLO-Customized Speculative Decoding","version":2},"cited_work":{"arxiv_id":"2501.12162","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2501.12162","snapshot_observed_at":"2026-06-30T00:24:04.305830Z","title":"Adaserve: Accelerating multi-slo llm serving with slo-customized speculative decoding.arXiv preprint arXiv:2501.12162","venue":null,"work_id":"29bbd187-e529-40c4-bb30-608f5bee82b7","year":2025},"citing_paper":{"arxiv_id":"2605.24832","last_updated":"2026-05-24T02:56:46Z","snapshot_observed_at":"2026-08-06T08:58:18.428802Z","submitted_at":"2026-05-24T02:56:46Z","title":"Optimus: Elastic Decoding for Efficient Diffusion LLM Serving","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-06-30T00:15:58.082725Z"},"links":{"cited_paper":"/paper/2501.12162","citing_paper":"/paper/2605.24832"},"observation_digest":"sha256:53015232364df901491419c62bb12c162f406f22e0d180157b7dfadbc0120b68","observation_id":"c93f4313-07d0-48b5-b1c2-56bc619e7ef3","resolution":{"observed_at":"2026-06-30T00:24:04.307474Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.12162","last_updated":"2025-05-17T07:09:10Z","snapshot_observed_at":"2026-08-05T21:38:55.998982Z","submitted_at":"2025-01-21T14:15:01Z","title":"AdaServe: Accelerating Multi-SLO LLM Serving with SLO-Customized Speculative Decoding","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.12162","snapshot_observed_at":"2026-08-03T01:35:54.834934Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.28848","last_updated":"2026-07-30T21:20:21Z","snapshot_observed_at":"2026-08-07T03:28:41.264411Z","submitted_at":"2026-07-30T21:20:21Z","title":"DeltaServe: Host-Agnostic Co-Serving of Inference and Fine-Tuning for LLMs","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-03T01:35:54.834934Z"},"links":{"cited_paper":"/paper/2501.12162","citing_paper":"/paper/2607.28848"},"observation_digest":"sha256:34a53ddd4549d0d384511cae12bfa29bab7e633b420f852754a23f0258f1a5de","observation_id":"6966aa44-9947-4cbb-85b0-c6916ee00f94","resolution":{"observed_at":"2026-08-03T01:35:54.834934Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2501.12162/citation-record","integrity":"/paper/2501.12162/integrity","json":"/paper/2501.12162/citation-record.json","paper":"/paper/2501.12162"},"outbound":[],"paper":{"arxiv_id":"2501.12162","last_updated":"2025-05-17T07:09:10Z","latest_version":2,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-05T21:38:55.998982Z","submitted_at":"2025-01-21T14:15:01Z","title":"AdaServe: Accelerating Multi-SLO LLM Serving with SLO-Customized Speculative Decoding"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"thesis":"As of 7 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 11 inbound Pith citation observations for arXiv:2501.12162."}