{"as_of":"2026-08-08T08:47:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:2727186d2e8652024a39c2b4c54cbc70e885827428529f072c771ccb118539f5","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":11,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":11,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-08T06:32:00.761636+00:00","state":"measured"},{"denominator":11,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":11,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T18:55:04.262738Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-07-04T05:09:36.948141Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2406.03243","last_updated":"2024-06-05T13:20:18Z","snapshot_observed_at":"2026-07-06T18:25:53.102819Z","submitted_at":"2024-06-05T13:20:18Z","title":"Llumnix: Dynamic Scheduling for Large Language Model Serving","version":1},"cited_work":{"arxiv_id":"2406.03243","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2406.03243","snapshot_observed_at":"2026-07-04T05:09:36.948141Z","title":"Llumnix: Dynamic scheduling for large language model serving","venue":null,"work_id":"04bf0f56-7543-46e3-af2b-01f5283e0f2a","year":2024},"citing_paper":{"arxiv_id":"2404.14294","last_updated":"2024-07-19T04:47:36Z","snapshot_observed_at":"2026-07-06T18:03:47.096406Z","submitted_at":"2024-04-22T15:53:08Z","title":"A Survey on Efficient Inference for Large Language Models","version":3},"reference_index":288,"source":"pdf_text","source_observed_at":"2026-05-15T02:39:33.007894Z"},"links":{"cited_paper":"/paper/2406.03243","citing_paper":"/paper/2404.14294"},"observation_digest":"sha256:9c03f794a826bcd83dac7629c4133dd41f209a4d9e6ff35180900df0d0346a2d","observation_id":"a8662879-5233-4297-891e-10034a42c6f2","resolution":{"observed_at":"2026-05-15T02:39:33.533295Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.03243","last_updated":"2024-06-05T13:20:18Z","snapshot_observed_at":"2026-07-06T18:25:53.102819Z","submitted_at":"2024-06-05T13:20:18Z","title":"Llumnix: Dynamic Scheduling for Large Language Model Serving","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.03243","snapshot_observed_at":"2026-08-07T18:55:04.262738Z","title":"Llumnix: Dynamic scheduling for large language model serving,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2502.15763","last_updated":"2025-02-14T16:00:00Z","snapshot_observed_at":"2026-08-07T18:45:27.225294Z","submitted_at":"2025-02-14T16:00:00Z","title":"Hybrid Offline-online Scheduling Method for Large Language Model Inference Optimization","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-07T18:55:04.262738Z"},"links":{"cited_paper":"/paper/2406.03243","citing_paper":"/paper/2502.15763"},"observation_digest":"sha256:d9d781cb496c03b8da8d3448825a1eeb1d28541d220260c82b866c81b4abf7b0","observation_id":"a2107f8d-b8b7-4e61-ab44-c7932de99b71","resolution":{"observed_at":"2026-08-07T18:55:04.262738Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.03243","last_updated":"2024-06-05T13:20:18Z","snapshot_observed_at":"2026-07-06T18:25:53.102819Z","submitted_at":"2024-06-05T13:20:18Z","title":"Llumnix: Dynamic Scheduling for Large Language Model Serving","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.03243","snapshot_observed_at":"2026-08-06T20:59:01.514175Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2507.01438","last_updated":"2025-07-02T07:47:28Z","snapshot_observed_at":"2026-08-06T20:49:01.786008Z","submitted_at":"2025-07-02T07:47:28Z","title":"EdgeLoRA: An Efficient Multi-Tenant LLM Serving System on Edge Devices","version":1},"reference_index":74,"source":"pdf_text","source_observed_at":"2026-08-06T20:59:01.514175Z"},"links":{"cited_paper":"/paper/2406.03243","citing_paper":"/paper/2507.01438"},"observation_digest":"sha256:327f9ff20c98acb5ce1ec03f9973cd217236bf59a22a657afd8184d6086d71d0","observation_id":"d580ae2e-c08a-4398-9254-a1a722f9e353","resolution":{"observed_at":"2026-08-06T20:59:01.514175Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.03243","last_updated":"2024-06-05T13:20:18Z","snapshot_observed_at":"2026-07-06T18:25:53.102819Z","submitted_at":"2024-06-05T13:20:18Z","title":"Llumnix: Dynamic Scheduling for Large Language Model Serving","version":1},"cited_work":{"arxiv_id":"2406.03243","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2406.03243","snapshot_observed_at":"2026-07-04T05:09:36.948141Z","title":"Llumnix: Dynamic scheduling for large language model serving","venue":null,"work_id":"04bf0f56-7543-46e3-af2b-01f5283e0f2a","year":2024},"citing_paper":{"arxiv_id":"2604.04750","last_updated":"2026-04-09T14:13:19Z","snapshot_observed_at":"2026-07-06T22:53:37.357996Z","submitted_at":"2026-04-06T15:16:35Z","title":"DeepStack: Scalable and Accurate Design Space Exploration for Distributed 3D-Stacked AI Accelerators","version":2},"reference_index":100,"source":"pdf_text","source_observed_at":"2026-05-10T19:04:17.725111Z"},"links":{"cited_paper":"/paper/2406.03243","citing_paper":"/paper/2604.04750"},"observation_digest":"sha256:0c1074ecad73d8fef191b2bd18281dd15c544c7c8238ee5ff624c91388e93d4a","observation_id":"ec890e84-bc22-4e75-a72c-e1508795c22d","resolution":{"observed_at":"2026-05-10T23:30:51.976790Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.03243","last_updated":"2024-06-05T13:20:18Z","snapshot_observed_at":"2026-07-06T18:25:53.102819Z","submitted_at":"2024-06-05T13:20:18Z","title":"Llumnix: Dynamic Scheduling for Large Language Model Serving","version":1},"cited_work":{"arxiv_id":"2406.03243","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2406.03243","snapshot_observed_at":"2026-07-04T05:09:36.948141Z","title":"Llumnix: Dynamic scheduling for large language model serving","venue":null,"work_id":"04bf0f56-7543-46e3-af2b-01f5283e0f2a","year":2024},"citing_paper":{"arxiv_id":"2605.20863","last_updated":"2026-05-20T07:55:06Z","snapshot_observed_at":"2026-07-06T23:31:23.407780Z","submitted_at":"2026-05-20T07:55:06Z","title":"PlexRL: Cluster-Level Orchestration of Serviceized LLM Execution for RLVR","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-05-21T02:24:48.872065Z"},"links":{"cited_paper":"/paper/2406.03243","citing_paper":"/paper/2605.20863"},"observation_digest":"sha256:ab7d1d45d80076bb09de97a4f56d5210c25f98da3ebcce327a6e9a464360d9a6","observation_id":"9b8ecfe1-4204-47c1-86a8-ff631983be04","resolution":{"observed_at":"2026-05-21T02:29:25.459494Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.03243","last_updated":"2024-06-05T13:20:18Z","snapshot_observed_at":"2026-07-06T18:25:53.102819Z","submitted_at":"2024-06-05T13:20:18Z","title":"Llumnix: Dynamic Scheduling for Large Language Model Serving","version":1},"cited_work":{"arxiv_id":"2406.03243","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2406.03243","snapshot_observed_at":"2026-07-04T05:09:36.948141Z","title":"Llumnix: Dynamic scheduling for large language model serving","venue":null,"work_id":"04bf0f56-7543-46e3-af2b-01f5283e0f2a","year":2024},"citing_paper":{"arxiv_id":"2606.20295","last_updated":"2026-07-24T07:54:17Z","snapshot_observed_at":"2026-08-02T10:48:55.986116Z","submitted_at":"2026-06-18T14:33:09Z","title":"Token-Operations-Oriented Inference Optimization Techniques for Large Models","version":1},"reference_index":209,"source":"pdf_text","source_observed_at":"2026-06-26T16:15:22.543601Z"},"links":{"cited_paper":"/paper/2406.03243","citing_paper":"/paper/2606.20295"},"observation_digest":"sha256:a0e57d2dba974d07d3cee06e1539218a349ac08a18fc64ec521808600e54d909","observation_id":"69c54834-2df4-4c20-99a8-fd16db5453c1","resolution":{"observed_at":"2026-07-04T05:09:36.950269Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.03243","last_updated":"2024-06-05T13:20:18Z","snapshot_observed_at":"2026-07-06T18:25:53.102819Z","submitted_at":"2024-06-05T13:20:18Z","title":"Llumnix: Dynamic Scheduling for Large Language Model Serving","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.03243","snapshot_observed_at":"2026-08-02T10:49:23.421625Z","title":"Llumnix: Dynamic Scheduling for Large Language Model Serving","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2606.20295","last_updated":"2026-07-24T07:54:17Z","snapshot_observed_at":"2026-08-02T10:48:55.986116Z","submitted_at":"2026-06-18T14:33:09Z","title":"Token-Operations-Oriented Inference Optimization Techniques for Large Models","version":2},"reference_index":209,"source":"pdf_text","source_observed_at":"2026-08-02T10:49:23.421625Z"},"links":{"cited_paper":"/paper/2406.03243","citing_paper":"/paper/2606.20295"},"observation_digest":"sha256:6b9670a94a82c347c2d1eab0536954edaf2a798c627498611483c7df2a201200","observation_id":"12f4dc19-0ea6-498b-9e11-7c6542185cf7","resolution":{"observed_at":"2026-08-02T10:49:23.421625Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.03243","last_updated":"2024-06-05T13:20:18Z","snapshot_observed_at":"2026-07-06T18:25:53.102819Z","submitted_at":"2024-06-05T13:20:18Z","title":"Llumnix: Dynamic Scheduling for Large Language Model Serving","version":1},"cited_work":{"arxiv_id":"2406.03243","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2406.03243","snapshot_observed_at":"2026-07-04T05:09:36.948141Z","title":"Llumnix: Dynamic scheduling for large language model serving","venue":null,"work_id":"04bf0f56-7543-46e3-af2b-01f5283e0f2a","year":2024},"citing_paper":{"arxiv_id":"2607.01579","last_updated":"2026-07-02T01:23:02Z","snapshot_observed_at":"2026-08-02T15:54:32.616830Z","submitted_at":"2026-07-02T01:23:02Z","title":"OmniPilot: An Uncertainty-Aware LLM Inference Advisor for Heterogeneous GPU Clusters","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-07-03T06:30:30.308713Z"},"links":{"cited_paper":"/paper/2406.03243","citing_paper":"/paper/2607.01579"},"observation_digest":"sha256:294e3c58a5a3be7ac8c46c76206e207d681eca971652a7a230d942f1fb20a8da","observation_id":"69adefae-59f4-4d84-b581-8b282b50b1bd","resolution":{"observed_at":"2026-07-03T06:37:42.178086Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.03243","last_updated":"2024-06-05T13:20:18Z","snapshot_observed_at":"2026-07-06T18:25:53.102819Z","submitted_at":"2024-06-05T13:20:18Z","title":"Llumnix: Dynamic Scheduling for Large Language Model Serving","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.03243","snapshot_observed_at":"2026-07-12T05:56:18.797766Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.02942","last_updated":"2026-07-03T04:28:49Z","snapshot_observed_at":"2026-08-08T04:24:02.724600Z","submitted_at":"2026-07-03T04:28:49Z","title":"A Workflow-Aware Serving Layer for Agentic Applications","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-07-12T05:56:18.797766Z"},"links":{"cited_paper":"/paper/2406.03243","citing_paper":"/paper/2607.02942"},"observation_digest":"sha256:e3d5bb0824d5267675d382bba6f4b373e93d3984cf6b655b82b6c7534ceea8e8","observation_id":"606ba1af-719c-4a2e-98f9-0ee96a7ab5f8","resolution":{"observed_at":"2026-07-12T05:56:18.797766Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.03243","last_updated":"2024-06-05T13:20:18Z","snapshot_observed_at":"2026-07-06T18:25:53.102819Z","submitted_at":"2024-06-05T13:20:18Z","title":"Llumnix: Dynamic Scheduling for Large Language Model Serving","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.03243","snapshot_observed_at":"2026-07-30T11:40:16.380604Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.23815","last_updated":"2026-07-26T19:29:31Z","snapshot_observed_at":"2026-08-08T08:44:41.089167Z","submitted_at":"2026-07-26T19:29:31Z","title":"Kalypso: Relational LLM Serving","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-07-30T11:40:16.380604Z"},"links":{"cited_paper":"/paper/2406.03243","citing_paper":"/paper/2607.23815"},"observation_digest":"sha256:24381e9228a7220947240c0f3b8046d5b6bb081854d46d486a6e3d308ec8b960","observation_id":"d9952405-2363-41ae-9c04-8732e4777c9a","resolution":{"observed_at":"2026-07-30T11:40:16.380604Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.03243","last_updated":"2024-06-05T13:20:18Z","snapshot_observed_at":"2026-07-06T18:25:53.102819Z","submitted_at":"2024-06-05T13:20:18Z","title":"Llumnix: Dynamic Scheduling for Large Language Model Serving","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.03243","snapshot_observed_at":"2026-08-06T00:33:27.295482Z","title":"Llumnix: Dynamic scheduling for large language model serving,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2608.01126","last_updated":"2026-08-02T09:51:15Z","snapshot_observed_at":"2026-08-07T16:53:54.514847Z","submitted_at":"2026-08-02T09:51:15Z","title":"Spatial Prefix Caching for Wireless Edge LLM Inference: A Stochastic-Geometry and Queueing Framework","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-06T00:33:27.295482Z"},"links":{"cited_paper":"/paper/2406.03243","citing_paper":"/paper/2608.01126"},"observation_digest":"sha256:69a9cac6fc05baf4f07cf7fd6afc1de074a5b1df9211c7c0fbd183ea741a0759","observation_id":"dee33baa-89dd-4005-8575-75202281c9fe","resolution":{"observed_at":"2026-08-06T00:33:27.295482Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2406.03243/citation-record","integrity":"/paper/2406.03243/integrity","json":"/paper/2406.03243/citation-record.json","paper":"/paper/2406.03243"},"outbound":[],"paper":{"arxiv_id":"2406.03243","last_updated":"2024-06-05T13:20:18Z","latest_version":1,"primary_category":"cs.AR","snapshot_observed_at":"2026-07-06T18:25:53.102819Z","submitted_at":"2024-06-05T13:20:18Z","title":"Llumnix: Dynamic Scheduling for Large Language Model Serving"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"thesis":"As of 8 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 11 inbound Pith citation observations for arXiv:2406.03243."}