{"as_of":"2026-08-12T05:52:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:56519f3bce43c13fa247e240a728633076ce98ccf4a5c3ff5c6282160b3dee20","coverage":[{"denominator":61,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":61,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-11T04:38:25.630673Z","state":"measured"},{"denominator":61,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":61,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-11T06:34:44.6726+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2412.18695/citation-record","integrity":"/paper/2412.18695/integrity","json":"/paper/2412.18695/citation-record.json","paper":"/paper/2412.18695"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T04:38:27.843340Z","title":"https://github.c om/vllm-project/vllm, 2024","venue":null,"work_id":"e7cd6599-2dfc-4b6d-9099-f27ad544f349","year":2024},"citing_paper":{"arxiv_id":"2412.18695","last_updated":"2024-12-24T22:51:29Z","snapshot_observed_at":"2026-08-11T23:49:11.223523Z","submitted_at":"2024-12-24T22:51:29Z","title":"TimelyLLM: Segmented LLM Serving System for Time-sensitive Robotic Applications","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-11T04:38:24.243032Z"},"links":{"citing_paper":"/paper/2412.18695"},"observation_digest":"sha256:16754998f9a86300ccf91bded919c1b67fb5a278d13efbabbc4b185521ba259b","observation_id":"99928eec-c128-4ca8-a7f7-e2a13fd6493a","resolution":{"observed_at":"2026-08-11T04:38:27.848114Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2404.14219","last_updated":"2024-08-30T21:17:17Z","snapshot_observed_at":"2026-08-10T14:07:02.234322Z","submitted_at":"2024-04-22T14:32:33Z","title":"Phi-3 Technical Report: A Highly Capable Language Model Locally on Your Phone","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.14219","snapshot_observed_at":"2026-08-11T04:38:24.272878Z","title":"Phi-3 technical report: A highly capable language model locally on your phone.arXiv preprint arXiv:2404.14219, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.18695","last_updated":"2024-12-24T22:51:29Z","snapshot_observed_at":"2026-08-11T23:49:11.223523Z","submitted_at":"2024-12-24T22:51:29Z","title":"TimelyLLM: Segmented LLM Serving System for Time-sensitive Robotic Applications","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-11T04:38:24.272878Z"},"links":{"cited_paper":"/paper/2404.14219","citing_paper":"/paper/2412.18695"},"observation_digest":"sha256:a2f0fa33cfe2a03f7d692c6a059170b5b3c1d13f6e3f0bf7da06f8357d593fe4","observation_id":"395a7211-d1e2-4263-8129-2699c69bebc3","resolution":{"observed_at":"2026-08-11T04:38:24.272878Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T04:38:27.832279Z","title":"Infercept: Efficient intercept support for augmented large language model inference","venue":null,"work_id":"fd083bcc-13bc-4457-b633-a251b9c9884a","year":null},"citing_paper":{"arxiv_id":"2412.18695","last_updated":"2024-12-24T22:51:29Z","snapshot_observed_at":"2026-08-11T23:49:11.223523Z","submitted_at":"2024-12-24T22:51:29Z","title":"TimelyLLM: Segmented LLM Serving System for Time-sensitive Robotic Applications","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-11T04:38:24.327030Z"},"links":{"citing_paper":"/paper/2412.18695"},"observation_digest":"sha256:760ca6c706b6d13122c9cbeeeb3d02448998621e7a8258ab7e6dd2f8bf18b80f","observation_id":"0aee1008-3364-4a96-970e-dc3bd3eb5095","resolution":{"observed_at":"2026-08-11T04:38:27.836564Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2303.08774","last_updated":"2024-03-04T06:01:33Z","snapshot_observed_at":"2026-08-07T07:30:12.213965Z","submitted_at":"2023-03-15T17:15:04Z","title":"GPT-4 Technical Report","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2303.08774","snapshot_observed_at":"2026-08-11T04:38:24.386208Z","title":"Gpt-4 technical report","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2412.18695","last_updated":"2024-12-24T22:51:29Z","snapshot_observed_at":"2026-08-11T23:49:11.223523Z","submitted_at":"2024-12-24T22:51:29Z","title":"TimelyLLM: Segmented LLM Serving System for Time-sensitive Robotic Applications","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-11T04:38:24.386208Z"},"links":{"cited_paper":"/paper/2303.08774","citing_paper":"/paper/2412.18695"},"observation_digest":"sha256:496f958a14b38363c02c0cd9e404d0e6372a45abdeda77df8dc4c84eeae1a106","observation_id":"dd836487-3e62-4897-addf-2675eafba803","resolution":{"observed_at":"2026-08-11T04:38:24.386208Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T04:38:27.677764Z","title":"Taming throughput-latency tradeoff in llm inference with sarathi-serve","venue":null,"work_id":"43b148a5-48f2-41fd-b178-e1dee0c7ac97","year":2024},"citing_paper":{"arxiv_id":"2412.18695","last_updated":"2024-12-24T22:51:29Z","snapshot_observed_at":"2026-08-11T23:49:11.223523Z","submitted_at":"2024-12-24T22:51:29Z","title":"TimelyLLM: Segmented LLM Serving System for Time-sensitive Robotic Applications","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-11T04:38:24.447658Z"},"links":{"citing_paper":"/paper/2412.18695"},"observation_digest":"sha256:671099ff06aadf821565e75a03b4f8ee8d9998113001b1d5cba00284fdf5d0b0","observation_id":"0b195bd6-3279-4460-9867-4f323721e703","resolution":{"observed_at":"2026-08-11T04:38:27.770990Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T04:38:27.614175Z","title":"Utility accrual real-time scheduling under variable cost functions","venue":null,"work_id":"e1b666c1-9c36-4957-a98a-7be1250b8093","year":2007},"citing_paper":{"arxiv_id":"2412.18695","last_updated":"2024-12-24T22:51:29Z","snapshot_observed_at":"2026-08-11T23:49:11.223523Z","submitted_at":"2024-12-24T22:51:29Z","title":"TimelyLLM: Segmented LLM Serving System for Time-sensitive Robotic Applications","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-11T04:38:24.484523Z"},"links":{"citing_paper":"/paper/2412.18695"},"observation_digest":"sha256:9d2d246dc6544b258c411495efbba6f5e4de9f87d4d772de0f22e5cf4fc18192","observation_id":"aaf54535-a78c-4272-8ee4-e5826e64771a","resolution":{"observed_at":"2026-08-11T04:38:27.628700Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T04:38:27.589224Z","title":"An evaluation model for information distribution in multi-robot systems","venue":null,"work_id":"a18a2df0-9234-4f79-9934-a721b5460eab","year":2019},"citing_paper":{"arxiv_id":"2412.18695","last_updated":"2024-12-24T22:51:29Z","snapshot_observed_at":"2026-08-11T23:49:11.223523Z","submitted_at":"2024-12-24T22:51:29Z","title":"TimelyLLM: Segmented LLM Serving System for Time-sensitive Robotic Applications","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-11T04:38:24.565139Z"},"links":{"citing_paper":"/paper/2412.18695"},"observation_digest":"sha256:d73d2f4cb8b091ca1d11fd178454d383a17ca78f3673f54a32010ca25332af55","observation_id":"68f8db8f-3032-4d1c-94b5-b78888e8a17f","resolution":{"observed_at":"2026-08-11T04:38:27.602025Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2005.14165","last_updated":"2020-07-22T19:47:17Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2020-05-28T17:29:03Z","title":"Language Models are Few-Shot Learners","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2005.14165","snapshot_observed_at":"2026-08-11T04:38:24.574581Z","title":"Language models are few-shot learners","venue":null,"work_id":null,"year":2005},"citing_paper":{"arxiv_id":"2412.18695","last_updated":"2024-12-24T22:51:29Z","snapshot_observed_at":"2026-08-11T23:49:11.223523Z","submitted_at":"2024-12-24T22:51:29Z","title":"TimelyLLM: Segmented LLM Serving System for Time-sensitive Robotic Applications","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-11T04:38:24.574581Z"},"links":{"cited_paper":"/paper/2005.14165","citing_paper":"/paper/2412.18695"},"observation_digest":"sha256:86274f504ea629ea9b55b25118b9662253a812b670e25b18a9dc6b66eee6701b","observation_id":"0c40ae0f-dee3-443b-a33c-99f958cbf217","resolution":{"observed_at":"2026-08-11T04:38:24.574581Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T04:38:27.575450Z","title":"Drone detection using depth maps","venue":null,"work_id":"e7f066dd-28a2-4723-a681-ace9eb96d40a","year":2018},"citing_paper":{"arxiv_id":"2412.18695","last_updated":"2024-12-24T22:51:29Z","snapshot_observed_at":"2026-08-11T23:49:11.223523Z","submitted_at":"2024-12-24T22:51:29Z","title":"TimelyLLM: Segmented LLM Serving System for Time-sensitive Robotic Applications","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-11T04:38:24.587951Z"},"links":{"citing_paper":"/paper/2412.18695"},"observation_digest":"sha256:f0bcce601cd3d44033754ed05fa007f458c4d17cf2e2e46080e106b34d0de1b6","observation_id":"6960a155-4972-4932-84b3-85a834670797","resolution":{"observed_at":"2026-08-11T04:38:27.581382Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2312.14950","last_updated":"2024-09-26T15:45:13Z","snapshot_observed_at":"2026-08-12T05:17:36.832579Z","submitted_at":"2023-12-08T15:57:18Z","title":"TypeFly: Flying Drones with Large Language Model","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.14950","snapshot_observed_at":"2026-08-11T04:38:24.608098Z","title":"Type- fly: Flying drones with large language model","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.18695","last_updated":"2024-12-24T22:51:29Z","snapshot_observed_at":"2026-08-11T23:49:11.223523Z","submitted_at":"2024-12-24T22:51:29Z","title":"TimelyLLM: Segmented LLM Serving System for Time-sensitive Robotic Applications","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-11T04:38:24.608098Z"},"links":{"cited_paper":"/paper/2312.14950","citing_paper":"/paper/2412.18695"},"observation_digest":"sha256:400c7ebca7343bc584de101e4a1673127d48d8cdae48fc3ec2a900f5ec9de3bc","observation_id":"25888fd1-9ea5-4c5a-888d-b9a38d38729a","resolution":{"observed_at":"2026-08-11T04:38:24.608098Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T04:38:27.555232Z","title":"A scheduling algorithm for tasks described by time value function","venue":null,"work_id":"3a421d3a-f571-4f1d-b373-98442a42ca27","year":1996},"citing_paper":{"arxiv_id":"2412.18695","last_updated":"2024-12-24T22:51:29Z","snapshot_observed_at":"2026-08-11T23:49:11.223523Z","submitted_at":"2024-12-24T22:51:29Z","title":"TimelyLLM: Segmented LLM Serving System for Time-sensitive Robotic Applications","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-11T04:38:24.612490Z"},"links":{"citing_paper":"/paper/2412.18695"},"observation_digest":"sha256:4c8f5b14f1fd2257550e75ddfab7ab535f30d936b8f850df3f7496b944ebd230","observation_id":"ecbe4b49-29e5-485f-a789-d190aa04502a","resolution":{"observed_at":"2026-08-11T04:38:27.563306Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T04:38:27.533016Z","title":"Robots that can chat","venue":null,"work_id":"843d08f6-d32e-446d-b8b1-609c4ae6559b","year":2024},"citing_paper":{"arxiv_id":"2412.18695","last_updated":"2024-12-24T22:51:29Z","snapshot_observed_at":"2026-08-11T23:49:11.223523Z","submitted_at":"2024-12-24T22:51:29Z","title":"TimelyLLM: Segmented LLM Serving System for Time-sensitive Robotic Applications","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-11T04:38:24.616334Z"},"links":{"citing_paper":"/paper/2412.18695"},"observation_digest":"sha256:a167d92b6a1786f3ed532461bbcdf10e955027a3cb48269d4e2153b4fb04afbd","observation_id":"b78ca992-48c8-4396-8f3a-d6d9c39cc375","resolution":{"observed_at":"2026-08-11T04:38:27.540107Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T04:38:27.379266Z","title":"Number of parameters in gpt-4","venue":null,"work_id":"80f9cb19-7fed-4a61-a66b-3d635f35efa1","year":2024},"citing_paper":{"arxiv_id":"2412.18695","last_updated":"2024-12-24T22:51:29Z","snapshot_observed_at":"2026-08-11T23:49:11.223523Z","submitted_at":"2024-12-24T22:51:29Z","title":"TimelyLLM: Segmented LLM Serving System for Time-sensitive Robotic Applications","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-11T04:38:24.620187Z"},"links":{"citing_paper":"/paper/2412.18695"},"observation_digest":"sha256:79fa86e1f1650cf2369792409222c3fd67f9d966f2a537adf1c0b83fe728902e","observation_id":"585e45ac-466f-4b8a-a364-60eb64186d6d","resolution":{"observed_at":"2026-08-11T04:38:27.487427Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T04:38:27.247483Z","title":"Transformers: State-of-the-art machine learning for pytorch, tensorflow, and jax","venue":null,"work_id":"28359e00-9bae-41bc-a6d4-5e544f7c68fa","year":2024},"citing_paper":{"arxiv_id":"2412.18695","last_updated":"2024-12-24T22:51:29Z","snapshot_observed_at":"2026-08-11T23:49:11.223523Z","submitted_at":"2024-12-24T22:51:29Z","title":"TimelyLLM: Segmented LLM Serving System for Time-sensitive Robotic Applications","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-11T04:38:24.623493Z"},"links":{"citing_paper":"/paper/2412.18695"},"observation_digest":"sha256:e8ef728f853017f5e12da5e877da4044c49fe370e60696132991b331b3b2f17d","observation_id":"e34ba0eb-10ed-4f8e-899d-8a047456c567","resolution":{"observed_at":"2026-08-11T04:38:27.311267Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1805.04833","last_updated":"2018-05-13T07:07:08Z","snapshot_observed_at":"2026-07-06T06:38:48.758485Z","submitted_at":"2018-05-13T07:07:08Z","title":"Hierarchical Neural Story Generation","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1805.04833","snapshot_observed_at":"2026-08-11T04:38:24.627595Z","title":"Hierarchical neural story generation","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2412.18695","last_updated":"2024-12-24T22:51:29Z","snapshot_observed_at":"2026-08-11T23:49:11.223523Z","submitted_at":"2024-12-24T22:51:29Z","title":"TimelyLLM: Segmented LLM Serving System for Time-sensitive Robotic Applications","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-11T04:38:24.627595Z"},"links":{"cited_paper":"/paper/1805.04833","citing_paper":"/paper/2412.18695"},"observation_digest":"sha256:967aa2550908c705c2101d05b1890484f2db39222d96f59a0e7ccd253aa80907","observation_id":"d8f2c8b7-6ac5-4b75-8ed5-fe8dedcd5281","resolution":{"observed_at":"2026-08-11T04:38:24.627595Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T04:38:27.212202Z","title":"Figure + openai allow speech-to-speech reasoning over learned behaviors","venue":null,"work_id":"4c58f0e1-90a5-492e-a59b-763e2adffd6f","year":2024},"citing_paper":{"arxiv_id":"2412.18695","last_updated":"2024-12-24T22:51:29Z","snapshot_observed_at":"2026-08-11T23:49:11.223523Z","submitted_at":"2024-12-24T22:51:29Z","title":"TimelyLLM: Segmented LLM Serving System for Time-sensitive Robotic Applications","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-11T04:38:24.637982Z"},"links":{"citing_paper":"/paper/2412.18695"},"observation_digest":"sha256:5b347de4c2d315d76c6a533b0db4f0a39b9531567b34eca897189e478d068508","observation_id":"55b13b88-5217-4d8a-abda-8393eeed75fa","resolution":{"observed_at":"2026-08-11T04:38:27.216278Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2408.15792","last_updated":"2024-08-28T13:35:54Z","snapshot_observed_at":"2026-08-10T07:06:10.343215Z","submitted_at":"2024-08-28T13:35:54Z","title":"Efficient LLM Scheduling by Learning to Rank","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2408.15792","snapshot_observed_at":"2026-08-11T04:38:24.653256Z","title":"Efficient llm scheduling by learning to rank","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.18695","last_updated":"2024-12-24T22:51:29Z","snapshot_observed_at":"2026-08-11T23:49:11.223523Z","submitted_at":"2024-12-24T22:51:29Z","title":"TimelyLLM: Segmented LLM Serving System for Time-sensitive Robotic Applications","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-11T04:38:24.653256Z"},"links":{"cited_paper":"/paper/2408.15792","citing_paper":"/paper/2412.18695"},"observation_digest":"sha256:b4773410048b32a19587640dd21abaf64bc7588d9118253d1a017d06385057ec","observation_id":"f3b5f765-da64-478d-b8b9-c2bb9a5cea34","resolution":{"observed_at":"2026-08-11T04:38:24.653256Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T04:38:27.201400Z","title":"Prompt cache: Modular attention reuse for low-latency inference","venue":null,"work_id":"fc421c2d-4329-459e-9abc-0ebd879e151a","year":2024},"citing_paper":{"arxiv_id":"2412.18695","last_updated":"2024-12-24T22:51:29Z","snapshot_observed_at":"2026-08-11T23:49:11.223523Z","submitted_at":"2024-12-24T22:51:29Z","title":"TimelyLLM: Segmented LLM Serving System for Time-sensitive Robotic Applications","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-11T04:38:24.676877Z"},"links":{"citing_paper":"/paper/2412.18695"},"observation_digest":"sha256:49fb238eb12c14b6270c6a1888c297e53754a59c8aa35df817f5720dbc2587f4","observation_id":"298ddeae-864a-4d2d-8a4d-261e569a04c3","resolution":{"observed_at":"2026-08-11T04:38:27.205256Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.21276","last_updated":"2024-10-25T17:43:01Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-10-25T17:43:01Z","title":"GPT-4o System Card","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.21276","snapshot_observed_at":"2026-08-11T04:38:24.724781Z","title":"Gpt-4o system card","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2412.18695","last_updated":"2024-12-24T22:51:29Z","snapshot_observed_at":"2026-08-11T23:49:11.223523Z","submitted_at":"2024-12-24T22:51:29Z","title":"TimelyLLM: Segmented LLM Serving System for Time-sensitive Robotic Applications","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-11T04:38:24.724781Z"},"links":{"cited_paper":"/paper/2410.21276","citing_paper":"/paper/2412.18695"},"observation_digest":"sha256:6e3f0b0f05803086ff7993ac02ba0580b822738c053293a5ccc8cf251fda5b3d","observation_id":"fac64749-c89e-4a29-bd29-5fc9c2ab5f2b","resolution":{"observed_at":"2026-08-11T04:38:24.724781Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T04:38:27.188946Z","title":"A time- driven scheduling model for real-time operating systems","venue":null,"work_id":"a7e0209b-8a55-4018-b479-732b590e16ca","year":1985},"citing_paper":{"arxiv_id":"2412.18695","last_updated":"2024-12-24T22:51:29Z","snapshot_observed_at":"2026-08-11T23:49:11.223523Z","submitted_at":"2024-12-24T22:51:29Z","title":"TimelyLLM: Segmented LLM Serving System for Time-sensitive Robotic Applications","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-11T04:38:24.766307Z"},"links":{"citing_paper":"/paper/2412.18695"},"observation_digest":"sha256:8347c566eb3f578ea210c21f0dc2aee3052e8e0ce0995ed5d8d2d5f618d4a890","observation_id":"e9cff7bf-3c15-45fa-ae34-2beb8418754e","resolution":{"observed_at":"2026-08-11T04:38:27.193263Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T04:38:27.165929Z","title":"Coedge: A co- operative edge system for distributed real-time deep learning tasks","venue":null,"work_id":"4bb3c0e6-b950-467c-86b3-3b5c09c43166","year":2023},"citing_paper":{"arxiv_id":"2412.18695","last_updated":"2024-12-24T22:51:29Z","snapshot_observed_at":"2026-08-11T23:49:11.223523Z","submitted_at":"2024-12-24T22:51:29Z","title":"TimelyLLM: Segmented LLM Serving System for Time-sensitive Robotic Applications","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-11T04:38:24.820179Z"},"links":{"citing_paper":"/paper/2412.18695"},"observation_digest":"sha256:4865d0dd40852533fc218f820ef182f4ba0430d85178760d245033811a4dc299","observation_id":"2f3eee56-fbb5-466e-8458-7189026e0a0e","resolution":{"observed_at":"2026-08-11T04:38:27.177228Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T04:38:27.148803Z","title":"𝑠3: Increasing gpu utilization during generative inference for higher throughput","venue":null,"work_id":"dc5fb7fc-a6f3-4e2c-90d0-56d7232b3531","year":2023},"citing_paper":{"arxiv_id":"2412.18695","last_updated":"2024-12-24T22:51:29Z","snapshot_observed_at":"2026-08-11T23:49:11.223523Z","submitted_at":"2024-12-24T22:51:29Z","title":"TimelyLLM: Segmented LLM Serving System for Time-sensitive Robotic Applications","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-11T04:38:24.870994Z"},"links":{"citing_paper":"/paper/2412.18695"},"observation_digest":"sha256:7d675f056cec936bd059ee6b35338607134f041063a577e46dc011c8aac2b074","observation_id":"43e01fc5-d675-465e-8e1a-1aee8c27e62e","resolution":{"observed_at":"2026-08-11T04:38:27.152609Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2312.04511","last_updated":"2024-06-05T03:53:10Z","snapshot_observed_at":"2026-08-10T17:07:28.597233Z","submitted_at":"2023-12-07T18:32:04Z","title":"An LLM Compiler for Parallel Function Calling","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.04511","snapshot_observed_at":"2026-08-11T04:38:24.901003Z","title":"An llm compiler for parallel function calling","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2412.18695","last_updated":"2024-12-24T22:51:29Z","snapshot_observed_at":"2026-08-11T23:49:11.223523Z","submitted_at":"2024-12-24T22:51:29Z","title":"TimelyLLM: Segmented LLM Serving System for Time-sensitive Robotic Applications","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-11T04:38:24.901003Z"},"links":{"cited_paper":"/paper/2312.04511","citing_paper":"/paper/2412.18695"},"observation_digest":"sha256:4855ddd118bc59414f629cd848bc747e021ca78e1841e0f1f4b04ed7c02ff36a","observation_id":"081ecc48-058f-4754-bec5-312bd046c9dc","resolution":{"observed_at":"2026-08-11T04:38:24.901003Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T04:38:24.920252Z","title":"Gonzalez, Hao Zhang, and Ion Stoica","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2412.18695","last_updated":"2024-12-24T22:51:29Z","snapshot_observed_at":"2026-08-11T23:49:11.223523Z","submitted_at":"2024-12-24T22:51:29Z","title":"TimelyLLM: Segmented LLM Serving System for Time-sensitive Robotic Applications","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-11T04:38:24.920252Z"},"links":{"citing_paper":"/paper/2412.18695"},"observation_digest":"sha256:db9b49196934edbcf89feb5cce264ed0d1c7fbf3e9dc1a7e130d2b4534fb316a","observation_id":"a61c66f3-5dc1-4a5b-80ed-df7299dd09ad","resolution":{"observed_at":"2026-08-11T04:38:24.920252Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T04:38:27.004160Z","title":"Mobilegpt: Augment- ing llm with human-like app memory for mobile task automation","venue":null,"work_id":"c7bc7b52-478e-4f0e-add8-683947719650","year":2024},"citing_paper":{"arxiv_id":"2412.18695","last_updated":"2024-12-24T22:51:29Z","snapshot_observed_at":"2026-08-11T23:49:11.223523Z","submitted_at":"2024-12-24T22:51:29Z","title":"TimelyLLM: Segmented LLM Serving System for Time-sensitive Robotic Applications","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-11T04:38:24.931709Z"},"links":{"citing_paper":"/paper/2412.18695"},"observation_digest":"sha256:845c8830da79d875d10a7ce8306108f58a66b900e943616a777d16afe9da2e15","observation_id":"f7ca6d68-ac17-4a0a-99a8-8a71c6413a0d","resolution":{"observed_at":"2026-08-11T04:38:27.086375Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T04:38:26.883962Z","title":"A utility accrual scheduling algorithm for real-time activities with mu- tual exclusion resource constraints","venue":null,"work_id":"d99d0a07-dc9d-4a89-9721-d825db5db44c","year":2006},"citing_paper":{"arxiv_id":"2412.18695","last_updated":"2024-12-24T22:51:29Z","snapshot_observed_at":"2026-08-11T23:49:11.223523Z","submitted_at":"2024-12-24T22:51:29Z","title":"TimelyLLM: Segmented LLM Serving System for Time-sensitive Robotic Applications","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-11T04:38:24.940518Z"},"links":{"citing_paper":"/paper/2412.18695"},"observation_digest":"sha256:6066dff59f1fc7dabb4debe0c571f2e76810dd1b4304911a530519432322736c","observation_id":"d80e1ed3-e31d-4cde-b746-0f0259d84b23","resolution":{"observed_at":"2026-08-11T04:38:26.964842Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2306.15724","last_updated":"2023-10-16T21:04:57Z","snapshot_observed_at":"2026-08-02T12:10:11.743351Z","submitted_at":"2023-06-27T18:03:15Z","title":"REFLECT: Summarizing Robot Experiences for Failure Explanation and Correction","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2306.15724","snapshot_observed_at":"2026-08-11T04:38:24.944589Z","title":"Reflect: Summarizing robot experiences for failure explanation and correction","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2412.18695","last_updated":"2024-12-24T22:51:29Z","snapshot_observed_at":"2026-08-11T23:49:11.223523Z","submitted_at":"2024-12-24T22:51:29Z","title":"TimelyLLM: Segmented LLM Serving System for Time-sensitive Robotic Applications","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-11T04:38:24.944589Z"},"links":{"cited_paper":"/paper/2306.15724","citing_paper":"/paper/2412.18695"},"observation_digest":"sha256:72a3d1b31a7894d1108395667a7c8734aa851fbec9013208484fd5af7be3d1c5","observation_id":"efdb42c8-0623-4f44-a707-e0771553472d","resolution":{"observed_at":"2026-08-11T04:38:24.944589Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T04:38:26.827159Z","title":"An example real-time command, control, and battle management application for alpha","venue":null,"work_id":"9e76228f-793f-4314-94d6-aa87b61ee11b","year":1988},"citing_paper":{"arxiv_id":"2412.18695","last_updated":"2024-12-24T22:51:29Z","snapshot_observed_at":"2026-08-11T23:49:11.223523Z","submitted_at":"2024-12-24T22:51:29Z","title":"TimelyLLM: Segmented LLM Serving System for Time-sensitive Robotic Applications","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-11T04:38:24.948919Z"},"links":{"citing_paper":"/paper/2412.18695"},"observation_digest":"sha256:0e46d0a28a23bfa61f36e4a41a2a3adac414224467ab59526aaabac456ca5730","observation_id":"bba860fa-66d2-45b7-9009-ebedc2b0f11c","resolution":{"observed_at":"2026-08-11T04:38:26.836295Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T04:38:26.797689Z","title":"The llama 3 herd of models","venue":null,"work_id":"6950295c-00d1-4924-89ce-26c8428725ee","year":2024},"citing_paper":{"arxiv_id":"2412.18695","last_updated":"2024-12-24T22:51:29Z","snapshot_observed_at":"2026-08-11T23:49:11.223523Z","submitted_at":"2024-12-24T22:51:29Z","title":"TimelyLLM: Segmented LLM Serving System for Time-sensitive Robotic Applications","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-11T04:38:24.957460Z"},"links":{"citing_paper":"/paper/2412.18695"},"observation_digest":"sha256:4ccd8f43138b3b6f06100855c85f6a23fb22a21d2ec15a3161a839f44eb4effa","observation_id":"ef50c6b3-049a-4522-83ac-d1f8bc0b38ed","resolution":{"observed_at":"2026-08-11T04:38:26.801059Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T04:38:26.764933Z","title":"Meta ai assistant built with llama 3","venue":null,"work_id":"360ecd7d-5b25-4bac-afde-8d906f9d9fd4","year":2024},"citing_paper":{"arxiv_id":"2412.18695","last_updated":"2024-12-24T22:51:29Z","snapshot_observed_at":"2026-08-11T23:49:11.223523Z","submitted_at":"2024-12-24T22:51:29Z","title":"TimelyLLM: Segmented LLM Serving System for Time-sensitive Robotic Applications","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-11T04:38:24.968585Z"},"links":{"citing_paper":"/paper/2412.18695"},"observation_digest":"sha256:e60ec81f358dbf4f5f5f6b6444802672b3a52df6a87253adf87e4646a3bbfe38","observation_id":"02d90645-8b53-4d0f-b5a8-bf00bbccae1f","resolution":{"observed_at":"2026-08-11T04:38:26.775272Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T04:38:26.746310Z","title":"Introducing meta llama 3: The most capable openly available llm to date","venue":null,"work_id":"81af8750-81bf-438b-a631-749df17411ad","year":2024},"citing_paper":{"arxiv_id":"2412.18695","last_updated":"2024-12-24T22:51:29Z","snapshot_observed_at":"2026-08-11T23:49:11.223523Z","submitted_at":"2024-12-24T22:51:29Z","title":"TimelyLLM: Segmented LLM Serving System for Time-sensitive Robotic Applications","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-11T04:38:24.975538Z"},"links":{"citing_paper":"/paper/2412.18695"},"observation_digest":"sha256:8edf66a7ad6863a9f4cad85d6e5c89f287168b5bc64efc7cadfb76614ef3ddc3","observation_id":"f5418bd6-6ed8-4275-b3f2-0de38cf0fa48","resolution":{"observed_at":"2026-08-11T04:38:26.753419Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T04:38:26.706664Z","title":"Llama 2 70B: An MLPerf Inference Benchmark for Large Language Models","venue":null,"work_id":"67b29ecc-9fcb-4e06-8b55-521fb185a998","year":2024},"citing_paper":{"arxiv_id":"2412.18695","last_updated":"2024-12-24T22:51:29Z","snapshot_observed_at":"2026-08-11T23:49:11.223523Z","submitted_at":"2024-12-24T22:51:29Z","title":"TimelyLLM: Segmented LLM Serving System for Time-sensitive Robotic Applications","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-11T04:38:24.979568Z"},"links":{"citing_paper":"/paper/2412.18695"},"observation_digest":"sha256:8c1e3a86ad987afac08a1a985fc35550a5cbac7ba983e1d8bd77d532e913fc46","observation_id":"3ccbaa84-3693-448d-8894-741444196932","resolution":{"observed_at":"2026-08-11T04:38:26.710242Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T04:38:26.698293Z","title":"Neuromeka indy","venue":null,"work_id":"66beb569-3aa2-4216-9a7b-682fa4c93009","year":2024},"citing_paper":{"arxiv_id":"2412.18695","last_updated":"2024-12-24T22:51:29Z","snapshot_observed_at":"2026-08-11T23:49:11.223523Z","submitted_at":"2024-12-24T22:51:29Z","title":"TimelyLLM: Segmented LLM Serving System for Time-sensitive Robotic Applications","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-11T04:38:24.983843Z"},"links":{"citing_paper":"/paper/2412.18695"},"observation_digest":"sha256:fdec3f7561191aa71563bc5b4d448dc2f64f0d430f53a2d1280fac4b2fb132e5","observation_id":"074bf257-205d-43c6-bdd3-558442656e2b","resolution":{"observed_at":"2026-08-11T04:38:26.701408Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T04:38:26.686198Z","title":"Tensorrt-llm","venue":null,"work_id":"72339dc6-7af9-4e0d-9cab-c65df24db792","year":2024},"citing_paper":{"arxiv_id":"2412.18695","last_updated":"2024-12-24T22:51:29Z","snapshot_observed_at":"2026-08-11T23:49:11.223523Z","submitted_at":"2024-12-24T22:51:29Z","title":"TimelyLLM: Segmented LLM Serving System for Time-sensitive Robotic Applications","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-11T04:38:24.987907Z"},"links":{"citing_paper":"/paper/2412.18695"},"observation_digest":"sha256:aeaf6172e97001d3e7032d915d0bf49df056024a212ec825fe8515f6efee9e0a","observation_id":"618f1dc9-a2a8-404f-8278-1b85c5e60eda","resolution":{"observed_at":"2026-08-11T04:38:26.691756Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T04:38:26.667894Z","title":"Exegpt: Constraint-aware resource scheduling for llm inference","venue":null,"work_id":"846e9d5e-8423-4e3a-ac5e-2de133d1cb58","year":2024},"citing_paper":{"arxiv_id":"2412.18695","last_updated":"2024-12-24T22:51:29Z","snapshot_observed_at":"2026-08-11T23:49:11.223523Z","submitted_at":"2024-12-24T22:51:29Z","title":"TimelyLLM: Segmented LLM Serving System for Time-sensitive Robotic Applications","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-11T04:38:24.992302Z"},"links":{"citing_paper":"/paper/2412.18695"},"observation_digest":"sha256:2d3483a7bb42c1f9ea84a4b9f72a44e53af02f47eac657c99af7f517f9dcf936","observation_id":"1a4ba292-3887-4393-a8be-041c9c5012eb","resolution":{"observed_at":"2026-08-11T04:38:26.680139Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2407.00047","last_updated":"2025-02-25T17:54:13Z","snapshot_observed_at":"2026-08-11T13:07:23.875983Z","submitted_at":"2024-06-05T21:17:34Z","title":"Queue management for slo-oriented large language model serving","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.00047","snapshot_observed_at":"2026-08-11T04:38:24.996442Z","title":"One queue is all you need: Resolving head-of-line blocking in large language model serving","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.18695","last_updated":"2024-12-24T22:51:29Z","snapshot_observed_at":"2026-08-11T23:49:11.223523Z","submitted_at":"2024-12-24T22:51:29Z","title":"TimelyLLM: Segmented LLM Serving System for Time-sensitive Robotic Applications","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-11T04:38:24.996442Z"},"links":{"cited_paper":"/paper/2407.00047","citing_paper":"/paper/2412.18695"},"observation_digest":"sha256:3133f6d3d97dfa4f9b21637ffb43a7842506a23cffb2f7980add8599170bb72b","observation_id":"f07e6141-ba24-4158-ac85-9556aacde1cb","resolution":{"observed_at":"2026-08-11T04:38:24.996442Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T04:38:26.651868Z","title":"Managing delays in human- robot interaction","venue":null,"work_id":"3b82b3b3-44ab-4f18-bd26-7aedc417a8aa","year":2023},"citing_paper":{"arxiv_id":"2412.18695","last_updated":"2024-12-24T22:51:29Z","snapshot_observed_at":"2026-08-11T23:49:11.223523Z","submitted_at":"2024-12-24T22:51:29Z","title":"TimelyLLM: Segmented LLM Serving System for Time-sensitive Robotic Applications","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-11T04:38:25.000726Z"},"links":{"citing_paper":"/paper/2412.18695"},"observation_digest":"sha256:d4781dcb4ce2966bd60fb3fbdfebad8df96af725d3ad71498a7e6c9a541f4acf","observation_id":"666d98ee-ddb4-4c64-b8cf-15ec1c8ccdc6","resolution":{"observed_at":"2026-08-11T04:38:26.656472Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T04:38:26.561789Z","title":"Efficiently scaling transformer inference","venue":null,"work_id":"bc210fea-5e54-41ac-853a-fcf972b0b518","year":2023},"citing_paper":{"arxiv_id":"2412.18695","last_updated":"2024-12-24T22:51:29Z","snapshot_observed_at":"2026-08-11T23:49:11.223523Z","submitted_at":"2024-12-24T22:51:29Z","title":"TimelyLLM: Segmented LLM Serving System for Time-sensitive Robotic Applications","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-11T04:38:25.005056Z"},"links":{"citing_paper":"/paper/2412.18695"},"observation_digest":"sha256:caca93c0c4dbe1f2cb043e14f8aa3ac12e99427461bee40a29f5124a65996cd4","observation_id":"7071a7fe-a514-4ba6-814e-9397c9703e09","resolution":{"observed_at":"2026-08-11T04:38:26.606296Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2404.08509","last_updated":"2024-11-25T17:35:07Z","snapshot_observed_at":"2026-08-09T09:17:59.811981Z","submitted_at":"2024-04-12T14:46:15Z","title":"Efficient Interactive LLM Serving with Proxy Model-based Sequence Length Prediction","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.08509","snapshot_observed_at":"2026-08-11T04:38:25.029210Z","title":"Efficient interactive llm serving with proxy model-based sequence length prediction","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.18695","last_updated":"2024-12-24T22:51:29Z","snapshot_observed_at":"2026-08-11T23:49:11.223523Z","submitted_at":"2024-12-24T22:51:29Z","title":"TimelyLLM: Segmented LLM Serving System for Time-sensitive Robotic Applications","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-11T04:38:25.029210Z"},"links":{"cited_paper":"/paper/2404.08509","citing_paper":"/paper/2412.18695"},"observation_digest":"sha256:06657fc3e2144b0a0c912cdae14d472f582109ce252679d39ea9cd63f9570ed8","observation_id":"d407896f-c54b-482a-ba84-0ef671fcfe52","resolution":{"observed_at":"2026-08-11T04:38:25.029210Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2307.06135","last_updated":"2023-09-27T23:17:28Z","snapshot_observed_at":"2026-07-06T15:53:11.310262Z","submitted_at":"2023-07-12T12:37:55Z","title":"SayPlan: Grounding Large Language Models using 3D Scene Graphs for Scalable Robot Task Planning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2307.06135","snapshot_observed_at":"2026-08-11T04:38:25.074571Z","title":"Sayplan: Grounding large language models using 3d scene graphs for scalable task planning.arXiv preprint arXiv:2307.06135, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2412.18695","last_updated":"2024-12-24T22:51:29Z","snapshot_observed_at":"2026-08-11T23:49:11.223523Z","submitted_at":"2024-12-24T22:51:29Z","title":"TimelyLLM: Segmented LLM Serving System for Time-sensitive Robotic Applications","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-11T04:38:25.074571Z"},"links":{"cited_paper":"/paper/2307.06135","citing_paper":"/paper/2412.18695"},"observation_digest":"sha256:39676460fb5ccf84bba171878efaffdb8fc2316b4fccaa0dc25a7e6510932bef","observation_id":"ab2a5bf0-36ef-4fa8-a527-f71db3dd0f19","resolution":{"observed_at":"2026-08-11T04:38:25.074571Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T04:38:26.450310Z","title":"Ravindran, E.D","venue":null,"work_id":"ff8fb07a-cab5-4053-acb8-04c753f1ee11","year":2005},"citing_paper":{"arxiv_id":"2412.18695","last_updated":"2024-12-24T22:51:29Z","snapshot_observed_at":"2026-08-11T23:49:11.223523Z","submitted_at":"2024-12-24T22:51:29Z","title":"TimelyLLM: Segmented LLM Serving System for Time-sensitive Robotic Applications","version":1},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-08-11T04:38:25.127352Z"},"links":{"citing_paper":"/paper/2412.18695"},"observation_digest":"sha256:a79ccdea176e416270ca8037ea6f5527869abb4014d6c16c55e9792d9ac6ccef","observation_id":"6f594cc3-1de5-4e98-b9f7-c7bd10736450","resolution":{"observed_at":"2026-08-11T04:38:26.518176Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2307.01928","last_updated":"2023-09-04T16:06:48Z","snapshot_observed_at":"2026-08-12T03:21:04.968416Z","submitted_at":"2023-07-04T21:25:12Z","title":"Robots That Ask For Help: Uncertainty Alignment for Large Language Model Planners","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2307.01928","snapshot_observed_at":"2026-08-11T04:38:25.173656Z","title":"Robots that ask for help: Uncertainty alignment for large language model planners","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2412.18695","last_updated":"2024-12-24T22:51:29Z","snapshot_observed_at":"2026-08-11T23:49:11.223523Z","submitted_at":"2024-12-24T22:51:29Z","title":"TimelyLLM: Segmented LLM Serving System for Time-sensitive Robotic Applications","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-11T04:38:25.173656Z"},"links":{"cited_paper":"/paper/2307.01928","citing_paper":"/paper/2412.18695"},"observation_digest":"sha256:1dff436075016fb1fa80a13a242c83618da274c846211574aea0d355b5d59fdf","observation_id":"03af85fc-1136-4001-bbfa-d85a3be6fcca","resolution":{"observed_at":"2026-08-11T04:38:25.173656Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T04:38:26.353785Z","title":"Average reading speed","venue":null,"work_id":"637e1546-7857-419c-bae1-e7e6057d5cc0","year":2024},"citing_paper":{"arxiv_id":"2412.18695","last_updated":"2024-12-24T22:51:29Z","snapshot_observed_at":"2026-08-11T23:49:11.223523Z","submitted_at":"2024-12-24T22:51:29Z","title":"TimelyLLM: Segmented LLM Serving System for Time-sensitive Robotic Applications","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-11T04:38:25.226187Z"},"links":{"citing_paper":"/paper/2412.18695"},"observation_digest":"sha256:4a905d39aec5657a035c4550c2e9d253f1fa0e32fab6f99395c7eb55d54d8a5c","observation_id":"adcd4f3d-3dd5-4fce-932b-d824da396f10","resolution":{"observed_at":"2026-08-11T04:38:26.375375Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T04:38:26.337120Z","title":"Revenue-driven scheduling in drone delivery networks with time-sensitive service level agreements","venue":null,"work_id":"f25757d0-4d7d-43b8-925a-5ffc85ab7061","year":2019},"citing_paper":{"arxiv_id":"2412.18695","last_updated":"2024-12-24T22:51:29Z","snapshot_observed_at":"2026-08-11T23:49:11.223523Z","submitted_at":"2024-12-24T22:51:29Z","title":"TimelyLLM: Segmented LLM Serving System for Time-sensitive Robotic Applications","version":1},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-11T04:38:25.246677Z"},"links":{"citing_paper":"/paper/2412.18695"},"observation_digest":"sha256:05de42b2c81179273438da7d16cee401bab7e3d3c02c4461905922bbf08e980b","observation_id":"189564b4-3e4e-4b9d-a134-53cb414100ee","resolution":{"observed_at":"2026-08-11T04:38:26.340384Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.01035","last_updated":"2024-10-01T19:51:07Z","snapshot_observed_at":"2026-08-12T04:34:58.964146Z","submitted_at":"2024-10-01T19:51:07Z","title":"Don't Stop Me Now: Embedding Based Scheduling for LLMs","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.01035","snapshot_observed_at":"2026-08-11T04:38:25.260335Z","title":"Don’t stop me now: Embedding based scheduling for llms","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.18695","last_updated":"2024-12-24T22:51:29Z","snapshot_observed_at":"2026-08-11T23:49:11.223523Z","submitted_at":"2024-12-24T22:51:29Z","title":"TimelyLLM: Segmented LLM Serving System for Time-sensitive Robotic Applications","version":1},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-11T04:38:25.260335Z"},"links":{"cited_paper":"/paper/2410.01035","citing_paper":"/paper/2412.18695"},"observation_digest":"sha256:f9caf3003168ec85192ae1b25b4411b4d867cab803f2586a766b2b2f218bfc4e","observation_id":"68210787-0793-4d55-ab2e-a2d45652f3d9","resolution":{"observed_at":"2026-08-11T04:38:25.260335Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T04:38:25.270207Z","title":"Hugginggpt: Solving ai tasks with chatgpt and its friends in hugging face","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.18695","last_updated":"2024-12-24T22:51:29Z","snapshot_observed_at":"2026-08-11T23:49:11.223523Z","submitted_at":"2024-12-24T22:51:29Z","title":"TimelyLLM: Segmented LLM Serving System for Time-sensitive Robotic Applications","version":1},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-11T04:38:25.270207Z"},"links":{"citing_paper":"/paper/2412.18695"},"observation_digest":"sha256:30d40ba4a392468dd5f750880c6c7286eb67df13c320d99ee75b5479fa9df9ae","observation_id":"73add5c7-41bf-41dc-8441-e310f172f34b","resolution":{"observed_at":"2026-08-11T04:38:25.270207Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T04:38:26.314541Z","title":"Response time and display rate in human perfor- mance with computers","venue":null,"work_id":"d02bb98d-a010-4ccd-9af6-851af628a3f7","year":1984},"citing_paper":{"arxiv_id":"2412.18695","last_updated":"2024-12-24T22:51:29Z","snapshot_observed_at":"2026-08-11T23:49:11.223523Z","submitted_at":"2024-12-24T22:51:29Z","title":"TimelyLLM: Segmented LLM Serving System for Time-sensitive Robotic Applications","version":1},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-11T04:38:25.279474Z"},"links":{"citing_paper":"/paper/2412.18695"},"observation_digest":"sha256:e34ad1907732a13f37e79dd6cdaa737c7bf97ad6b85ee6e1e0331aab623d94c8","observation_id":"b7e64dd8-5970-4aa4-bd03-ae62b73acada","resolution":{"observed_at":"2026-08-11T04:38:26.323463Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2303.14100","last_updated":"2023-03-24T16:06:11Z","snapshot_observed_at":"2026-08-09T06:19:35.380893Z","submitted_at":"2023-03-24T16:06:11Z","title":"Errors are Useful Prompts: Instruction Guided Task Programming with Verifier-Assisted Iterative Prompting","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2303.14100","snapshot_observed_at":"2026-08-11T04:38:25.285326Z","title":"Errors are useful prompts: Instruction guided task programming with verifier-assisted iterative prompting","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2412.18695","last_updated":"2024-12-24T22:51:29Z","snapshot_observed_at":"2026-08-11T23:49:11.223523Z","submitted_at":"2024-12-24T22:51:29Z","title":"TimelyLLM: Segmented LLM Serving System for Time-sensitive Robotic Applications","version":1},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-11T04:38:25.285326Z"},"links":{"cited_paper":"/paper/2303.14100","citing_paper":"/paper/2412.18695"},"observation_digest":"sha256:1b2261ccc6ef9825921fd1c95756f196ec5646b2be95e084a4fe52310bc2f044","observation_id":"8acdd9c5-ed87-4dd8-9d3b-06e3de664305","resolution":{"observed_at":"2026-08-11T04:38:25.285326Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T04:38:26.302012Z","title":"Tello sdk user guide, 2023","venue":null,"work_id":"c4ad2b6d-1020-46b4-bca2-47926b978622","year":2023},"citing_paper":{"arxiv_id":"2412.18695","last_updated":"2024-12-24T22:51:29Z","snapshot_observed_at":"2026-08-11T23:49:11.223523Z","submitted_at":"2024-12-24T22:51:29Z","title":"TimelyLLM: Segmented LLM Serving System for Time-sensitive Robotic Applications","version":1},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-11T04:38:25.296865Z"},"links":{"citing_paper":"/paper/2412.18695"},"observation_digest":"sha256:494aba26c8a75fb3d0f79e10cba364b5128a58a9a226a06e28e3b24d44e58183","observation_id":"5fe17180-805a-4a5c-82c6-a64ab40bbcbd","resolution":{"observed_at":"2026-08-11T04:38:26.305980Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T04:38:26.276684Z","title":"Tidwell, R","venue":null,"work_id":"aa6a2883-8c4f-486c-a7fb-3c042f4eb97a","year":2010},"citing_paper":{"arxiv_id":"2412.18695","last_updated":"2024-12-24T22:51:29Z","snapshot_observed_at":"2026-08-11T23:49:11.223523Z","submitted_at":"2024-12-24T22:51:29Z","title":"TimelyLLM: Segmented LLM Serving System for Time-sensitive Robotic Applications","version":1},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-08-11T04:38:25.309491Z"},"links":{"citing_paper":"/paper/2412.18695"},"observation_digest":"sha256:670fb2cebdc12857062f6448e5e82b9a3f1cf9fa8ad8a61572d9dc77beb3abcf","observation_id":"60cf42c0-46e3-4f84-a1da-c7457c6f731f","resolution":{"observed_at":"2026-08-11T04:38:26.285433Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T04:38:26.245089Z","title":"Optimizing expected time utility in cyber-physical systems schedulers","venue":null,"work_id":"5f1b01d1-29cb-40cd-bb2b-70bfa00e7b9c","year":2010},"citing_paper":{"arxiv_id":"2412.18695","last_updated":"2024-12-24T22:51:29Z","snapshot_observed_at":"2026-08-11T23:49:11.223523Z","submitted_at":"2024-12-24T22:51:29Z","title":"TimelyLLM: Segmented LLM Serving System for Time-sensitive Robotic Applications","version":1},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-08-11T04:38:25.324476Z"},"links":{"citing_paper":"/paper/2412.18695"},"observation_digest":"sha256:61d020bc07bb569223da9e51580c78810c18c874afae297f99054d9a1f2e9f3d","observation_id":"641bd73b-b0b4-4ce9-94f5-eb741e2e057b","resolution":{"observed_at":"2026-08-11T04:38:26.254036Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T04:38:25.333152Z","title":"Chatgpt for robotics: Design principles and model abilities","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.18695","last_updated":"2024-12-24T22:51:29Z","snapshot_observed_at":"2026-08-11T23:49:11.223523Z","submitted_at":"2024-12-24T22:51:29Z","title":"TimelyLLM: Segmented LLM Serving System for Time-sensitive Robotic Applications","version":1},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-08-11T04:38:25.333152Z"},"links":{"citing_paper":"/paper/2412.18695"},"observation_digest":"sha256:60462a3705ca692cbb8eca156687c49413ad85afd4b85bf5f409efd200ea5837","observation_id":"bd6a5605-a431-4b60-969e-5fc87d3bc5bd","resolution":{"observed_at":"2026-08-11T04:38:25.333152Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2403.11552","last_updated":"2024-08-21T09:46:35Z","snapshot_observed_at":"2026-08-10T16:27:27.347496Z","submitted_at":"2024-03-18T08:03:47Z","title":"LLM3:Large Language Model-based Task and Motion Planning with Motion Failure Reasoning","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.11552","snapshot_observed_at":"2026-08-11T04:38:25.337040Z","title":"Llmˆ 3: Large language model-based task and motion planning with motion failure reasoning","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.18695","last_updated":"2024-12-24T22:51:29Z","snapshot_observed_at":"2026-08-11T23:49:11.223523Z","submitted_at":"2024-12-24T22:51:29Z","title":"TimelyLLM: Segmented LLM Serving System for Time-sensitive Robotic Applications","version":1},"reference_index":53,"source":"pdf_text","source_observed_at":"2026-08-11T04:38:25.337040Z"},"links":{"cited_paper":"/paper/2403.11552","citing_paper":"/paper/2412.18695"},"observation_digest":"sha256:ce31c2b24bb30a790247084c075918621f869184a483c63d7fcca7660f1bc1a8","observation_id":"cfd1c4c5-107c-49e1-86eb-b704ad565627","resolution":{"observed_at":"2026-08-11T04:38:25.337040Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T04:38:26.205403Z","title":"Autodroid: Llm-powered task automation in android","venue":null,"work_id":"07548fad-4580-4f1a-9b52-1164155666c4","year":2024},"citing_paper":{"arxiv_id":"2412.18695","last_updated":"2024-12-24T22:51:29Z","snapshot_observed_at":"2026-08-11T23:49:11.223523Z","submitted_at":"2024-12-24T22:51:29Z","title":"TimelyLLM: Segmented LLM Serving System for Time-sensitive Robotic Applications","version":1},"reference_index":54,"source":"pdf_text","source_observed_at":"2026-08-11T04:38:25.357810Z"},"links":{"citing_paper":"/paper/2412.18695"},"observation_digest":"sha256:51bd96560862c627ce85432ee77f88d89a687ca55c4bc94812959c98fd5e65d1","observation_id":"c62736ec-0002-4176-aa75-9fb2af639824","resolution":{"observed_at":"2026-08-11T04:38:26.209410Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T04:38:26.166438Z","title":"Time-utility function — Wikipedia, 2024","venue":null,"work_id":"b6db63ed-89d1-4ce3-80a7-c8e65f2c9ef5","year":2024},"citing_paper":{"arxiv_id":"2412.18695","last_updated":"2024-12-24T22:51:29Z","snapshot_observed_at":"2026-08-11T23:49:11.223523Z","submitted_at":"2024-12-24T22:51:29Z","title":"TimelyLLM: Segmented LLM Serving System for Time-sensitive Robotic Applications","version":1},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-08-11T04:38:25.396599Z"},"links":{"citing_paper":"/paper/2412.18695"},"observation_digest":"sha256:f72afb9c141c7e9a61067924e836c6a68110eb42908894f9c7c1ff0cecc00c0e","observation_id":"60f7203e-25b4-4c43-aed9-380f931f13e6","resolution":{"observed_at":"2026-08-11T04:38:26.174380Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2305.05920","last_updated":"2024-09-25T05:57:51Z","snapshot_observed_at":"2026-08-09T14:08:47.340904Z","submitted_at":"2023-05-10T06:17:50Z","title":"Fast Distributed Inference Serving for Large Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.05920","snapshot_observed_at":"2026-08-11T04:38:25.427260Z","title":"Fast distributed inference serving for large language models","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2412.18695","last_updated":"2024-12-24T22:51:29Z","snapshot_observed_at":"2026-08-11T23:49:11.223523Z","submitted_at":"2024-12-24T22:51:29Z","title":"TimelyLLM: Segmented LLM Serving System for Time-sensitive Robotic Applications","version":1},"reference_index":56,"source":"pdf_text","source_observed_at":"2026-08-11T04:38:25.427260Z"},"links":{"cited_paper":"/paper/2305.05920","citing_paper":"/paper/2412.18695"},"observation_digest":"sha256:4386ac4803d1ec44bb049fe88c0d629474fe698a044df76358fa60472c1c3c08","observation_id":"39ed668f-4302-4242-9041-26e49bacdcbf","resolution":{"observed_at":"2026-08-11T04:38:25.427260Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T04:38:26.054839Z","title":"Utility accrual scheduling under arbitrary time/utility functions and multi-unit resource constraints","venue":null,"work_id":"d82a062f-9026-4e1f-8f25-01819646ecfb","year":2004},"citing_paper":{"arxiv_id":"2412.18695","last_updated":"2024-12-24T22:51:29Z","snapshot_observed_at":"2026-08-11T23:49:11.223523Z","submitted_at":"2024-12-24T22:51:29Z","title":"TimelyLLM: Segmented LLM Serving System for Time-sensitive Robotic Applications","version":1},"reference_index":57,"source":"pdf_text","source_observed_at":"2026-08-11T04:38:25.453147Z"},"links":{"citing_paper":"/paper/2412.18695"},"observation_digest":"sha256:c98af1ca7cf46491782f2110f534778d0f3349df4ef352897eb5eb28ccfa9084","observation_id":"40058a5d-c392-4189-88e8-91c0bc70e301","resolution":{"observed_at":"2026-08-11T04:38:26.111952Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T04:38:25.499266Z","title":"Orca: A distributed serving system for {Transformer-Based} generative models","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2412.18695","last_updated":"2024-12-24T22:51:29Z","snapshot_observed_at":"2026-08-11T23:49:11.223523Z","submitted_at":"2024-12-24T22:51:29Z","title":"TimelyLLM: Segmented LLM Serving System for Time-sensitive Robotic Applications","version":1},"reference_index":58,"source":"pdf_text","source_observed_at":"2026-08-11T04:38:25.499266Z"},"links":{"citing_paper":"/paper/2412.18695"},"observation_digest":"sha256:f90bfd9e53cdb2f74b75d2d2cbb4c8fa7b924220dcb4d9303cc2dfdd8f4617c2","observation_id":"cb0dde80-536e-49f3-bd9d-6f46ab2a5349","resolution":{"observed_at":"2026-08-11T04:38:25.499266Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T04:38:25.965450Z","title":"{SHEPHERD}: Serving{DNNs} in the wild","venue":null,"work_id":"2b3a7dc3-37fe-42d8-a93b-4aad31feff6c","year":2023},"citing_paper":{"arxiv_id":"2412.18695","last_updated":"2024-12-24T22:51:29Z","snapshot_observed_at":"2026-08-11T23:49:11.223523Z","submitted_at":"2024-12-24T22:51:29Z","title":"TimelyLLM: Segmented LLM Serving System for Time-sensitive Robotic Applications","version":1},"reference_index":59,"source":"pdf_text","source_observed_at":"2026-08-11T04:38:25.534928Z"},"links":{"citing_paper":"/paper/2412.18695"},"observation_digest":"sha256:9e6790d7c728e6330a91fec6a78b8aa99ad025abe1c20948f3226953d078ee02","observation_id":"8dd46fec-4422-40e8-a7a5-12061308528c","resolution":{"observed_at":"2026-08-11T04:38:25.994566Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.10021","last_updated":"2023-10-17T12:01:17Z","snapshot_observed_at":"2026-07-06T16:33:30.170449Z","submitted_at":"2023-10-16T02:43:47Z","title":"Bootstrap Your Own Skills: Learning to Solve New Tasks with Large Language Model Guidance","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.10021","snapshot_observed_at":"2026-08-11T04:38:25.558727Z","title":"Bootstrap your own skills: Learning to solve new tasks with large language model guidance","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2412.18695","last_updated":"2024-12-24T22:51:29Z","snapshot_observed_at":"2026-08-11T23:49:11.223523Z","submitted_at":"2024-12-24T22:51:29Z","title":"TimelyLLM: Segmented LLM Serving System for Time-sensitive Robotic Applications","version":1},"reference_index":60,"source":"pdf_text","source_observed_at":"2026-08-11T04:38:25.558727Z"},"links":{"cited_paper":"/paper/2310.10021","citing_paper":"/paper/2412.18695"},"observation_digest":"sha256:186db38322cf75deb3278ae9f58b3fece0ab1997358c7030f6400eb3095551d2","observation_id":"8b4a0efb-aace-4fd7-bd4b-de69f139f743","resolution":{"observed_at":"2026-08-11T04:38:25.558727Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-11T04:38:25.926772Z","title":"Response length perception and sequence scheduling: An llm-empowered llm inference pipeline.Advances in Neural Information Processing Systems, 36, 2024","venue":null,"work_id":"ff68964f-fd56-48e1-b79b-977d9a41e49a","year":2024},"citing_paper":{"arxiv_id":"2412.18695","last_updated":"2024-12-24T22:51:29Z","snapshot_observed_at":"2026-08-11T23:49:11.223523Z","submitted_at":"2024-12-24T22:51:29Z","title":"TimelyLLM: Segmented LLM Serving System for Time-sensitive Robotic Applications","version":1},"reference_index":61,"source":"pdf_text","source_observed_at":"2026-08-11T04:38:25.630673Z"},"links":{"citing_paper":"/paper/2412.18695"},"observation_digest":"sha256:4fbb5101539a31cf1501ab864323bda5a0ccb304736e67d5913611253740d094","observation_id":"63756280-cb92-4452-b8d8-8f7b72a0b5fc","resolution":{"observed_at":"2026-08-11T04:38:25.932349Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2412.18695","last_updated":"2024-12-24T22:51:29Z","latest_version":1,"primary_category":"cs.RO","snapshot_observed_at":"2026-08-11T23:49:11.223523Z","submitted_at":"2024-12-24T22:51:29Z","title":"TimelyLLM: Segmented LLM Serving System for Time-sensitive Robotic Applications"},"reference_resolution":{"displayed":61,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":22,"verified_exact":0,"verified_fuzzy":39},"total_outbound_references":61},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"thesis":"As of 12 August 2026, this Paper Citation Record lists 61 of 61 outbound references and 0 inbound Pith citation observations for arXiv:2412.18695."}