{"as_of":"2026-08-04T22:03:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:1889d3d6c24f4cc702f803fd6f1f9191fe971b6c942c9f954f8c39e61faaa70a","coverage":[{"denominator":48,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":48,"source":"paper_references, paper_reference_links","source_observed_at":"2026-05-20T02:10:57.582345Z","state":"measured"},{"denominator":48,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":48,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-04T06:34:03.388597+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2605.19481/citation-record","integrity":"/paper/2605.19481/integrity","json":"/paper/2605.19481/citation-record.json","paper":"/paper/2605.19481"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"115aa4d2-653d-4e65-b35a-e563679ad00a","year":null},"citing_paper":{"arxiv_id":"2605.19481","last_updated":"2026-05-19T07:34:08Z","snapshot_observed_at":"2026-08-02T01:46:46.351301Z","submitted_at":"2026-05-19T07:34:08Z","title":"C2CServe: Leveraging NVLink-C2C for Elastic Serverless LLM Serving on MIG","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-05-20T02:10:57.582345Z"},"links":{"citing_paper":"/paper/2605.19481"},"observation_digest":"sha256:c1c7ef03a2d4b1ba8e1e492e125bfb78771a619fa43927c191f0b7d990771a82","observation_id":"f5980d03-ca8d-4ab2-846f-b84f7f6c72a9","resolution":{"observed_at":"2026-05-20T02:12:58.948956Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"f159faac-d0e0-4024-bf58-3a294be3593b","year":null},"citing_paper":{"arxiv_id":"2605.19481","last_updated":"2026-05-19T07:34:08Z","snapshot_observed_at":"2026-08-02T01:46:46.351301Z","submitted_at":"2026-05-19T07:34:08Z","title":"C2CServe: Leveraging NVLink-C2C for Elastic Serverless LLM Serving on MIG","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-05-20T02:10:57.582345Z"},"links":{"citing_paper":"/paper/2605.19481"},"observation_digest":"sha256:56bef176b4a2f2c610e62dc20e1b4a48dd7bfd2ec5800909c42d3da804eed8ea","observation_id":"dcd468b4-cd41-4d27-88bf-79bb80c3a088","resolution":{"observed_at":"2026-05-20T02:12:59.030834Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"5046d6a4-2e4c-4d26-a549-c64de717a417","year":null},"citing_paper":{"arxiv_id":"2605.19481","last_updated":"2026-05-19T07:34:08Z","snapshot_observed_at":"2026-08-02T01:46:46.351301Z","submitted_at":"2026-05-19T07:34:08Z","title":"C2CServe: Leveraging NVLink-C2C for Elastic Serverless LLM Serving on MIG","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-05-20T02:10:57.582345Z"},"links":{"citing_paper":"/paper/2605.19481"},"observation_digest":"sha256:8e6c9d236fa238267a31a3a65bf42262fda9642442f3b7251007f5f1936c51e1","observation_id":"eb10be50-c8d0-4246-a1a7-100240ca58f3","resolution":{"observed_at":"2026-05-20T02:12:59.010585Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"4effd58b-9d81-43c9-a802-dea0310357d3","year":null},"citing_paper":{"arxiv_id":"2605.19481","last_updated":"2026-05-19T07:34:08Z","snapshot_observed_at":"2026-08-02T01:46:46.351301Z","submitted_at":"2026-05-19T07:34:08Z","title":"C2CServe: Leveraging NVLink-C2C for Elastic Serverless LLM Serving on MIG","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-05-20T02:10:57.582345Z"},"links":{"citing_paper":"/paper/2605.19481"},"observation_digest":"sha256:0e4e8756fb108aa33d4fa862726c55ad880951ad69068d00d3c68955214e021b","observation_id":"1af3945f-6239-4c07-b127-347e915af9e6","resolution":{"observed_at":"2026-05-20T02:12:59.074434Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"bcf0f55b-4f6b-46f9-a790-ed2516d425ac","year":null},"citing_paper":{"arxiv_id":"2605.19481","last_updated":"2026-05-19T07:34:08Z","snapshot_observed_at":"2026-08-02T01:46:46.351301Z","submitted_at":"2026-05-19T07:34:08Z","title":"C2CServe: Leveraging NVLink-C2C for Elastic Serverless LLM Serving on MIG","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-05-20T02:10:57.582345Z"},"links":{"citing_paper":"/paper/2605.19481"},"observation_digest":"sha256:8d421a2873347b0b9d77adee7379f0333f504424aa77d4ccf0dbf98240759b0e","observation_id":"7bc1e05e-2561-4925-a0a2-30b590e3b57a","resolution":{"observed_at":"2026-05-20T02:12:58.952931Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"9b117c8b-b161-4b24-9377-d437493db6a1","year":2022},"citing_paper":{"arxiv_id":"2605.19481","last_updated":"2026-05-19T07:34:08Z","snapshot_observed_at":"2026-08-02T01:46:46.351301Z","submitted_at":"2026-05-19T07:34:08Z","title":"C2CServe: Leveraging NVLink-C2C for Elastic Serverless LLM Serving on MIG","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-05-20T02:10:57.582345Z"},"links":{"citing_paper":"/paper/2605.19481"},"observation_digest":"sha256:b96d2ae4bab1e456ab67516b489ae1f2e30ba07697395df0feb6c383e9307b81","observation_id":"cfc20a51-cb30-4498-9ee8-d50b3a845661","resolution":{"observed_at":"2026-05-20T02:12:58.971958Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"2ac8f12d-be1e-4434-af39-84bcdf0ddc86","year":2022},"citing_paper":{"arxiv_id":"2605.19481","last_updated":"2026-05-19T07:34:08Z","snapshot_observed_at":"2026-08-02T01:46:46.351301Z","submitted_at":"2026-05-19T07:34:08Z","title":"C2CServe: Leveraging NVLink-C2C for Elastic Serverless LLM Serving on MIG","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-05-20T02:10:57.582345Z"},"links":{"citing_paper":"/paper/2605.19481"},"observation_digest":"sha256:ef142ceaccb12b297604c0d56f39036503c7a6eb3ab8656442d6e6e844353aba","observation_id":"ea277010-fcb3-4964-845c-54b960a4123f","resolution":{"observed_at":"2026-05-20T02:12:58.956956Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"cfb56cef-b1fd-4f18-8b0c-68cd34c6c3a1","year":2023},"citing_paper":{"arxiv_id":"2605.19481","last_updated":"2026-05-19T07:34:08Z","snapshot_observed_at":"2026-08-02T01:46:46.351301Z","submitted_at":"2026-05-19T07:34:08Z","title":"C2CServe: Leveraging NVLink-C2C for Elastic Serverless LLM Serving on MIG","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-05-20T02:10:57.582345Z"},"links":{"citing_paper":"/paper/2605.19481"},"observation_digest":"sha256:2804a42f77eeced5dbed985958080c1c6989a20ffcba118b7760c36c0cb6645e","observation_id":"62305bc0-e690-4634-9225-f936e5ae4491","resolution":{"observed_at":"2026-05-20T02:12:58.975423Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"8d263d5e-f22b-4cc1-bca3-bf18135c4d91","year":2023},"citing_paper":{"arxiv_id":"2605.19481","last_updated":"2026-05-19T07:34:08Z","snapshot_observed_at":"2026-08-02T01:46:46.351301Z","submitted_at":"2026-05-19T07:34:08Z","title":"C2CServe: Leveraging NVLink-C2C for Elastic Serverless LLM Serving on MIG","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-05-20T02:10:57.582345Z"},"links":{"citing_paper":"/paper/2605.19481"},"observation_digest":"sha256:33b3788f763bfabb66af0bbd25ab4f8235e9be377739292e2203ef56dbb30ee7","observation_id":"e1fee6fa-ea3e-4732-be29-a30610397563","resolution":{"observed_at":"2026-05-20T02:12:58.978895Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"deacd3fb-c893-4639-885a-9c53941f5e5f","year":2025},"citing_paper":{"arxiv_id":"2605.19481","last_updated":"2026-05-19T07:34:08Z","snapshot_observed_at":"2026-08-02T01:46:46.351301Z","submitted_at":"2026-05-19T07:34:08Z","title":"C2CServe: Leveraging NVLink-C2C for Elastic Serverless LLM Serving on MIG","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-05-20T02:10:57.582345Z"},"links":{"citing_paper":"/paper/2605.19481"},"observation_digest":"sha256:52e7f4e90a70c3e931d4e21a5cbd07bb1605c224168188551e6e676f3036659e","observation_id":"30855cbb-4f89-4bc9-8428-47d10ab9d20a","resolution":{"observed_at":"2026-05-20T02:12:58.999300Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"9f13fbfe-8a25-4e61-a0ac-961a50d58387","year":2025},"citing_paper":{"arxiv_id":"2605.19481","last_updated":"2026-05-19T07:34:08Z","snapshot_observed_at":"2026-08-02T01:46:46.351301Z","submitted_at":"2026-05-19T07:34:08Z","title":"C2CServe: Leveraging NVLink-C2C for Elastic Serverless LLM Serving on MIG","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-05-20T02:10:57.582345Z"},"links":{"citing_paper":"/paper/2605.19481"},"observation_digest":"sha256:f0a1c4fec0378ffa611b63025cb8b4492028211f401d4d62df3dfd13303011f3","observation_id":"0d5712f5-2faf-414a-abb6-aa9bc7fee4bb","resolution":{"observed_at":"2026-05-20T02:12:58.982567Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"21bc19bd-a08f-405c-baaa-bfd6064aec8d","year":2025},"citing_paper":{"arxiv_id":"2605.19481","last_updated":"2026-05-19T07:34:08Z","snapshot_observed_at":"2026-08-02T01:46:46.351301Z","submitted_at":"2026-05-19T07:34:08Z","title":"C2CServe: Leveraging NVLink-C2C for Elastic Serverless LLM Serving on MIG","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-05-20T02:10:57.582345Z"},"links":{"citing_paper":"/paper/2605.19481"},"observation_digest":"sha256:4eac524ae176034748aa5c8977534f16c4fb904643df10ef36c314b1e5821b3b","observation_id":"2f4b166c-b7e6-4551-a6ea-45abec11b94a","resolution":{"observed_at":"2026-05-20T02:12:58.987071Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"dc917a80-1257-4350-917a-8f3089785291","year":2025},"citing_paper":{"arxiv_id":"2605.19481","last_updated":"2026-05-19T07:34:08Z","snapshot_observed_at":"2026-08-02T01:46:46.351301Z","submitted_at":"2026-05-19T07:34:08Z","title":"C2CServe: Leveraging NVLink-C2C for Elastic Serverless LLM Serving on MIG","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-05-20T02:10:57.582345Z"},"links":{"citing_paper":"/paper/2605.19481"},"observation_digest":"sha256:26862fb0c4e632cf63f55861b599008e8bf0067925d4a93aa0dffdcac3473c20","observation_id":"689c1e4d-6ea5-43c3-8eb8-e7b8d364f974","resolution":{"observed_at":"2026-05-20T02:12:58.964412Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"com/cublas","venue":null,"work_id":"3436ce4a-810e-4d0f-896c-1f5f1b848f24","year":2026},"citing_paper":{"arxiv_id":"2605.19481","last_updated":"2026-05-19T07:34:08Z","snapshot_observed_at":"2026-08-02T01:46:46.351301Z","submitted_at":"2026-05-19T07:34:08Z","title":"C2CServe: Leveraging NVLink-C2C for Elastic Serverless LLM Serving on MIG","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-05-20T02:10:57.582345Z"},"links":{"citing_paper":"/paper/2605.19481"},"observation_digest":"sha256:3f2e57227281cd3a4f168084cce3ebd9569ead55b3263c39fe27bfe35636f0b9","observation_id":"ddab7fa1-46ba-43b8-883b-869b079a8d9e","resolution":{"observed_at":"2026-05-20T02:12:58.944998Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"870aad57-a888-4592-9c9a-11bbdf3d2d73","year":2026},"citing_paper":{"arxiv_id":"2605.19481","last_updated":"2026-05-19T07:34:08Z","snapshot_observed_at":"2026-08-02T01:46:46.351301Z","submitted_at":"2026-05-19T07:34:08Z","title":"C2CServe: Leveraging NVLink-C2C for Elastic Serverless LLM Serving on MIG","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-05-20T02:10:57.582345Z"},"links":{"citing_paper":"/paper/2605.19481"},"observation_digest":"sha256:529a623ec24f9525db5de977ddc5e17b7baf815ce7b099675226e6821cd8f0a2","observation_id":"11d77897-3c98-4fbe-9409-5bb47335d683","resolution":{"observed_at":"2026-05-20T02:12:58.941423Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2303.08774","last_updated":"2024-03-04T06:01:33Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-03-15T17:15:04Z","title":"GPT-4 Technical Report","version":6},"cited_work":{"arxiv_id":"2303.08774","doi":"10.1002/tea.20265","metadata_source":"pith","pith_arxiv_id":"2303.08774","snapshot_observed_at":"2026-07-11T11:50:26.030339Z","title":"GPT-4 Technical Report","venue":"cs.CL","work_id":"b928e041-6991-4c08-8c81-0359e4097c7b","year":2023},"citing_paper":{"arxiv_id":"2605.19481","last_updated":"2026-05-19T07:34:08Z","snapshot_observed_at":"2026-08-02T01:46:46.351301Z","submitted_at":"2026-05-19T07:34:08Z","title":"C2CServe: Leveraging NVLink-C2C for Elastic Serverless LLM Serving on MIG","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-05-20T02:10:57.582345Z"},"links":{"cited_paper":"/paper/2303.08774","citing_paper":"/paper/2605.19481"},"observation_digest":"sha256:781f72839bbd5172c1f2ce7e432e0e4ed3c921f5a6bd5f6e6e306c017c135812","observation_id":"9e80f8cf-27d3-4f73-b116-1a27ec1f7310","resolution":{"observed_at":"2026-05-20T02:12:58.297659Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Taming {Throughput-Latency} tradeoff in {LLM} inference with {Sarathi-Serve}","venue":null,"work_id":"0a8227cf-1856-4a39-b25c-bffa3a4ae643","year":2024},"citing_paper":{"arxiv_id":"2605.19481","last_updated":"2026-05-19T07:34:08Z","snapshot_observed_at":"2026-08-02T01:46:46.351301Z","submitted_at":"2026-05-19T07:34:08Z","title":"C2CServe: Leveraging NVLink-C2C for Elastic Serverless LLM Serving on MIG","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-05-20T02:10:57.582345Z"},"links":{"citing_paper":"/paper/2605.19481"},"observation_digest":"sha256:e163ffbbbed1ed35596aecf7c5222335dca6eb3bc24b73c26419dbb453439757","observation_id":"1c0d5284-8341-49b3-8e2b-ad1765209da4","resolution":{"observed_at":"2026-05-20T02:12:59.014792Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Serving heterogeneous machine learning models on {Multi-GPU } servers with {Spatio-Temporal} sharing","venue":null,"work_id":"f414e4e0-a915-4a2e-a004-dbc88f9fe8f4","year":2022},"citing_paper":{"arxiv_id":"2605.19481","last_updated":"2026-05-19T07:34:08Z","snapshot_observed_at":"2026-08-02T01:46:46.351301Z","submitted_at":"2026-05-19T07:34:08Z","title":"C2CServe: Leveraging NVLink-C2C for Elastic Serverless LLM Serving on MIG","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-05-20T02:10:57.582345Z"},"links":{"citing_paper":"/paper/2605.19481"},"observation_digest":"sha256:7fec2d694550925823fcbfa1a02ebf1b5156833c40b88b153da2ed7d53ef9cc7","observation_id":"362563db-b81a-4b88-a5ea-b7a0dc81d108","resolution":{"observed_at":"2026-05-20T02:12:59.018991Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Muxserve: flexible spatial-temporal multiplexing for multiple llm serving","venue":null,"work_id":"f8142707-b015-46fd-b39a-5df601c7a68f","year":2024},"citing_paper":{"arxiv_id":"2605.19481","last_updated":"2026-05-19T07:34:08Z","snapshot_observed_at":"2026-08-02T01:46:46.351301Z","submitted_at":"2026-05-19T07:34:08Z","title":"C2CServe: Leveraging NVLink-C2C for Elastic Serverless LLM Serving on MIG","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-05-20T02:10:57.582345Z"},"links":{"citing_paper":"/paper/2605.19481"},"observation_digest":"sha256:dee71516e9a508524c97bc3a5c9fe36a9f382f8b0d04d7225b3ca6b5dd02ccb5","observation_id":"82f6895f-8e66-4ed9-9b74-94ce3561c402","resolution":{"observed_at":"2026-05-20T02:12:59.022898Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"The llama 3 herd of models.arXiv e-prints","venue":null,"work_id":"cd9fee59-9ed8-4913-9203-9e83f63bd920","year":2024},"citing_paper":{"arxiv_id":"2605.19481","last_updated":"2026-05-19T07:34:08Z","snapshot_observed_at":"2026-08-02T01:46:46.351301Z","submitted_at":"2026-05-19T07:34:08Z","title":"C2CServe: Leveraging NVLink-C2C for Elastic Serverless LLM Serving on MIG","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-05-20T02:10:57.582345Z"},"links":{"citing_paper":"/paper/2605.19481"},"observation_digest":"sha256:ff856e1ae2103e6a218d1a4a8e8a602ff8b8b4cf532a73efa39de043c437d02b","observation_id":"36b0ef04-771a-4cd9-b6d7-53133977825a","resolution":{"observed_at":"2026-05-20T02:12:59.077922Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"InProceedings of OSDI","venue":null,"work_id":"f0a2b6c4-47e8-48a1-b178-50dc8d13ed76","year":2024},"citing_paper":{"arxiv_id":"2605.19481","last_updated":"2026-05-19T07:34:08Z","snapshot_observed_at":"2026-08-02T01:46:46.351301Z","submitted_at":"2026-05-19T07:34:08Z","title":"C2CServe: Leveraging NVLink-C2C for Elastic Serverless LLM Serving on MIG","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-05-20T02:10:57.582345Z"},"links":{"citing_paper":"/paper/2605.19481"},"observation_digest":"sha256:a648bd5e7433197c7851c26a67e4e62e2a0afaa82a3ec3483e974bfa583532d9","observation_id":"1a59ca42-1517-471e-8642-71f4517f83f3","resolution":{"observed_at":"2026-05-20T02:12:59.003304Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"53c7144e-d832-4aea-aba3-b46b5972c50b","year":2022},"citing_paper":{"arxiv_id":"2605.19481","last_updated":"2026-05-19T07:34:08Z","snapshot_observed_at":"2026-08-02T01:46:46.351301Z","submitted_at":"2026-05-19T07:34:08Z","title":"C2CServe: Leveraging NVLink-C2C for Elastic Serverless LLM Serving on MIG","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-05-20T02:10:57.582345Z"},"links":{"citing_paper":"/paper/2605.19481"},"observation_digest":"sha256:1df1921aa40730df232eb59d7a82dd20844c3f6b51d757ef0e4bff3556fe940c","observation_id":"8c155c6c-916d-4c5a-8dba-2b48d8daeca8","resolution":{"observed_at":"2026-05-20T02:12:59.006857Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2404.14527","last_updated":"2024-07-22T10:56:19Z","snapshot_observed_at":"2026-07-06T18:04:00.217896Z","submitted_at":"2024-04-22T18:56:18Z","title":"M\\'elange: Cost Efficient Large Language Model Serving by Exploiting GPU Heterogeneity","version":4},"cited_work":{"arxiv_id":"2404.14527","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2404.14527","snapshot_observed_at":"2026-07-03T12:48:11.776825Z","title":"M \\’elange: Cost efficient large language model serving by exploiting gpu heterogeneity","venue":null,"work_id":"bc062082-bb69-4c79-9bda-625aac03fada","year":2024},"citing_paper":{"arxiv_id":"2605.19481","last_updated":"2026-05-19T07:34:08Z","snapshot_observed_at":"2026-08-02T01:46:46.351301Z","submitted_at":"2026-05-19T07:34:08Z","title":"C2CServe: Leveraging NVLink-C2C for Elastic Serverless LLM Serving on MIG","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-05-20T02:10:57.582345Z"},"links":{"cited_paper":"/paper/2404.14527","citing_paper":"/paper/2605.19481"},"observation_digest":"sha256:6465ca266e0baef2a862021ac1ecf86c9590a589c95e14deaa5a789b6e92d77e","observation_id":"5ad6e453-c711-4d7e-ba5f-f05139539f09","resolution":{"observed_at":"2026-05-20T02:12:58.302834Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Microsecond-scale preemption for concurrent {GPU- accelerated} {DNN}inferences","venue":null,"work_id":"342c21a2-2602-47bb-a56f-ff2cd99aab67","year":2022},"citing_paper":{"arxiv_id":"2605.19481","last_updated":"2026-05-19T07:34:08Z","snapshot_observed_at":"2026-08-02T01:46:46.351301Z","submitted_at":"2026-05-19T07:34:08Z","title":"C2CServe: Leveraging NVLink-C2C for Elastic Serverless LLM Serving on MIG","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-05-20T02:10:57.582345Z"},"links":{"citing_paper":"/paper/2605.19481"},"observation_digest":"sha256:761ac2b26907387c7de93aac47cfa17e2d3d4e75bf9fb2d432b702f1f72d6e59","observation_id":"b94f6beb-1af9-4535-a223-8a5297684439","resolution":{"observed_at":"2026-05-20T02:12:58.960867Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Resource multiplexing in tuning and serving large language models","venue":null,"work_id":"42ea69eb-51b8-4aaf-bd48-9291ff07aef4","year":2025},"citing_paper":{"arxiv_id":"2605.19481","last_updated":"2026-05-19T07:34:08Z","snapshot_observed_at":"2026-08-02T01:46:46.351301Z","submitted_at":"2026-05-19T07:34:08Z","title":"C2CServe: Leveraging NVLink-C2C for Elastic Serverless LLM Serving on MIG","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-05-20T02:10:57.582345Z"},"links":{"citing_paper":"/paper/2605.19481"},"observation_digest":"sha256:126339eae1a811139b1264ed6ae88982f82e893a6dcbe6d574ce336135ab16dc","observation_id":"477a9d29-f1cc-40ef-98a5-0839d0890610","resolution":{"observed_at":"2026-05-20T02:12:58.968239Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"{DEEPSERVE}: Serverless large language model serving at scale","venue":null,"work_id":"2e77fc51-40b7-4da1-866c-faab88c884fd","year":2025},"citing_paper":{"arxiv_id":"2605.19481","last_updated":"2026-05-19T07:34:08Z","snapshot_observed_at":"2026-08-02T01:46:46.351301Z","submitted_at":"2026-05-19T07:34:08Z","title":"C2CServe: Leveraging NVLink-C2C for Elastic Serverless LLM Serving on MIG","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-05-20T02:10:57.582345Z"},"links":{"citing_paper":"/paper/2605.19481"},"observation_digest":"sha256:dae49a6f5de7d4ccc64e7664e062a3da721c8c698d4e0ad683ba62b6fca7de38","observation_id":"a2e6cda0-470e-42da-828e-308aabb2e00c","resolution":{"observed_at":"2026-05-20T02:12:58.994936Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2401.04088","last_updated":"2024-01-08T18:47:34Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-01-08T18:47:34Z","title":"Mixtral of Experts","version":1},"cited_work":{"arxiv_id":"2401.04088","doi":"10.48550/arxiv.2401.04088","metadata_source":"pith","pith_arxiv_id":"2401.04088","snapshot_observed_at":"2026-07-11T03:27:47.028679Z","title":"Mixtral of Experts","venue":"cs.LG","work_id":"0de8c352-9daa-4e1e-8c7b-3d0dec69f369","year":2024},"citing_paper":{"arxiv_id":"2605.19481","last_updated":"2026-05-19T07:34:08Z","snapshot_observed_at":"2026-08-02T01:46:46.351301Z","submitted_at":"2026-05-19T07:34:08Z","title":"C2CServe: Leveraging NVLink-C2C for Elastic Serverless LLM Serving on MIG","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-05-20T02:10:57.582345Z"},"links":{"cited_paper":"/paper/2401.04088","citing_paper":"/paper/2605.19481"},"observation_digest":"sha256:7a662a011514b0812c067a041040a9def66ff109a1d2fa01eac7d8e73aec1103","observation_id":"abb78f65-7401-4db6-a437-aa0f679a7dd2","resolution":{"observed_at":"2026-05-20T02:12:58.307856Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-07-09T08:48:39.110013+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-09T08:48:39.110013+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2001.08361","last_updated":"2020-01-23T03:59:20Z","snapshot_observed_at":"2026-07-06T08:52:12.656082Z","submitted_at":"2020-01-23T03:59:20Z","title":"Scaling Laws for Neural Language Models","version":1},"cited_work":{"arxiv_id":"2001.08361","doi":"10.1145/3616855.3635845","metadata_source":"pith","pith_arxiv_id":"2001.08361","snapshot_observed_at":"2026-07-11T11:50:26.030339Z","title":"Scaling Laws for Neural Language Models","venue":"cs.LG","work_id":"b7dd8749-9c45-4977-ab9b-64478dce1ae8","year":2020},"citing_paper":{"arxiv_id":"2605.19481","last_updated":"2026-05-19T07:34:08Z","snapshot_observed_at":"2026-08-02T01:46:46.351301Z","submitted_at":"2026-05-19T07:34:08Z","title":"C2CServe: Leveraging NVLink-C2C for Elastic Serverless LLM Serving on MIG","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-05-20T02:10:57.582345Z"},"links":{"cited_paper":"/paper/2001.08361","citing_paper":"/paper/2605.19481"},"observation_digest":"sha256:e49c0fe0766c3e3c4ae3c8351339c4bc8d77c603c2001cba71a6a8787962db61","observation_id":"d6343828-bed5-4e23-9254-e2c40da025e3","resolution":{"observed_at":"2026-05-20T02:12:58.343677Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Tetris: Memory-efficient serverless inference through tensor sharing","venue":null,"work_id":"d2020f99-7495-49ac-ab56-74aa92547530","year":2022},"citing_paper":{"arxiv_id":"2605.19481","last_updated":"2026-05-19T07:34:08Z","snapshot_observed_at":"2026-08-02T01:46:46.351301Z","submitted_at":"2026-05-19T07:34:08Z","title":"C2CServe: Leveraging NVLink-C2C for Elastic Serverless LLM Serving on MIG","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-05-20T02:10:57.582345Z"},"links":{"citing_paper":"/paper/2605.19481"},"observation_digest":"sha256:abd31bec5274f464f7b8f34fe0bcacfc18c8492a0e12e149da3c3966df6806e9","observation_id":"9fcd2bd3-e1ab-48ea-b727-487eda578b4b","resolution":{"observed_at":"2026-05-20T02:12:58.990854Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Oneiros: Kv cache opti- mization through parameter remapping for multi-tenant llm serving","venue":null,"work_id":"f4ebb5dd-88bd-467d-98af-1512b8556e89","year":2025},"citing_paper":{"arxiv_id":"2605.19481","last_updated":"2026-05-19T07:34:08Z","snapshot_observed_at":"2026-08-02T01:46:46.351301Z","submitted_at":"2026-05-19T07:34:08Z","title":"C2CServe: Leveraging NVLink-C2C for Elastic Serverless LLM Serving on MIG","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-05-20T02:10:57.582345Z"},"links":{"citing_paper":"/paper/2605.19481"},"observation_digest":"sha256:d2aa42d8944b10d15064e713df416f3eb04d61f21e59498e83658a298ed8f32a","observation_id":"3b0842c5-01e0-4d7f-b0f7-84dc77388576","resolution":{"observed_at":"2026-05-20T02:12:59.070668Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Superoffload: Unleashing the power of large-scale llm training on superchips","venue":null,"work_id":"8176f659-44d7-4242-90ca-44f3b0784603","year":2026},"citing_paper":{"arxiv_id":"2605.19481","last_updated":"2026-05-19T07:34:08Z","snapshot_observed_at":"2026-08-02T01:46:46.351301Z","submitted_at":"2026-05-19T07:34:08Z","title":"C2CServe: Leveraging NVLink-C2C for Elastic Serverless LLM Serving on MIG","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-05-20T02:10:57.582345Z"},"links":{"citing_paper":"/paper/2605.19481"},"observation_digest":"sha256:8b9e9f223565d08c5c2a807d449f1142280e988226e50fbac0aa8a6bea23d816","observation_id":"d1a75fb5-64f3-47f4-9078-10636f649472","resolution":{"observed_at":"2026-05-20T02:12:59.081732Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Flexpipe: Adapting dynamic llm serving through inflight pipeline refactoring in fragmented serverless clusters","venue":null,"work_id":"9222af20-7b9e-4b07-8053-5ada0adce396","year":2026},"citing_paper":{"arxiv_id":"2605.19481","last_updated":"2026-05-19T07:34:08Z","snapshot_observed_at":"2026-08-02T01:46:46.351301Z","submitted_at":"2026-05-19T07:34:08Z","title":"C2CServe: Leveraging NVLink-C2C for Elastic Serverless LLM Serving on MIG","version":1},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-05-20T02:10:57.582345Z"},"links":{"citing_paper":"/paper/2605.19481"},"observation_digest":"sha256:c90e3af95cdc92cb4444d64e68dd8d23278655aa0b24d66724c1c3d43e0c953d","observation_id":"f716d51e-0d4d-4a59-aceb-d058644301cf","resolution":{"observed_at":"2026-05-20T02:12:59.085653Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Under- standing diffusion model serving in production: A top-down analysis of workload, scheduling, and resource efficiency","venue":null,"work_id":"3a0f59fa-f763-43d0-bd1c-5c02b44f9efe","year":2025},"citing_paper":{"arxiv_id":"2605.19481","last_updated":"2026-05-19T07:34:08Z","snapshot_observed_at":"2026-08-02T01:46:46.351301Z","submitted_at":"2026-05-19T07:34:08Z","title":"C2CServe: Leveraging NVLink-C2C for Elastic Serverless LLM Serving on MIG","version":1},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-05-20T02:10:57.582345Z"},"links":{"citing_paper":"/paper/2605.19481"},"observation_digest":"sha256:0bec767db3e4c7ce1ffccca7c4b105210d72bdee41a733c9282e099e0de67c0b","observation_id":"1a6ad5ee-d283-4490-9212-5b8d765e1668","resolution":{"observed_at":"2026-05-20T02:12:59.058546Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2604.06664","last_updated":"2026-04-08T04:31:34Z","snapshot_observed_at":"2026-08-03T07:58:59.160028Z","submitted_at":"2026-04-08T04:31:34Z","title":"Foundry: Template-Based CUDA Graph Context Materialization for Fast LLM Serving Cold Start","version":1},"cited_work":{"arxiv_id":"2604.06664","doi":null,"metadata_source":"pith","pith_arxiv_id":"2604.06664","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Foundry: Template-Based CUDA Graph Context Materialization for Fast LLM Serving Cold Start","venue":"cs.DC","work_id":"f6681bbf-6189-4fa9-ad45-b91094f65fa3","year":2026},"citing_paper":{"arxiv_id":"2605.19481","last_updated":"2026-05-19T07:34:08Z","snapshot_observed_at":"2026-08-02T01:46:46.351301Z","submitted_at":"2026-05-19T07:34:08Z","title":"C2CServe: Leveraging NVLink-C2C for Elastic Serverless LLM Serving on MIG","version":1},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-05-20T02:10:57.582345Z"},"links":{"cited_paper":"/paper/2604.06664","citing_paper":"/paper/2605.19481"},"observation_digest":"sha256:378fd87a65f92c9f25a8f783c3e020b7c0ac205578dcbadf2fbbe5eff7c1357c","observation_id":"d036bfed-c54f-4977-9299-1608fca2e7af","resolution":{"observed_at":"2026-05-20T02:12:58.337827Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Sky- serve: Serving ai models across regions and clouds with spot instances","venue":null,"work_id":"5d8c4cdd-08ee-4d7c-b3d7-5d86b03e58ae","year":2025},"citing_paper":{"arxiv_id":"2605.19481","last_updated":"2026-05-19T07:34:08Z","snapshot_observed_at":"2026-08-02T01:46:46.351301Z","submitted_at":"2026-05-19T07:34:08Z","title":"C2CServe: Leveraging NVLink-C2C for Elastic Serverless LLM Serving on MIG","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-05-20T02:10:57.582345Z"},"links":{"citing_paper":"/paper/2605.19481"},"observation_digest":"sha256:2a4d8416aad31eb7d216d113c0a5bfad93f8aa09475ca1bb291e63cffc4f8505","observation_id":"ecaa3fae-1868-4623-b0db-30116de0791d","resolution":{"observed_at":"2026-05-20T02:12:59.062459Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"S-lora: Serving thousands of concurrent lora adapters","venue":null,"work_id":"648c368b-f527-45f6-ad66-53417f25cdd3","year":2023},"citing_paper":{"arxiv_id":"2605.19481","last_updated":"2026-05-19T07:34:08Z","snapshot_observed_at":"2026-08-02T01:46:46.351301Z","submitted_at":"2026-05-19T07:34:08Z","title":"C2CServe: Leveraging NVLink-C2C for Elastic Serverless LLM Serving on MIG","version":1},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-05-20T02:10:57.582345Z"},"links":{"citing_paper":"/paper/2605.19481"},"observation_digest":"sha256:febca736c31368be08e77c4bff435d689ba07d52135b41505f2feed596bdacd6","observation_id":"cb9e339e-1dc0-4759-af09-7e2b1f5b92da","resolution":{"observed_at":"2026-05-20T02:12:59.051045Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Orion: Interference- aware, fine-grained gpu sharing for ml applications","venue":null,"work_id":"587b229c-6dec-43a5-9581-447434e817be","year":2024},"citing_paper":{"arxiv_id":"2605.19481","last_updated":"2026-05-19T07:34:08Z","snapshot_observed_at":"2026-08-02T01:46:46.351301Z","submitted_at":"2026-05-19T07:34:08Z","title":"C2CServe: Leveraging NVLink-C2C for Elastic Serverless LLM Serving on MIG","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-05-20T02:10:57.582345Z"},"links":{"citing_paper":"/paper/2605.19481"},"observation_digest":"sha256:7a3dc553c3e7509ef40d099e829b979c0eefa4d329acafdeac87928871327b2b","observation_id":"0b01c14e-3f75-44b6-9342-8b793f143d5e","resolution":{"observed_at":"2026-05-20T02:12:59.054741Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2302.13971","last_updated":"2023-02-27T17:11:15Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-02-27T17:11:15Z","title":"LLaMA: Open and Efficient Foundation Language Models","version":1},"cited_work":{"arxiv_id":"2302.13971","doi":"10.48550/arxiv.2302.13971","metadata_source":"pith","pith_arxiv_id":"2302.13971","snapshot_observed_at":"2026-07-11T03:47:48.508181Z","title":"LLaMA: Open and Efficient Foundation Language Models","venue":"cs.CL","work_id":"c018fc23-6f3f-4035-9d02-28a2173b2b9d","year":2023},"citing_paper":{"arxiv_id":"2605.19481","last_updated":"2026-05-19T07:34:08Z","snapshot_observed_at":"2026-08-02T01:46:46.351301Z","submitted_at":"2026-05-19T07:34:08Z","title":"C2CServe: Leveraging NVLink-C2C for Elastic Serverless LLM Serving on MIG","version":1},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-05-20T02:10:57.582345Z"},"links":{"cited_paper":"/paper/2302.13971","citing_paper":"/paper/2605.19481"},"observation_digest":"sha256:9da396b3a7542e2816e16e64046563ba3766e492488d4fa02c9fd9d682ae4a82","observation_id":"920c90b3-2f35-49e0-95f9-f4dabf415dc1","resolution":{"observed_at":"2026-05-20T02:12:58.332068Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-01T11:08:05.851253+00:00","source":"crossref_status_cache"},{"observed_at":"2026-08-01T11:08:05.851253+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Attention is all you need","venue":null,"work_id":"e71454ba-f7cf-47ec-a730-67eb2eabe465","year":2017},"citing_paper":{"arxiv_id":"2605.19481","last_updated":"2026-05-19T07:34:08Z","snapshot_observed_at":"2026-08-02T01:46:46.351301Z","submitted_at":"2026-05-19T07:34:08Z","title":"C2CServe: Leveraging NVLink-C2C for Elastic Serverless LLM Serving on MIG","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-05-20T02:10:57.582345Z"},"links":{"citing_paper":"/paper/2605.19481"},"observation_digest":"sha256:fdb82f86628cb20f11571f7b40c2b9ec81c2c072b5a00b8e122c3d7338db072a","observation_id":"eb283536-f141-47d6-a77e-ef566bfacf72","resolution":{"observed_at":"2026-05-20T02:12:59.066949Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Zorua: A holistic approach to resource virtualization in gpus","venue":null,"work_id":"2143386b-e61a-4a2c-9249-64c1638e7b21","year":2016},"citing_paper":{"arxiv_id":"2605.19481","last_updated":"2026-05-19T07:34:08Z","snapshot_observed_at":"2026-08-02T01:46:46.351301Z","submitted_at":"2026-05-19T07:34:08Z","title":"C2CServe: Leveraging NVLink-C2C for Elastic Serverless LLM Serving on MIG","version":1},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-05-20T02:10:57.582345Z"},"links":{"citing_paper":"/paper/2605.19481"},"observation_digest":"sha256:77f46f3593ffb25e589082a354c176d6e92d655ff1ffb26b20e6c9567c51fa7e","observation_id":"eab6ec7e-a728-4912-8bb6-4d7be7423d71","resolution":{"observed_at":"2026-05-20T02:12:59.089649Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"{ByteCheckpoint}: A unified checkpointing system for large foundation model development","venue":null,"work_id":"08a312d5-eb9d-407e-b7f3-872bc5ca124b","year":2025},"citing_paper":{"arxiv_id":"2605.19481","last_updated":"2026-05-19T07:34:08Z","snapshot_observed_at":"2026-08-02T01:46:46.351301Z","submitted_at":"2026-05-19T07:34:08Z","title":"C2CServe: Leveraging NVLink-C2C for Elastic Serverless LLM Serving on MIG","version":1},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-05-20T02:10:57.582345Z"},"links":{"citing_paper":"/paper/2605.19481"},"observation_digest":"sha256:098978db4a2465649446ffd40257cf89edfd955dc264ce3788091d0b0dc4fe75","observation_id":"d2737fbc-1204-4385-8390-0f7c9df4914c","resolution":{"observed_at":"2026-05-20T02:12:59.043423Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Aegaeon: Effective gpu pooling for concurrent llm serving on the market","venue":null,"work_id":"39a5eca2-ef02-44ee-9860-5ed4a361194a","year":2025},"citing_paper":{"arxiv_id":"2605.19481","last_updated":"2026-05-19T07:34:08Z","snapshot_observed_at":"2026-08-02T01:46:46.351301Z","submitted_at":"2026-05-19T07:34:08Z","title":"C2CServe: Leveraging NVLink-C2C for Elastic Serverless LLM Serving on MIG","version":1},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-05-20T02:10:57.582345Z"},"links":{"citing_paper":"/paper/2605.19481"},"observation_digest":"sha256:0e57fa11f72fa186bed2ceb0ef9b4eb2f714a532c80366971ac97605bb230ad6","observation_id":"5dbbb3ba-1bf1-4174-85bd-9f47babcb9d7","resolution":{"observed_at":"2026-05-20T02:12:59.046880Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2411.09317","last_updated":"2024-11-14T09:50:41Z","snapshot_observed_at":"2026-07-06T19:50:15.546603Z","submitted_at":"2024-11-14T09:50:41Z","title":"Pie: Pooling CPU Memory for LLM Inference","version":1},"cited_work":{"arxiv_id":"2411.09317","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2411.09317","snapshot_observed_at":"2026-07-01T21:06:13.603806Z","title":"Pie: Pooling cpu memory for llm inference.arXiv preprint arXiv:2411.09317","venue":null,"work_id":"c65f2db1-e070-47b4-8bb0-5206762f13a5","year":2024},"citing_paper":{"arxiv_id":"2605.19481","last_updated":"2026-05-19T07:34:08Z","snapshot_observed_at":"2026-08-02T01:46:46.351301Z","submitted_at":"2026-05-19T07:34:08Z","title":"C2CServe: Leveraging NVLink-C2C for Elastic Serverless LLM Serving on MIG","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-05-20T02:10:57.582345Z"},"links":{"cited_paper":"/paper/2411.09317","citing_paper":"/paper/2605.19481"},"observation_digest":"sha256:0900a231ef45bb20ffa5cf2e34f0f7b8866ff65b640830d3e1a0befa53f00418","observation_id":"8d9f84d8-5603-4c8e-954b-1abe07e21e15","resolution":{"observed_at":"2026-05-20T02:12:58.319737Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2401.14361","last_updated":"2025-03-12T18:14:21Z","snapshot_observed_at":"2026-08-03T09:09:47.667294Z","submitted_at":"2024-01-25T18:07:50Z","title":"MoE-Infinity: Efficient MoE Inference on Personal Machines with Sparsity-Aware Expert Cache","version":3},"cited_work":{"arxiv_id":"2401.14361","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2401.14361","snapshot_observed_at":"2026-07-04T07:59:39.550340Z","title":"Xue, L., Fu, Y ., Lu, Z., Mai, L., and Marina, M","venue":null,"work_id":"2fcfe188-dd65-40ca-b9f1-f10a58f699ae","year":2025},"citing_paper":{"arxiv_id":"2605.19481","last_updated":"2026-05-19T07:34:08Z","snapshot_observed_at":"2026-08-02T01:46:46.351301Z","submitted_at":"2026-05-19T07:34:08Z","title":"C2CServe: Leveraging NVLink-C2C for Elastic Serverless LLM Serving on MIG","version":1},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-05-20T02:10:57.582345Z"},"links":{"cited_paper":"/paper/2401.14361","citing_paper":"/paper/2605.19481"},"observation_digest":"sha256:649519a0757675ec0e82a72269730c3aa8897a7ea1f7ab788508a0112b8186a2","observation_id":"81743ddd-ab30-4031-9717-3c676997fcbc","resolution":{"observed_at":"2026-05-20T02:12:58.326430Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.09388","last_updated":"2025-05-14T13:41:34Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-05-14T13:41:34Z","title":"Qwen3 Technical Report","version":1},"cited_work":{"arxiv_id":"2505.09388","doi":"10.1016/j.aiopen.2022.12","metadata_source":"pith","pith_arxiv_id":"2505.09388","snapshot_observed_at":"2026-07-11T11:50:26.030339Z","title":"Qwen3 Technical Report","venue":"cs.CL","work_id":"25a4e30c-1232-48e7-9925-02fa12ba7c9e","year":2025},"citing_paper":{"arxiv_id":"2605.19481","last_updated":"2026-05-19T07:34:08Z","snapshot_observed_at":"2026-08-02T01:46:46.351301Z","submitted_at":"2026-05-19T07:34:08Z","title":"C2CServe: Leveraging NVLink-C2C for Elastic Serverless LLM Serving on MIG","version":1},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-05-20T02:10:57.582345Z"},"links":{"cited_paper":"/paper/2505.09388","citing_paper":"/paper/2605.19481"},"observation_digest":"sha256:e0de3afebcbf7398be77d436b0ad63d39f079e2c32699da7f546cf1eb71df6ce","observation_id":"3eb6b05e-238f-433c-a3f4-92da10f9495a","resolution":{"observed_at":"2026-05-20T02:12:58.313173Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Taming latency- memory trade-off in moe-based llm serving via fine-grained expert offloading","venue":null,"work_id":"e62389c9-86e0-44d9-a74f-51c6fc6bb07f","year":2026},"citing_paper":{"arxiv_id":"2605.19481","last_updated":"2026-05-19T07:34:08Z","snapshot_observed_at":"2026-08-02T01:46:46.351301Z","submitted_at":"2026-05-19T07:34:08Z","title":"C2CServe: Leveraging NVLink-C2C for Elastic Serverless LLM Serving on MIG","version":1},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-05-20T02:10:57.582345Z"},"links":{"citing_paper":"/paper/2605.19481"},"observation_digest":"sha256:c4202b5943715bbb6ece2c3eeea706bf1bd4443aaaa883b015374f3d343a34b0","observation_id":"0aff380f-fdc2-4738-93c8-879d679b2078","resolution":{"observed_at":"2026-05-20T02:12:59.035065Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Superinfer: Slo- aware rotary scheduling and memory management for llm inference on superchips","venue":null,"work_id":"eb1a4726-bbeb-407f-8fe3-64cff19b933f","year":2026},"citing_paper":{"arxiv_id":"2605.19481","last_updated":"2026-05-19T07:34:08Z","snapshot_observed_at":"2026-08-02T01:46:46.351301Z","submitted_at":"2026-05-19T07:34:08Z","title":"C2CServe: Leveraging NVLink-C2C for Elastic Serverless LLM Serving on MIG","version":1},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-05-20T02:10:57.582345Z"},"links":{"citing_paper":"/paper/2605.19481"},"observation_digest":"sha256:b0b965549848a9dd3bd9dcc54f265122e244ea692bfc081637a8dd2a374e5fa2","observation_id":"11be7229-e4e3-41bc-afad-b105efcd4a12","resolution":{"observed_at":"2026-05-20T02:12:59.039513Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Medusa: Accelerating serverless llm inference with materialization","venue":null,"work_id":"274fa6a8-504d-4695-9fe2-b6de6392ed6f","year":2025},"citing_paper":{"arxiv_id":"2605.19481","last_updated":"2026-05-19T07:34:08Z","snapshot_observed_at":"2026-08-02T01:46:46.351301Z","submitted_at":"2026-05-19T07:34:08Z","title":"C2CServe: Leveraging NVLink-C2C for Elastic Serverless LLM Serving on MIG","version":1},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-05-20T02:10:57.582345Z"},"links":{"citing_paper":"/paper/2605.19481"},"observation_digest":"sha256:90abc88a014139843183653f6f6b5aa2404850d282d84d43525b42ba926f1002","observation_id":"538bb82d-8842-4946-b2da-75eefdf75bcc","resolution":{"observed_at":"2026-05-20T02:12:59.026972Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-04T06:34:03.388597+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2605.19481","last_updated":"2026-05-19T07:34:08Z","latest_version":1,"primary_category":"cs.OS","snapshot_observed_at":"2026-08-02T01:46:46.351301Z","submitted_at":"2026-05-19T07:34:08Z","title":"C2CServe: Leveraging NVLink-C2C for Elastic Serverless LLM Serving on MIG"},"reference_resolution":{"displayed":48,"state_counts":{"malformed_identifier":0,"metadata_mismatch":1,"parse_uncertain":0,"unresolved":15,"verified_exact":8,"verified_fuzzy":24},"total_outbound_references":48},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-04T06:34:03.388597+00:00","source":"crossref"},{"observed_at":"2026-08-04T06:33:57.428241+00:00","source":"retraction_watch"}],"thesis":"As of 4 August 2026, this Paper Citation Record lists 48 of 48 outbound references and 0 inbound Pith citation observations for arXiv:2605.19481."}