{"as_of":"2026-08-07T23:28:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:8bf995d980ac3c43350b4f7e8f2027682d9ae54571ac4f0cfca16c719eba9a72","coverage":[{"denominator":96,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":96,"source":"paper_references, paper_reference_links","source_observed_at":"2026-05-21T08:39:31.911497Z","state":"measured"},{"denominator":98,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":98,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-07T06:34:17.273281+00:00","state":"measured"},{"denominator":2,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":2,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-02T11:13:59.715326Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"pith","source_observed_at":"2026-07-02T02:16:27.208025Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"cited_work":{"arxiv_id":"2605.06534","doi":null,"metadata_source":"pith","pith_arxiv_id":"2605.06534","snapshot_observed_at":"2026-07-02T02:16:27.208025Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","venue":"cs.DC","work_id":"d4e4bc9c-c747-4e84-9433-546227fd5365","year":2026},"citing_paper":{"arxiv_id":"2606.03077","last_updated":"2026-06-10T06:28:18Z","snapshot_observed_at":"2026-07-06T23:43:22.513369Z","submitted_at":"2026-06-02T03:09:13Z","title":"Libra: Efficient Resource Management for Agentic RL Post-Training","version":2},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-06-28T11:02:00.385932Z"},"links":{"cited_paper":"/paper/2605.06534","citing_paper":"/paper/2606.03077"},"observation_digest":"sha256:c75da990a51cff203f83091ec0f85676c76263190a7c6cccf55c533fdc7699f6","observation_id":"13355f68-f65e-441d-8bb5-0c4141360bf7","resolution":{"observed_at":"2026-07-02T02:16:27.209193Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2605.06534","snapshot_observed_at":"2026-08-02T11:13:59.715326Z","title":"2026.ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2607.22614","last_updated":"2026-06-15T09:41:09Z","snapshot_observed_at":"2026-08-07T04:23:01.926858Z","submitted_at":"2026-06-15T09:41:09Z","title":"DynaResize: Runtime GPU Reallocation for Disaggregated LLM Post-Training","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-02T11:13:59.715326Z"},"links":{"cited_paper":"/paper/2605.06534","citing_paper":"/paper/2607.22614"},"observation_digest":"sha256:7c2355b0eb951fb42385a164f28d1c35d850fda8b9585883b2f4701c17e1b6ad","observation_id":"a36a1558-768b-40de-9ae5-7aceb0913ee1","resolution":{"observed_at":"2026-08-02T11:13:59.715326Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2605.06534/citation-record","integrity":"/paper/2605.06534/integrity","json":"/paper/2605.06534/citation-record.json","paper":"/paper/2605.06534"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"f595e682-fae1-4b9e-849d-08ea4273c3d4","year":2026},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:a30b6b9f33b7adfef7d13718a74c0844ead9d9bd1a20ed88532366effe625683","observation_id":"3349a32d-f86a-4dba-ace2-14acadab4302","resolution":{"observed_at":"2026-05-21T08:39:53.878845Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Li, Ryota Tomioka, and Milan Vojnovic","venue":null,"work_id":"69321d8c-dfd9-4a1d-877a-25fa31aa4596","year":2017},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:9e6b7971661d2f29a0d1d1285731d82f2673e17084819e190bddd6c893737e37","observation_id":"8ed6382b-2ac2-4e62-975e-2d375f549d36","resolution":{"observed_at":"2026-05-21T08:39:53.871505Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"d2319df8-f4bd-4cbe-88fc-d7bdd06ba877","year":2022},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:3e8295e9a4ed125d974f9b6f6534521df2af4469c5c645e337815af3e05e8adf","observation_id":"1d90ea0c-1c10-47c8-9e93-ff37258de499","resolution":{"observed_at":"2026-05-21T08:39:53.873579Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Hegde, Connor Chen, Charlie Ruan, Tyler Griggs, Shu Liu, Eric Tang, Richard Liaw, Philipp Moritz, Matei Zaharia, Joseph E","venue":null,"work_id":"54e74c47-4eb6-4cb9-b149-89a0902cadad","year":null},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:84d9236e131e9665dd3cddeedcbc297530e8f0b0416e76cf3b0f736a8077016a","observation_id":"13a92ec5-dd67-4210-9455-541276dc9aa2","resolution":{"observed_at":"2026-05-21T08:39:53.883615Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2511.16108","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-03T21:28:58.506549Z","title":"arXiv preprint arXiv:2511.16108(2025)","venue":null,"work_id":"cd691b36-09ff-4fda-9a35-a206d861e133","year":2025},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:ef6b3755d1a04449f138545b92ca602f8cbd5e63133beeb1a4bc375c9b85a1ef","observation_id":"79ec03cb-c279-484f-8458-da9c4a46a58a","resolution":{"observed_at":"2026-05-21T08:39:53.183793Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2511.16193","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"2b8e8d75-887e-48f7-b243-8eec9f925636","year":2025},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:c442e46253c86b113f33b334148ce5d64f54e404ed1730d033fb94d88743c0eb","observation_id":"07faa37e-2442-405a-9ca9-8375c583dcd8","resolution":{"observed_at":"2026-05-21T08:39:53.180529Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-06T13:42:37.037119Z","title":null,"venue":null,"work_id":"878d291e-4563-4328-b990-a8234c4ca86e","year":null},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:e8783da0cb6993f0cb49a0e1e277572bbea8157db79eced81f68accfe1c3634d","observation_id":"77d40d98-79fe-4953-a270-05202872ba14","resolution":{"observed_at":"2026-05-21T08:39:53.881126Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2510.09665","doi":"10.48550/arxiv.2510.09665","metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Lmcache: An efficient kv cache layer for enterprise-scale llm inference","venue":"arXiv (Cornell University)","work_id":"089b937e-f3f9-4525-a792-524ef0f7db1d","year":2025},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:9c072680e8dce84c8f8cb201a58766a1bf8e30ea3e9f741590d35ed27230cd2c","observation_id":"7b6feaa1-de39-4be6-b913-be134af9fc20","resolution":{"observed_at":"2026-05-21T08:39:53.351843Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"1d484c88-43cf-4f72-b19c-6348e043c58e","year":2024},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:109bc72e0567a69f1fea6ee82e0a8b2a536dbaf13e78f8719eac04ce3f9f43e4","observation_id":"7c41b9ce-aa19-40de-9e81-9d6b3325e23b","resolution":{"observed_at":"2026-05-21T08:39:53.875983Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"cfd17f06-ef71-4f51-a818-9133ebd19a99","year":2022},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:98323c7fb35c496e726942df3f4b694d5b1c581fc5e0f0d62ece6d6173e60169","observation_id":"a6258ea0-4391-44f3-9917-9a06a4b0c87b","resolution":{"observed_at":"2026-05-21T08:39:53.855575Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"9f191fa9-8641-4241-a0b6-e718763b5f36","year":2024},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:d18d0bfd1fbccf0c1bb174325848218dcf6b2d647799b375d9378bc4aff23172","observation_id":"6ca2e847-022a-4b2b-af0f-431df60dfd0f","resolution":{"observed_at":"2026-05-21T08:39:53.851666Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"67f410ad-b00b-4cd4-a804-8235cc8ab275","year":null},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:fe5d761294099dbce1d9398fc08f8c782be7f6fad0240882e0e7633a20afc2cd","observation_id":"a0b35315-4e73-4c24-8468-442ff27ad9bd","resolution":{"observed_at":"2026-05-21T08:39:53.793415Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"InProceedings of the 2021 ACM SIGCOMM 2021 Conference","venue":null,"work_id":"e70e444a-8954-4e93-b08d-3ab2b385c48c","year":2021},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:d358f9458a219ea2c4527b67c0599f14a077ca715f75ba4ccb2cb108b4647d3e","observation_id":"4954c3d8-c85b-4a3f-bc84-64fb8342e5c9","resolution":{"observed_at":"2026-05-21T08:39:53.785712Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.10978","last_updated":"2025-10-28T15:11:36Z","snapshot_observed_at":"2026-07-29T19:20:21.974239Z","submitted_at":"2025-05-16T08:26:59Z","title":"Group-in-Group Policy Optimization for LLM Agent Training","version":3},"cited_work":{"arxiv_id":"2505.10978","doi":"10.48550/arxiv.2505.10978","metadata_source":"pith","pith_arxiv_id":"2505.10978","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Group-in-Group Policy Optimization for LLM Agent Training","venue":"cs.LG","work_id":"bc65d492-e6ba-4522-874c-43d2f4fc5191","year":2025},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"cited_paper":"/paper/2505.10978","citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:8332c00bbb28b3ff46ae4e2deb9b68a9d0ada7364080d2c8d27ee6ec286fb0f9","observation_id":"1c034139-ad48-4e83-8dd0-bdb1208ffa75","resolution":{"observed_at":"2026-05-21T08:39:53.188337Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"abf486e4-83c6-4f61-8512-17d624aa6dee","year":2024},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:a75bc9babe1faff8c27bdd2c3d4f2bafd3665569374a5100c2e3943a21d7c402","observation_id":"c8d41c7b-54d7-49c3-956c-9f7b5f1f58f5","resolution":{"observed_at":"2026-05-21T08:39:53.862397Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2025.35531","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"2730b816-f693-4cdb-9172-9745e2842e01","year":2025},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:84786a3d324f76a2e88b83a4e0e53c631bab32c742b4971fb07bc6a6c033a566","observation_id":"7954dc54-1c83-47d4-9ed3-518e56ccac53","resolution":{"observed_at":"2026-05-21T08:39:53.337920Z","resolver_source":"arxiv_id","status":"malformed_identifier"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2509.21009","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-04T16:29:56.817095Z","title":"Rollpacker: Mitigating long-tail rollouts for fast, synchronous rl post-training","venue":null,"work_id":"5a098eae-1c24-47d3-8e41-f32bd7dbd571","year":2025},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:d74585b793f9cad33a650aea67a03fef13d446531595200aff028f6b1c6ea3b5","observation_id":"9cbd6f9b-8f3e-4349-8e24-2d4e9cdbb3ac","resolution":{"observed_at":"2026-05-21T08:39:53.322937Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2512.22560","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-03T21:28:58.512098Z","title":"Rollart: Scaling agentic rl training via disaggregated infrastructure.arXiv preprint arXiv:2512.22560","venue":null,"work_id":"21e4b275-4c2c-4896-adaf-a7001b4799c7","year":2026},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:71ec25016de005e33214eac414e1f1abb16a5d726726ea5761376cff2e299b09","observation_id":"f2dec488-6840-4491-b238-4a68e916fd8e","resolution":{"observed_at":"2026-05-21T08:39:53.297334Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"15cf978b-403f-46ba-8cfe-ce1c4cde32c7","year":null},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:6c68783535cd29c67037290608d7aa093e48f9419ad4f4b74a0155bd93dabd21","observation_id":"3a6dcea3-9979-4a8a-80c7-25e9f0f4647f","resolution":{"observed_at":"2026-05-21T08:39:53.778185Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"In16th USENIX Symposium on Oper- ating Systems Design and Implementation (OSDI 22)","venue":null,"work_id":"a2be7671-005f-4fa6-9a75-3f1cb96a6872","year":null},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:a7f365a9379eddef461cdce2bf854a8993238b8fc9dfddf71f10f79c9f32856c","observation_id":"3a5b35f5-75b6-4b8f-83cb-0abadb379eba","resolution":{"observed_at":"2026-05-21T08:39:53.781065Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2507.01663","last_updated":"2025-07-02T12:45:34Z","snapshot_observed_at":"2026-08-06T20:43:52.277899Z","submitted_at":"2025-07-02T12:45:34Z","title":"AsyncFlow: An Asynchronous Streaming RL Framework for Efficient LLM Post-Training","version":1},"cited_work":{"arxiv_id":"2507.01663","doi":"10.48550/arxiv.2507.01663","metadata_source":"arxiv_reference","pith_arxiv_id":"2507.01663","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Asyncflow: An asynchronous streaming rl framework for efficient llm post-training","venue":"ArXiv.org","work_id":"c78dd69f-0c92-450b-8e2b-dec1ded07146","year":2025},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"cited_paper":"/paper/2507.01663","citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:6d89e540c08d6919f68ad31a2a8bf8ff22d7eda4591b939c783ea41c6da2d458","observation_id":"e5ed996a-e8b6-4c17-a9f2-f131e885f42d","resolution":{"observed_at":"2026-05-21T08:39:53.267980Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2508.05118","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-30T19:05:00.371021Z","title":"Reasoning through exploration: A reinforcement learning framework for robust function calling","venue":null,"work_id":"17ac5f84-5524-4185-9df0-e9be7d77267f","year":2025},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:7722541f9aa2eb2c4e33724cb07ce06761f293d2f210480e4463e010cb6d321c","observation_id":"f3e0f8f4-5489-489c-8e95-fec71d144528","resolution":{"observed_at":"2026-05-21T08:39:53.260846Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"7412.777462","doi":"10.1145/777412.777462","metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Squillante","venue":null,"work_id":"30775000-bdac-4cde-8354-25a63739b772","year":2003},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:e3f28edf1eb634fd6aab15f65e2ba23c737593bca0313c63faf5825547372059","observation_id":"b06f79a5-9744-48b9-92ff-60ae59334d14","resolution":{"observed_at":"2026-05-21T08:39:53.019142Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"28145c48-cca1-4271-a486-f843e25b846a","year":2025},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:30d29214dd0adf65dbb77a41ae114a8769db9d0ecd46b6b1219bce9011727a6b","observation_id":"532853be-3f57-45bc-b651-5d8fc21e6eb4","resolution":{"observed_at":"2026-05-21T08:39:53.833993Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2508.18588","last_updated":"2025-08-26T01:42:46Z","snapshot_observed_at":"2026-08-05T16:20:23.166001Z","submitted_at":"2025-08-26T01:42:46Z","title":"History Rhymes: Accelerating LLM Reinforcement Learning with RhymeRL","version":1},"cited_work":{"arxiv_id":"2508.18588","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2508.18588","snapshot_observed_at":"2026-07-03T23:59:07.264671Z","title":"History rhymes: Accelerating llm reinforcement learning with rhymerl","venue":null,"work_id":"2b77db42-385a-4db7-8b75-d164235015ec","year":2025},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"cited_paper":"/paper/2508.18588","citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:6be944ecc782ef41d303efefc3b494e35e7fabf8a64f15b7f4a48bd2c4363e9c","observation_id":"2cbc7f6b-7da4-453e-aaa2-1a48b797e015","resolution":{"observed_at":"2026-05-21T08:39:53.249603Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2405.11143","last_updated":"2025-10-09T12:22:46Z","snapshot_observed_at":"2026-07-31T12:28:37.704994Z","submitted_at":"2024-05-20T01:04:40Z","title":"OpenRLHF: An Easy-to-use, Scalable and High-performance RLHF Framework","version":6},"cited_work":{"arxiv_id":"2405.11143","doi":"10.48550/arxiv.2405.11143","metadata_source":"pith","pith_arxiv_id":"2405.11143","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"OpenRLHF: An Easy-to-use, Scalable and High-performance RLHF Framework","venue":"cs.AI","work_id":"70fa48c9-2f84-49f6-9aca-37476e021fc3","year":2024},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"cited_paper":"/paper/2405.11143","citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:3da0435ce60b04aca9a0cff08892cbcbeafa533d65316e66851b126157e65545","observation_id":"2a6ed7cb-93a5-42a9-84b2-9d9424fb38b8","resolution":{"observed_at":"2026-05-21T08:39:53.278478Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"cd0ec959-ac32-41c0-9937-4921c5812d89","year":2024},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:1e9d90cea4de2ba90a6b8aba4d7e6a03dee63a404c8eb07e9474f67c5254d440","observation_id":"d28ac56f-5e79-46ca-a868-240486f2cdb6","resolution":{"observed_at":"2026-05-21T08:39:53.866965Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2310.06770","last_updated":"2024-11-11T23:05:04Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-10-10T16:47:29Z","title":"SWE-bench: Can Language Models Resolve Real-World GitHub Issues?","version":3},"cited_work":{"arxiv_id":"2310.06770","doi":"10.1145/512927.512945","metadata_source":"pith","pith_arxiv_id":"2310.06770","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SWE-bench: Can Language Models Resolve Real-World GitHub Issues?","venue":"cs.CL","work_id":"d0effe15-a689-441a-8e3f-ea35f1c4e4b1","year":2023},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"cited_paper":"/paper/2310.06770","citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:99280d325438be554ad0b09d7d44f5f73923dcb9b014c08300cc4c6dc4e141c3","observation_id":"7b734131-75c3-4b5f-97fe-b5bb97a97f36","resolution":{"observed_at":"2026-05-21T08:39:53.347767Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"2b83a09e-e0ae-4958-88cd-4df7bb3d4915","year":2023},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:3fa09211f3d21ee264f7ef66601b511c5277a54baba3f64dd9bd63b1f0e159ff","observation_id":"4cea67d3-dbd9-48e5-8b05-08c7965e3481","resolution":{"observed_at":"2026-05-21T08:39:53.829779Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Gonzalez, Hao Zhang, and Ion Sto- ica","venue":null,"work_id":"15d0d506-824d-4d41-b86e-1e4fc065401a","year":2023},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:83e1425e507dde92711690f028304e7586ba6bde97d7f7521e682cffb88a104a","observation_id":"935c6620-bd04-4b0b-9819-286198344990","resolution":{"observed_at":"2026-05-21T08:39:53.823039Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2326.358744","doi":"10.1145/3552326.3587442","metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Lyra: Elastic scheduling for deep learning clusters","venue":null,"work_id":"e1307e19-6ca6-403c-82bc-1fce2a41d52a","year":2023},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:ef756c3dbb4577d56eafb5977b805cc1eaa450f33de81ac908705b3e3ce32c9f","observation_id":"488a1f7d-04fe-44d0-9060-d8e5e212aa49","resolution":{"observed_at":"2026-05-21T08:39:53.007806Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-07-11T03:19:20.798897+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T03:19:20.798897+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"1945bde8-9a5d-411f-b031-bdadfa293642","year":2025},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:2c4924c58fc155cd078b3a2f5904ff43a151d0d726f4e4c0aa27f29c7a76ab71","observation_id":"e3c70477-f1f8-4673-8a77-dc756219ff73","resolution":{"observed_at":"2026-05-21T08:39:53.820922Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2508.19598","last_updated":"2025-08-27T06:19:50Z","snapshot_observed_at":"2026-08-05T15:44:13.449330Z","submitted_at":"2025-08-27T06:19:50Z","title":"Encouraging Good Processes Without the Need for Good Answers: Reinforcement Learning for LLM Agent Planning","version":1},"cited_work":{"arxiv_id":"2508.19598","doi":null,"metadata_source":"pith","pith_arxiv_id":"2508.19598","snapshot_observed_at":"2026-07-09T11:16:11.340947Z","title":"Encouraging good processes without the need for good answers: Reinforcement learning for llm agent planning","venue":"cs.LG","work_id":"9d394028-e798-472d-8802-1ff36b428157","year":2025},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"cited_paper":"/paper/2508.19598","citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:22a7979e2c69ae2c6981223fe9d100d942e1546656bda517e7548b25f858adb5","observation_id":"a4c0696b-4908-4827-8f28-deeddd9db38d","resolution":{"observed_at":"2026-05-21T08:39:53.355030Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"b943d61f-e644-4324-965b-b977c977236e","year":2023},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:bbf30c764ef74a9c63186e5abd86f9c37d5c59082fd0d89e2d78d173f9987913","observation_id":"1edfcfbf-f389-4c0e-bca3-04a74ca3b040","resolution":{"observed_at":"2026-05-21T08:39:53.831862Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"66ae6a41-1cc0-424c-b442-cd23fea1fa1b","year":null},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:f76e4bd73d997c8901ca09e0ea2bd8f01ac95e9015445e314b5091c0a272ea1b","observation_id":"89169ff8-d568-4f42-a044-c90539807a97","resolution":{"observed_at":"2026-05-21T08:39:53.860455Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"InProceedings of the ACM SIGCOMM 2024 Con- ference","venue":null,"work_id":"6e747e44-d156-4ba0-808d-1302c5adaf79","year":2024},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:2e08b78d385fa98c5d8c08adb317be5e780da24d73a383dd41a2c4950c8d2cce","observation_id":"faa9b3c0-a1ba-4dba-b3f6-2dd607b8e6bf","resolution":{"observed_at":"2026-05-21T08:39:53.840566Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.14239","last_updated":"2025-04-19T09:25:55Z","snapshot_observed_at":"2026-08-05T23:52:52.919433Z","submitted_at":"2025-04-19T09:25:55Z","title":"InfiGUI-R1: Advancing Multimodal GUI Agents from Reactive Actors to Deliberative Reasoners","version":1},"cited_work":{"arxiv_id":"2504.14239","doi":null,"metadata_source":"pith","pith_arxiv_id":"2504.14239","snapshot_observed_at":"2026-07-09T18:26:26.304911Z","title":"InfiGUI-R1: Advancing Multimodal GUI Agents from Reactive Actors to Deliberative Reasoners","venue":"cs.AI","work_id":"a7d6f6c7-8a57-4a83-9c4e-892cdc190280","year":2025},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"cited_paper":"/paper/2504.14239","citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:cb2b89d4782d76c017e044208b9ec7e58b75caac4feec6124d0a621a373e971f","observation_id":"d5410e0e-dcf0-450b-9c96-ca8b218fb5a8","resolution":{"observed_at":"2026-05-21T08:39:53.204732Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2510.11345","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-09T08:56:06.417978Z","title":"Part ii: Roll flash–accelerating rlvr and agentic training with asynchrony","venue":null,"work_id":"810c144e-2d15-4352-a654-0f8af536eb98","year":2025},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:53fe50e502f7b682abae149ce5ebe889b512cebb9cdcfa0a4ab7497027ccd15a","observation_id":"3bac49d6-de77-4a27-9739-59c3cea682c7","resolution":{"observed_at":"2026-05-21T08:39:53.293753Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2503.21620","last_updated":"2025-05-24T08:46:08Z","snapshot_observed_at":"2026-08-01T20:06:59.931737Z","submitted_at":"2025-03-27T15:39:30Z","title":"UI-R1: Enhancing Efficient Action Prediction of GUI Agents by Reinforcement Learning","version":5},"cited_work":{"arxiv_id":"2503.21620","doi":"10.48550/arxiv.2503.21620","metadata_source":"pith","pith_arxiv_id":"2503.21620","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"UI-R1: Enhancing Efficient Action Prediction of GUI Agents by Reinforcement Learning","venue":"cs.AI","work_id":"4637c89b-db94-4e6f-8bf2-030bea2fdd6e","year":2025},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":41,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"cited_paper":"/paper/2503.21620","citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:24754e6e0b1e473dc528f40322a79da6889320944d602c96b66e9428bd4cfe71","observation_id":"2a90d8cb-c65c-44d8-840e-40caa402f528","resolution":{"observed_at":"2026-05-21T08:39:53.245323Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"13e6416f-7fe4-4139-824c-3d6ac80432aa","year":2025},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:0efdaa3074b4cf14cb139463971a022e532bb2995b728c97fd39d0833e845f5c","observation_id":"e26c7267-d37a-493f-8e80-818c0a4f712c","resolution":{"observed_at":"2026-05-21T08:39:53.838527Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.10458","last_updated":"2025-10-01T04:55:20Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-04-14T17:45:54Z","title":"GUI-R1 : A Generalist R1-Style Vision-Language Action Model For GUI Agents","version":4},"cited_work":{"arxiv_id":"2504.10458","doi":null,"metadata_source":"pith","pith_arxiv_id":"2504.10458","snapshot_observed_at":"2026-07-04T14:19:55.071333Z","title":"GUI-R1 : A Generalist R1-Style Vision-Language Action Model For GUI Agents","venue":"cs.CV","work_id":"5e82d316-7129-4f55-9c00-0d7fcbcea139","year":2025},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"cited_paper":"/paper/2504.10458","citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:d062ecb9d917beee59c6180a920acb12d081468ebd85cf3df09e4590dc32a456","observation_id":"8f5dd2b7-abcf-4f5d-bfbc-623316fb8260","resolution":{"observed_at":"2026-05-21T08:39:53.344575Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"3e70475a-989c-4509-b5bb-0fe3ba3b54c0","year":2018},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:9de0b5721df2b2a532ddb9d26501a7fa36c30a8398c030305aa0c27264e95608","observation_id":"e3da2429-473e-4494-9dea-c1f64d4d5b12","resolution":{"observed_at":"2026-05-21T08:39:53.836286Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"fd315633-4061-4ae1-8738-0143b37e19ac","year":2024},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:05414d65785087f020167e458bc7d81e3a7b84e33be4560a50a3921db64c342d","observation_id":"53e25f4c-542d-48b5-901b-2d66438f1853","resolution":{"observed_at":"2026-05-21T08:39:53.845158Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"4d737962-ce2b-4d44-be9f-13c62990746d","year":2025},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:b998ea7a51e080eed301c38a715cafa936c54d36307d13603ed781521d142bdc","observation_id":"aee24785-ea41-45d9-93e1-b7a00b666255","resolution":{"observed_at":"2026-05-21T08:39:53.864580Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"9077.2024","doi":"10.1109/isca59077.2024","metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Reiner Pope, Sholto Douglas, Aakanksha Chowdhery, Jacob Devlin, James Bradbury, Anselm Lev- skaya, Jonathan Heek, Kefan Xiao, Shivani Agrawal, and Jeff Dean","venue":null,"work_id":"6a69794e-e9cd-48c9-9b7d-502b131ff6db","year":2024},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:3d2d1de399d2f42de74e44744a73e31cfdd86babc56aece39fffbb81e4e78328","observation_id":"d3b1c5be-ed10-4f7d-abe2-126030fe07ac","resolution":{"observed_at":"2026-05-21T08:39:53.024345Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.01228","last_updated":"2025-09-03T20:54:57Z","snapshot_observed_at":"2026-07-06T19:25:42.156795Z","submitted_at":"2024-10-02T04:12:13Z","title":"ConServe: Fine-Grained GPU Harvesting for LLM Online and Offline Co-Serving","version":2},"cited_work":{"arxiv_id":"2410.01228","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2410.01228","snapshot_observed_at":"2026-07-02T15:27:06.032587Z","title":"Gon- zalez, Ion Stoica, and Harry Xu","venue":null,"work_id":"8267ce10-9156-4dba-a201-7567a179fc73","year":2024},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"cited_paper":"/paper/2410.01228","citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:d245fab8eea2f83421d891034b9c559f1311418289a9f0a4107d823e2abd5731","observation_id":"dfbeedee-3cf8-47fd-8389-fc1a6f40e3ee","resolution":{"observed_at":"2026-05-21T08:39:53.282201Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"2ce8b22f-b8ac-4be3-bcb0-150558d7656a","year":null},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:80e4f4574e88ca67d56f3034545e4e8e534f3220a214321a45283f9171404cf6","observation_id":"67669c22-8f76-4493-a170-77c262721234","resolution":{"observed_at":"2026-05-21T08:39:53.869263Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2511.14617","last_updated":"2026-04-03T12:47:37Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-11-18T16:12:21Z","title":"Seer: Online Context Learning for Fast Synchronous LLM Reinforcement Learning","version":3},"cited_work":{"arxiv_id":"2511.14617","doi":null,"metadata_source":"pith","pith_arxiv_id":"2511.14617","snapshot_observed_at":"2026-07-04T06:49:38.047014Z","title":"Seer: Online Context Learning for Fast Synchronous LLM Reinforcement Learning","venue":"cs.DC","work_id":"03b265b6-317b-4574-a145-82e7e81c521d","year":2025},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"cited_paper":"/paper/2511.14617","citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:caad6c66901ff22fc6d312efb841136fc8814a1c1f318493650049a9d9e9c829","observation_id":"ffbfff9d-7a2c-4cfd-a8d2-3cf46caa70ab","resolution":{"observed_at":"2026-05-21T08:39:53.333910Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-08T09:14:47.291345Z","title":null,"venue":null,"work_id":"f06f75fe-b457-4ebd-a734-a6011bfb16c2","year":2024},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:8c6c850266fbb8d49965eeadc4257719e9575694f1d00122cf024769ba668639","observation_id":"d09091f9-49fd-45d4-a9e7-6b99992b5f7a","resolution":{"observed_at":"2026-05-21T08:39:53.847487Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"7928a99e-3bac-433e-8955-3f98a2357ac1","year":2025},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:b17aae4720843e7e6ed17d8b0c35dfd059437f65fc7b7a4cf36df55673673f76","observation_id":"8b82e5ed-2774-4890-a75d-f7685859c731","resolution":{"observed_at":"2026-05-21T08:39:53.827499Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.09970","last_updated":"2025-05-19T03:17:21Z","snapshot_observed_at":"2026-08-07T15:45:13.424803Z","submitted_at":"2025-05-15T05:17:47Z","title":"Pre-Act: Multi-Step Planning and Reasoning Improves Acting in LLM Agents","version":2},"cited_work":{"arxiv_id":"2505.09970","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2505.09970","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Pre-Act: Multi-step planning and reasoning improves acting in LLM agents","venue":null,"work_id":"db137848-b22e-4f62-bcfb-2762dd0e04c9","year":2025},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":53,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"cited_paper":"/paper/2505.09970","citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:3c58cbe1b9cd2a01b1274cb07bc5323a89a79d5f58d294f9cc214cb00977d8ac","observation_id":"5412bea4-b3ca-4ee3-980e-7c41b3819e9b","resolution":{"observed_at":"2026-05-21T08:39:53.221001Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2508.15706","last_updated":"2026-07-23T09:18:31Z","snapshot_observed_at":"2026-08-07T19:57:01.269382Z","submitted_at":"2025-08-21T16:48:19Z","title":"Overcoming the Communication-Performance Tradeoff in LLM Pretraining","version":3},"cited_work":{"arxiv_id":"2508.15706","doi":"10.48550/arxiv.2508.15706","metadata_source":"arxiv_reference","pith_arxiv_id":"2508.15706","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Communication efficient LLM pre-training with SparseLoCo","venue":"ArXiv.org","work_id":"55ed9ad5-b2b5-4b30-85f8-9c932ee25977","year":2025},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":54,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"cited_paper":"/paper/2508.15706","citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:5cb7048619cbc5e675b229378875f64ff6a8df5b5c07a483662948059145a9e6","observation_id":"9fb95a3d-52ce-4a82-ad20-b73f8a39e7da","resolution":{"observed_at":"2026-07-24T02:22:54.639277Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1802.05799","last_updated":"2018-02-21T04:30:30Z","snapshot_observed_at":"2026-07-06T06:23:45.215820Z","submitted_at":"2018-02-15T23:36:51Z","title":"Horovod: fast and easy distributed deep learning in TensorFlow","version":3},"cited_work":{"arxiv_id":"1802.05799","doi":"10.48550/arxiv.1802.05799","metadata_source":"pith","pith_arxiv_id":"1802.05799","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Horovod: fast and easy distributed deep learning in TensorFlow","venue":"cs.LG","work_id":"1437d05a-4ec5-4ff0-afdb-fd503478751e","year":2018},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"cited_paper":"/paper/1802.05799","citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:6c8581464ec76970936a296e13cbe6d4cdaf3ee23cc320c93d43b2831bce0923","observation_id":"f374cf90-96ae-46ac-a8e1-2baf92111f14","resolution":{"observed_at":"2026-05-21T08:39:53.228630Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"a1cfb5b7-68f7-4bc5-ac5d-d34804500e4b","year":2025},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":56,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:8179e0490f3c10ee93b662e3818e01c364c4e68c09803381f327ea2412511159","observation_id":"fcc3354a-5429-4fa3-8f4f-57c506373649","resolution":{"observed_at":"2026-05-21T08:39:53.808107Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2511.13841","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-03T12:58:08.719085Z","title":"Beat the long tail: Distribution-aware speculative decoding for rl training","venue":null,"work_id":"26b4c627-17fc-4861-837f-357caa283cf2","year":2025},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":57,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:4cdb972f3e5f96f5095bf563aa46d2d2b7b8ec931c54ec82036ef3b1e8b0df09","observation_id":"317cb998-9f95-4219-8efb-72cae6b151e7","resolution":{"observed_at":"2026-05-21T08:39:53.232970Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.03300","last_updated":"2024-04-27T15:25:53Z","snapshot_observed_at":"2026-08-06T14:58:42.911363Z","submitted_at":"2024-02-05T18:55:32Z","title":"DeepSeekMath: Pushing the Limits of Mathematical Reasoning in Open Language Models","version":3},"cited_work":{"arxiv_id":"2402.03300","doi":"10.1016/0004-3702(73)90011-8","metadata_source":"pith","pith_arxiv_id":"2402.03300","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"DeepSeekMath: Pushing the Limits of Mathematical Reasoning in Open Language Models","venue":"cs.CL","work_id":"c5006563-f3ec-438a-9e35-b7b484f34828","year":2024},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":58,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"cited_paper":"/paper/2402.03300","citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:4f09c2a3d326c48567b2e3effef52287fe20cf0e259278183efb4e222dc2a6b2","observation_id":"1b4dd8b8-85a4-40d2-967c-bc41d1c95734","resolution":{"observed_at":"2026-05-21T08:39:53.307881Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2510.12633","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-04T16:29:56.839451Z","title":"Laminar: A scalable asyn- chronous rl post-training framework","venue":null,"work_id":"bcbca89e-28ee-4440-9d4e-eb8a792cce4d","year":2025},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":59,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:486e1c8d7fbfa85a80abf0ff416546e0981b821e3295344e8f98e2b795521cd5","observation_id":"e31dceb3-4787-4f5d-946d-76d96bac048d","resolution":{"observed_at":"2026-05-21T08:39:53.304646Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2409.19256","last_updated":"2024-10-02T04:01:47Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-09-28T06:20:03Z","title":"HybridFlow: A Flexible and Efficient RLHF Framework","version":2},"cited_work":{"arxiv_id":"2409.19256","doi":"10.1145/3689031.3696075.url:","metadata_source":"pith","pith_arxiv_id":"2409.19256","snapshot_observed_at":"2026-07-11T11:50:26.030339Z","title":"HybridFlow: A Flexible and Efficient RLHF Framework","venue":"cs.LG","work_id":"7eb9c9f4-b322-4bba-8011-09ff8d6ad801","year":2024},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":60,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"cited_paper":"/paper/2409.19256","citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:f4133f58d47145e9fce2cb6e49432a83c32141477a102f17e35cc32364529b1e","observation_id":"d94bd0d9-9332-4869-8cde-9a871cc3d6bf","resolution":{"observed_at":"2026-05-21T08:39:53.212493Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"6396ae4a-1dbe-43ca-9783-381f91da5a89","year":2024},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":61,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:b2c3d2cb3933a1a51a7ceab0cee06e2d2f9f28b047a217b9a21bd443eac7eac5","observation_id":"d1479110-4e15-44e6-baf9-4788c6469cc1","resolution":{"observed_at":"2026-05-21T08:39:53.849518Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"1909.08053","last_updated":"2020-03-13T23:45:18Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2019-09-17T19:42:54Z","title":"Megatron-LM: Training Multi-Billion Parameter Language Models Using Model Parallelism","version":4},"cited_work":{"arxiv_id":"1909.08053","doi":"10.48550/arxiv.1909.08053","metadata_source":"pith","pith_arxiv_id":"1909.08053","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Megatron-LM: Training Multi-Billion Parameter Language Models Using Model Parallelism","venue":"cs.CL","work_id":"c888e6d1-0b1d-43d6-9ef5-f0912a0efa1b","year":2019},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":62,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"cited_paper":"/paper/1909.08053","citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:e5700e541a10289bd53a9b1929aac233b1d55991640d761b7634d5bb28b9f40b","observation_id":"104d01d8-9f95-41d1-8650-54029876d281","resolution":{"observed_at":"2026-05-21T08:39:53.285969Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-07-09T10:48:33.392193+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-09T10:48:33.392193+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2010.03768","last_updated":"2021-03-14T22:44:38Z","snapshot_observed_at":"2026-07-06T10:02:33.297722Z","submitted_at":"2020-10-08T05:13:36Z","title":"ALFWorld: Aligning Text and Embodied Environments for Interactive Learning","version":2},"cited_work":{"arxiv_id":"2010.03768","doi":"10.1109/cvpr.2001.990517","metadata_source":"pith","pith_arxiv_id":"2010.03768","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"ALFWorld: Aligning Text and Embodied Environments for Interactive Learning","venue":"cs.CL","work_id":"fa436f46-ec0a-4d2e-a0ff-e697def4a7be","year":2020},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":63,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"cited_paper":"/paper/2010.03768","citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:fed6af73968ca55ec32534747aa006bd2c5b11ccf9de8812f56e5c10e4104bf7","observation_id":"4d855fa2-a882-4567-a143-be86c2834a3a","resolution":{"observed_at":"2026-05-21T08:39:53.315343Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-05-20T11:23:33.289391+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-20T11:23:33.289391+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"0f1cf410-b9f3-48e2-9645-7df158816bb9","year":null},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":64,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:46e422042aeccedd7936ff0e8ba52dfee6313c98317adaab31aba8284a053178","observation_id":"921708fe-d39f-486d-a71a-3ef94d47c76b","resolution":{"observed_at":"2026-05-21T08:39:53.795951Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.01441","last_updated":"2025-04-28T10:42:49Z","snapshot_observed_at":"2026-08-05T20:36:54.358831Z","submitted_at":"2025-04-28T10:42:49Z","title":"Agentic Reasoning and Tool Integration for LLMs via Reinforcement Learning","version":1},"cited_work":{"arxiv_id":"2505.01441","doi":"10.48550/arxiv.2505.01441","metadata_source":"pith","pith_arxiv_id":"2505.01441","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Agentic Reasoning and Tool Integration for LLMs via Reinforcement Learning","venue":"cs.AI","work_id":"6d9faf3b-cd39-48ff-8918-a9379716abd2","year":2025},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":65,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"cited_paper":"/paper/2505.01441","citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:553ae598b56693d5d8a1b5825ddce55dd7ec7fec6be884c4a46a87e9cbc48887","observation_id":"b532f11c-e48a-4a89-9802-8d45e03d1b65","resolution":{"observed_at":"2026-05-21T08:39:53.264247Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"b0fbbabd-b5f4-456a-b64a-f345630bdb13","year":2025},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":66,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:6090d3a54be11bbdf271c72b26dd3ab99e467d320b6d92e8154899087d4f2742","observation_id":"8c76313f-bd34-4fc1-988e-b97743da1a34","resolution":{"observed_at":"2026-05-21T08:39:53.843097Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"216b39d3-fe27-41eb-a43b-826554da1d30","year":2025},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":67,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:f1792d671f712358569cd5367ae690be51babd5e2971b76b0c3a57dd212f7aee","observation_id":"05def486-d3a4-4b0e-8b25-28e5ea7c9d97","resolution":{"observed_at":"2026-05-21T08:39:53.790760Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"cb3ab048-77d0-492a-841f-9ee458a8a3ac","year":2025},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":68,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:47ab90506aabc00b280686ab0a467457299413a831514f88a2715826d8bcbae9","observation_id":"1405bbff-9061-4362-aac0-643eb4127f38","resolution":{"observed_at":"2026-05-21T08:39:53.800756Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"286cfcb7-d7f9-4eb7-b83c-956f9481faf4","year":2025},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":69,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:b3ccb17c2f3502c265a9439a3dcadd5fa652c7769ebe0274d8a8515fe83713ce","observation_id":"322a9c6b-976c-4533-aa96-d91e5967aeaf","resolution":{"observed_at":"2026-05-21T08:39:53.803790Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2506.06122","last_updated":"2025-06-06T14:33:56Z","snapshot_observed_at":"2026-08-07T21:55:43.371585Z","submitted_at":"2025-06-06T14:33:56Z","title":"Reinforcement Learning Optimization for Large-Scale Learning: An Efficient and User-Friendly Scaling Library","version":1},"cited_work":{"arxiv_id":"2506.06122","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2506.06122","snapshot_observed_at":"2026-07-04T10:29:45.796847Z","title":"Reinforcement learning optimization for large-scale learning: An efficient and user-friendly scaling library.arXiv preprint arXiv:2506.06122, 2025a","venue":null,"work_id":"37914127-0152-4730-8654-8b15734bd5f9","year":2025},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":70,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"cited_paper":"/paper/2506.06122","citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:039f815cf7d763476d1657b5c68fb850832eb67b61a80ab554e1a28d9e6210dc","observation_id":"d2f7b234-e637-4a7a-8c4c-15610da0c84e","resolution":{"observed_at":"2026-05-21T08:39:53.289906Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2512.24873","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-01T21:26:14.207661Z","title":"Let it flow: Agentic crafting on rock and roll, building the rome model within an open agentic learning ecosystem","venue":null,"work_id":"a1afde43-96e3-49c9-af14-25f128d65fe3","year":2025},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":71,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:7dd5f6e4d94fd75f2eacbdf7aef9030721df33a20a35b7a783222dea3591ae18","observation_id":"b562ace0-60ba-4843-97d6-141678789042","resolution":{"observed_at":"2026-05-21T08:39:53.301248Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2401.17644","last_updated":"2025-05-26T16:16:43Z","snapshot_observed_at":"2026-08-06T05:24:39.144385Z","submitted_at":"2024-01-31T07:52:48Z","title":"BurstGPT: A Real-world Workload Dataset to Optimize LLM Serving Systems","version":5},"cited_work":{"arxiv_id":"2401.17644","doi":null,"metadata_source":"pith","pith_arxiv_id":"2401.17644","snapshot_observed_at":"2026-07-07T20:34:09.759421Z","title":"Burstgpt: A real-world workload dataset to optimize llm serving systems","venue":"cs.DC","work_id":"34577abd-227c-452e-afb3-e3922e1bda04","year":2024},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":72,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"cited_paper":"/paper/2401.17644","citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:422c95a069d92ad97db8486a090f3209d71f0d1ba50802fb04a6cfaeaac36f3d","observation_id":"34eb5215-4a35-4947-8f3d-1b84e4ed133e","resolution":{"observed_at":"2026-05-21T08:39:53.241704Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"6572ec8d-e34d-4d44-8716-3f8da446d263","year":2025},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":73,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:921cbf437d9de8f8ad2189f89993928b6ad89797a5e1fe702a704a67dd189f96","observation_id":"bde49e18-d5bd-44c1-a0df-353fa0ad7f91","resolution":{"observed_at":"2026-05-21T08:39:53.818676Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"0077de2b-4eab-4412-8d82-defbe02745da","year":2025},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":74,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:a2b8fb926f8f76ff8fc0f506a7effafe8bd0880776cce3bac7896196bc3ad478","observation_id":"e8f98d5e-a0f1-424e-aa0f-bb946ff89df3","resolution":{"observed_at":"2026-05-21T08:39:53.810675Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2512.11306","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-04T02:49:24.704621Z","title":"Rollmux: Phase-level multiplex- ing for disaggregated rl post-training","venue":null,"work_id":"b6e4e5d2-205f-4779-bf37-49c151f932d7","year":2025},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":75,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:f6564f63e0737b3e52633a7bd5ddb437ba438d58e3a2c47f4ae4e644b3a5ba07","observation_id":"a969a22d-1863-4f37-9879-4282b2d72c0e","resolution":{"observed_at":"2026-05-21T08:39:53.271520Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2510.19225","last_updated":"2026-04-08T03:04:11Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-10-22T04:19:37Z","title":"RLBoost: Harvesting Preemptible Resources for Cost-Efficient Reinforcement Learning on LLMs","version":3},"cited_work":{"arxiv_id":"2510.19225","doi":null,"metadata_source":"pith","pith_arxiv_id":"2510.19225","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"RLBoost: Harvesting Preemptible Resources for Cost-Efficient Reinforcement Learning on LLMs","venue":"cs.DC","work_id":"307defa2-eea6-4ad3-96bf-98aa7ec127c9","year":2025},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":76,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"cited_paper":"/paper/2510.19225","citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:1921d61c01babe9fa2fdc82283a13088897bf1c183e8e32d711250ed03152dab","observation_id":"07c4c85e-3489-4c37-935a-494b0781c983","resolution":{"observed_at":"2026-05-21T08:39:53.192515Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.07608","last_updated":"2025-06-05T11:49:09Z","snapshot_observed_at":"2026-08-07T15:47:50.795694Z","submitted_at":"2025-05-12T14:30:11Z","title":"MiMo: Unlocking the Reasoning Potential of Language Model -- From Pretraining to Posttraining","version":2},"cited_work":{"arxiv_id":"2505.07608","doi":"10.48550/arxiv.2505.07608","metadata_source":"arxiv_reference","pith_arxiv_id":"2505.07608","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"MiMo: Unlocking the reasoning potential of language model–from pretraining to posttraining","venue":"ArXiv.org","work_id":"c6dfcaff-bd30-4b3d-a25f-f9a09176940c","year":2025},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":77,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"cited_paper":"/paper/2505.07608","citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:80887086af6ac4089def3ab93dd332417ef71b769d89316007cd1d45f3cb85fd","observation_id":"e05724d1-6c81-4d84-b112-d55252974934","resolution":{"observed_at":"2026-05-21T08:39:53.208971Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"d531a101-e4d4-44a6-92e6-03230a083078","year":2025},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":78,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:87a0e3275f0c05a991ee7defd723a8d9dcf6799d9ed3125290b337371212ffa0","observation_id":"1fb5416a-7216-4e53-a161-402da07cac86","resolution":{"observed_at":"2026-05-21T08:39:53.798185Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"faa89716-daf0-4720-8ca1-d705c2531350","year":2020},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":79,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:f752e9e4647749c692e759067506d618f9cb7393a7bdcb75c9a28004bbdd79a2","observation_id":"6039143f-e61f-4a55-812f-c917ec5e296c","resolution":{"observed_at":"2026-05-21T08:39:53.805992Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2508.08448","last_updated":"2025-08-11T20:11:43Z","snapshot_observed_at":"2026-08-05T21:32:53.021643Z","submitted_at":"2025-08-11T20:11:43Z","title":"Towards Efficient and Practical GPU Multitasking in the Era of LLM","version":1},"cited_work":{"arxiv_id":"2508.08448","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2508.08448","snapshot_observed_at":"2026-06-28T17:02:24.558075Z","title":"Towards efficient and practical gpu multitasking in the era of llm.arXiv preprint arXiv:2508.08448, 2025","venue":null,"work_id":"5c3a90c4-a851-492a-883c-24bffb0afb63","year":2025},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":80,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"cited_paper":"/paper/2508.08448","citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:55a84e6bb362e0752181555b5d2343abd96cc237db27ac29a47072594c2393b1","observation_id":"f861fc33-b81a-48f0-ac73-82ebe9c41a0d","resolution":{"observed_at":"2026-05-21T08:39:53.318914Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-05T12:51:04.391284Z","title":null,"venue":null,"work_id":"fb58cc3c-673e-425c-a8dd-71dbaea7d5b0","year":null},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":81,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:61b096fd250d1a00ba3a417ac304cf4c5d41fa1378247aa99c927f1124f7c9a9","observation_id":"225e304b-7a05-45c6-a4bd-c0a475a35641","resolution":{"observed_at":"2026-05-21T08:39:53.813203Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.09388","last_updated":"2025-05-14T13:41:34Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-05-14T13:41:34Z","title":"Qwen3 Technical Report","version":1},"cited_work":{"arxiv_id":"2505.09388","doi":"10.1016/j.aiopen.2022.12","metadata_source":"pith","pith_arxiv_id":"2505.09388","snapshot_observed_at":"2026-07-11T11:50:26.030339Z","title":"Qwen3 Technical Report","venue":"cs.CL","work_id":"25a4e30c-1232-48e7-9925-02fa12ba7c9e","year":2025},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":82,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"cited_paper":"/paper/2505.09388","citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:78e58d3e1b15cb91ab312d2b657738019c9ce4107a9b79ef342e18e9be31b9d1","observation_id":"3b750505-78dd-45e4-bef0-e2c9a67e22c0","resolution":{"observed_at":"2026-05-21T08:39:53.311663Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.00311","last_updated":"2025-02-01T04:18:28Z","snapshot_observed_at":"2026-07-06T20:29:26.524162Z","submitted_at":"2025-02-01T04:18:28Z","title":"Sparse Gradient Compression for Fine-Tuning Large Language Models","version":1},"cited_work":{"arxiv_id":"2502.00311","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2502.00311","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Yang, Mohammad Mohammadi Amiri, Tejaswini Pedap- ati, Subhajit Chaudhury, and Pin-Yu Chen","venue":null,"work_id":"4de0301e-9e1b-44e3-9029-d066edf4670b","year":2025},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":83,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"cited_paper":"/paper/2502.00311","citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:9b0b873cd663d0727f9bef1739951f36140aef35ade360a93833e75c535d539d","observation_id":"189fef69-51a0-4caf-9c2d-a20fc04556ed","resolution":{"observed_at":"2026-05-21T08:39:53.216763Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.21798","last_updated":"2025-05-21T17:21:45Z","snapshot_observed_at":"2026-07-06T21:17:09.837496Z","submitted_at":"2025-04-30T16:56:06Z","title":"SWE-smith: Scaling Data for Software Engineering Agents","version":2},"cited_work":{"arxiv_id":"2504.21798","doi":"10.48550/arxiv.2504.21798","metadata_source":"pith","pith_arxiv_id":"2504.21798","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"SWE-smith: Scaling Data for Software Engineering Agents","venue":"cs.SE","work_id":"6a906763-2e4e-4cea-a19c-d7a169c9376b","year":2025},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":84,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"cited_paper":"/paper/2504.21798","citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:4b4e71b8415349f48d264fec1245d4eb0fd7f79ae8fa931537b130e987e77bc0","observation_id":"dcca0704-087c-4188-9651-9ef24a9068d1","resolution":{"observed_at":"2026-05-21T08:39:53.253434Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"5b6350a4-5160-41b3-b8d2-24cf903057fd","year":2025},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":85,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:233723f4d63f84eb3010d27ff12dd26394fdfc54413c48a879349e5c22968a73","observation_id":"ccb98eff-26f8-4e40-9502-b36564ea5905","resolution":{"observed_at":"2026-05-21T08:39:53.853563Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2308.01320","last_updated":"2023-08-02T18:49:57Z","snapshot_observed_at":"2026-07-06T16:01:45.794350Z","submitted_at":"2023-08-02T18:49:57Z","title":"DeepSpeed-Chat: Easy, Fast and Affordable RLHF Training of ChatGPT-like Models at All Scales","version":1},"cited_work":{"arxiv_id":"2308.01320","doi":"10.48550/arxiv.2308.01320","metadata_source":"arxiv_reference","pith_arxiv_id":"2308.01320","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Deepspeed-chat: Easy, fast and affordable rlhf training of chatgpt-like models at all scales","venue":"arXiv (Cornell University)","work_id":"3d244678-a776-4393-8ca9-6b3c9a8e3fde","year":2023},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":86,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"cited_paper":"/paper/2308.01320","citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:492c2fbffebb9c4ef902e3b018829849764012fa2619495a32c6055b125cffa0","observation_id":"81aa4c7c-fa99-4278-bbf4-f768a620a14b","resolution":{"observed_at":"2026-05-21T08:39:53.327060Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2502.09922","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"Lambdas- cale: Enabling fast scaling for serverless large language model inference","venue":null,"work_id":"cf2e378f-e444-4d62-98ef-e22b591f4962","year":2025},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":87,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:ef30e54469a1197401fc518a8fc1f0d0387173246ef979f325959d07f9def662","observation_id":"2ab696df-27d7-417e-b41f-ffbad4d82d3d","resolution":{"observed_at":"2026-05-21T08:39:53.237338Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2503.14476","last_updated":"2025-05-20T01:37:34Z","snapshot_observed_at":"2026-08-02T01:40:54.187278Z","submitted_at":"2025-03-18T17:49:06Z","title":"DAPO: An Open-Source LLM Reinforcement Learning System at Scale","version":2},"cited_work":{"arxiv_id":"2503.14476","doi":"10.48550/arxiv.2503.14476","metadata_source":"pith","pith_arxiv_id":"2503.14476","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"DAPO: An Open-Source LLM Reinforcement Learning System at Scale","venue":"cs.LG","work_id":"64019d00-0b11-4bbd-b173-b46c8fad0157","year":2025},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":88,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"cited_paper":"/paper/2503.14476","citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:3cddaeebd3c7b46bd9129e7284883d9ad63dd6fbc02d38675c265c74dad4fcb9","observation_id":"91064bc4-92e4-42e6-a56a-e8aac892641b","resolution":{"observed_at":"2026-05-21T08:39:53.225109Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-05-24T09:23:06.254602+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-24T09:23:06.254602+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.04021","last_updated":"2026-06-10T23:04:53Z","snapshot_observed_at":"2026-08-07T15:49:22.835029Z","submitted_at":"2025-05-06T23:38:33Z","title":"Prism: Cost-Efficient Multi-LLM Serving via GPU Memory Ballooning","version":3},"cited_work":{"arxiv_id":"2505.04021","doi":null,"metadata_source":"pith","pith_arxiv_id":"2505.04021","snapshot_observed_at":"2026-07-03T17:18:43.284498Z","title":"Prism: Unleashing gpu sharing for cost-efficient multi- llm serving","venue":"cs.DC","work_id":"b06d977c-d2bf-4367-bdeb-cd88d2f6fa4f","year":2025},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":89,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"cited_paper":"/paper/2505.04021","citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:3a58932c4ccea09c6c9d2bc31bf75e32289cc43174089007a660820b57280dd6","observation_id":"aa4aec34-c9de-49e8-9f2e-7c79b40d25c5","resolution":{"observed_at":"2026-06-12T02:08:18.534775Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"cc88c1d8-b15a-4400-8e38-fe0d537c675c","year":2025},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":90,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:a9ba0b81635d4c7bb028871a6017ea9f1ee244f325068ee5d6726d40bcddf42e","observation_id":"a7a534cc-e79c-4edd-b9af-1ea11752ac13","resolution":{"observed_at":"2026-05-21T08:39:53.858036Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2506.09016","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-02T12:06:55.810752Z","title":"Speed-rl: Faster training of reasoning models via online curriculum learning","venue":null,"work_id":"a55d1c0a-b469-4666-b17d-28fb95a2e4e8","year":2025},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":91,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:21a73945f49fdb82e8510a32f707fadd2f83dd623d823afdd68a2eb350f5c0a7","observation_id":"4aeffa84-095f-45cf-8f3c-4cec2db7a293","resolution":{"observed_at":"2026-05-21T08:39:53.341207Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.15720","last_updated":"2025-04-22T09:08:46Z","snapshot_observed_at":"2026-08-07T16:00:18.626052Z","submitted_at":"2025-04-22T09:08:46Z","title":"SeaLLM: Service-Aware and Latency-Optimized Resource Sharing for Large Language Model Inference","version":1},"cited_work":{"arxiv_id":"2504.15720","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2504.15720","snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"ea3138fc-a8be-4f50-8b8c-977e7c708a1a","year":2025},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":92,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"cited_paper":"/paper/2504.15720","citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:57fea7d4c2b24466ab4989eb217b086de85ee4e983308820a12760cae04c107b","observation_id":"e5689416-4417-4ade-b42a-f2ea026da362","resolution":{"observed_at":"2026-05-21T08:39:53.275321Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"3f6f9970-c433-4b7a-88c9-b2a39454a699","year":2022},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":93,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:bede98b2054cb91f8f439c38709d4c527afb086d4307cc46c7abc845958843cb","observation_id":"0ecd1741-07da-4d45-bbd1-0bd2f68aacd5","resolution":{"observed_at":"2026-05-21T08:39:53.825293Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2506.02177","last_updated":"2025-06-02T19:03:00Z","snapshot_observed_at":"2026-08-07T16:09:18.755046Z","submitted_at":"2025-06-02T19:03:00Z","title":"Act Only When It Pays: Efficient Reinforcement Learning for LLM Reasoning via Selective Rollouts","version":1},"cited_work":{"arxiv_id":"2506.02177","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2506.02177","snapshot_observed_at":"2026-07-03T09:47:59.849796Z","title":"Bartoldson, Bhavya Kailkhura, Fan Lai, Jiawei Zhao, and Beidi Chen","venue":null,"work_id":"d2eaa80f-327c-4ce1-925b-baaff8681593","year":2025},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":94,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"cited_paper":"/paper/2506.02177","citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:b7d7b9a7b14f17f4423f49e9dd305a9f9514fe1a8da09ddebb118681e13db0d3","observation_id":"99e4c427-1350-4aaa-bef3-db35fa1bc34b","resolution":{"observed_at":"2026-05-21T08:39:53.200705Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"05eeeb8f-dc2f-40f9-86ea-be0718e13fd4","year":2024},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":95,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:4d687ecbabd76d78ea9c4b8daa75c0c6b4fa7cfa63a0052ceabff357dbe30402","observation_id":"9e3e8fec-70f5-4cbf-ba12-ca1aef2bab4c","resolution":{"observed_at":"2026-05-21T08:39:53.816441Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2504.15930","last_updated":"2025-04-22T14:19:06Z","snapshot_observed_at":"2026-08-07T16:00:12.969923Z","submitted_at":"2025-04-22T14:19:06Z","title":"StreamRL: Scalable, Heterogeneous, and Elastic RL for LLMs with Disaggregated Stream Generation","version":1},"cited_work":{"arxiv_id":"2504.15930","doi":"10.48550/arxiv.2504.15930","metadata_source":"arxiv_reference","pith_arxiv_id":"2504.15930","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Streamrl: Scalable, heterogeneous, and elastic rl for llms with disaggregated stream generation","venue":"ArXiv.org","work_id":"63df3644-702d-4ae6-8e4a-954a41887545","year":2025},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":96,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"cited_paper":"/paper/2504.15930","citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:4c81df4fcee1f714d9e9de32ecd89dc84f24317c6e86964477712b4f2972f5ab","observation_id":"3a096d9f-ad4a-47de-855b-debc9c8ae7bf","resolution":{"observed_at":"2026-05-21T08:39:53.330649Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-07-13T08:50:16.320291+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-13T08:50:16.320291+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":null,"venue":null,"work_id":"22dc6622-323e-44d9-815f-cab89c7d44ed","year":null},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":97,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:6d35fb8c1aef9413fa4855a38c423b73cb97a3caa27628eea3efc2e5db7facec","observation_id":"daf46e83-8827-46b9-bd25-d3fb79fb07e7","resolution":{"observed_at":"2026-05-21T08:39:53.788222Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-06-05T21:23:00.469572Z","title":"In22nd USENIX Symposium on Networked Systems Design and Implementation (NSDI 25)","venue":null,"work_id":"e9322f36-1cc2-483a-8431-72d1d2a51a7d","year":null},"citing_paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL","version":2},"reference_index":98,"source":"pdf_text","source_observed_at":"2026-05-21T08:39:31.911497Z"},"links":{"citing_paper":"/paper/2605.06534"},"observation_digest":"sha256:3027a8f14893d620a1e449f9abd5f6ef2c17e6e623d64ce71b530af648c3c62c","observation_id":"2a964a8a-6faf-4232-bbe3-c5a2f2162b95","resolution":{"observed_at":"2026-05-21T08:39:53.783285Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2605.06534","last_updated":"2026-05-20T11:37:51Z","latest_version":2,"primary_category":"cs.DC","snapshot_observed_at":"2026-08-03T06:43:16.351179Z","submitted_at":"2026-05-07T16:33:40Z","title":"ROSE: Rollout On Serving GPUs via Cooperative Elasticity for Agentic RL"},"reference_resolution":{"displayed":96,"state_counts":{"malformed_identifier":1,"metadata_mismatch":4,"parse_uncertain":0,"unresolved":39,"verified_exact":45,"verified_fuzzy":7},"total_outbound_references":96},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"thesis":"As of 7 August 2026, this Paper Citation Record lists 96 of 96 outbound references and 2 inbound Pith citation observations for arXiv:2605.06534."}