{"as_of":"2026-08-08T03:37:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:01c60d13df4be373272c10738f140152888a8052b8446e2a4e94060b6878cc8b","coverage":[{"denominator":31,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":31,"source":"paper_references, paper_reference_links","source_observed_at":"2026-07-14T06:25:23.264527Z","state":"measured"},{"denominator":31,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":31,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-07T06:34:17.273281+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2607.11185/citation-record","integrity":"/paper/2607.11185/integrity","json":"/paper/2607.11185/citation-record.json","paper":"/paper/2607.11185"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2504.00906","last_updated":"2025-04-01T15:40:27Z","snapshot_observed_at":"2026-07-06T21:02:28.100677Z","submitted_at":"2025-04-01T15:40:27Z","title":"Agent S2: A Compositional Generalist-Specialist Framework for Computer Use Agents","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2504.00906","snapshot_observed_at":"2026-07-14T06:25:23.264527Z","title":"Agent s2: A compositional generalist-specialist framework for computer use agents.arXiv preprint arXiv:2504.00906,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.11185","last_updated":"2026-07-13T07:32:35Z","snapshot_observed_at":"2026-07-16T23:19:14.868685Z","submitted_at":"2026-07-13T07:32:35Z","title":"SCALECUA: Scaling Computer Use Agents with Verifiable Task Synthesis and Efficient Online RL","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-07-14T06:25:23.264527Z"},"links":{"cited_paper":"/paper/2504.00906","citing_paper":"/paper/2607.11185"},"observation_digest":"sha256:1eaa14c08f299cf2669d12e3abcd635654d7758f0a66b6249d54c4ba666b600d","observation_id":"0b86aae1-f8bf-4270-860a-14b421dadb7a","resolution":{"observed_at":"2026-07-14T06:25:23.264527Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2511.21631","last_updated":"2025-11-27T12:16:54Z","snapshot_observed_at":"2026-07-06T22:37:03.716474Z","submitted_at":"2025-11-26T17:59:08Z","title":"Qwen3-VL Technical Report","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2511.21631","snapshot_observed_at":"2026-07-14T06:25:23.264527Z","title":"Qwen3-vl technical report.arXiv preprint arXiv:2511.21631,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.11185","last_updated":"2026-07-13T07:32:35Z","snapshot_observed_at":"2026-07-16T23:19:14.868685Z","submitted_at":"2026-07-13T07:32:35Z","title":"SCALECUA: Scaling Computer Use Agents with Verifiable Task Synthesis and Efficient Online RL","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-07-14T06:25:23.264527Z"},"links":{"cited_paper":"/paper/2511.21631","citing_paper":"/paper/2607.11185"},"observation_digest":"sha256:cfd392ee26f89cef2f5b3307b8cb6b189e48d469352c2cce863bdcafabdc8513","observation_id":"4db3b6de-c2c8-4d0c-93ae-b18c8da5366f","resolution":{"observed_at":"2026-07-14T06:25:23.264527Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.08264","last_updated":"2024-09-13T20:17:13Z","snapshot_observed_at":"2026-08-05T14:11:56.506073Z","submitted_at":"2024-09-12T17:56:43Z","title":"Windows Agent Arena: Evaluating Multi-Modal OS Agents at Scale","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.08264","snapshot_observed_at":"2026-07-14T06:25:23.264527Z","title":"Windows agent arena: Evaluating multi-modal os agents at scale.arXiv preprint arXiv:2409.08264,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.11185","last_updated":"2026-07-13T07:32:35Z","snapshot_observed_at":"2026-07-16T23:19:14.868685Z","submitted_at":"2026-07-13T07:32:35Z","title":"SCALECUA: Scaling Computer Use Agents with Verifiable Task Synthesis and Efficient Online RL","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-07-14T06:25:23.264527Z"},"links":{"cited_paper":"/paper/2409.08264","citing_paper":"/paper/2607.11185"},"observation_digest":"sha256:c064017600445e253e5d35e5074b1acc06982a35015bab449c68bce559ddab1e","observation_id":"64a4e43f-3499-4e10-bb6b-12f54f590a3c","resolution":{"observed_at":"2026-07-14T06:25:23.264527Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-14T06:25:23.264527Z","title":"Gui-genesis: Automated synthesis of efficient environments with verifiable rewards for gui agent post-training.arXiv preprint arXiv:2602.14093,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.11185","last_updated":"2026-07-13T07:32:35Z","snapshot_observed_at":"2026-07-16T23:19:14.868685Z","submitted_at":"2026-07-13T07:32:35Z","title":"SCALECUA: Scaling Computer Use Agents with Verifiable Task Synthesis and Efficient Online RL","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-07-14T06:25:23.264527Z"},"links":{"citing_paper":"/paper/2607.11185"},"observation_digest":"sha256:6e1241889e8ad7ff86d2e8117e99a3f35f4f6d72c6854bf822701d75648e202b","observation_id":"4f1dacf1-6acf-45ba-96f4-b2408baecb7b","resolution":{"observed_at":"2026-07-14T06:25:23.264527Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2507.19849","last_updated":"2025-07-26T07:53:11Z","snapshot_observed_at":"2026-07-06T22:03:15.296567Z","submitted_at":"2025-07-26T07:53:11Z","title":"Agentic Reinforced Policy Optimization","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2507.19849","snapshot_observed_at":"2026-07-14T06:25:23.264527Z","title":"Agentic reinforced policy optimization.arXiv preprint arXiv:2507.19849,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.11185","last_updated":"2026-07-13T07:32:35Z","snapshot_observed_at":"2026-07-16T23:19:14.868685Z","submitted_at":"2026-07-13T07:32:35Z","title":"SCALECUA: Scaling Computer Use Agents with Verifiable Task Synthesis and Efficient Online RL","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-07-14T06:25:23.264527Z"},"links":{"cited_paper":"/paper/2507.19849","citing_paper":"/paper/2607.11185"},"observation_digest":"sha256:cc20502a7103420e584edfe7aea6fff5836ed880be810fd1f031bc5b9fa5a965","observation_id":"a7fe5b11-3441-4c07-b2bc-5d76ee54b478","resolution":{"observed_at":"2026-07-14T06:25:23.264527Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.05243","last_updated":"2025-06-17T15:06:02Z","snapshot_observed_at":"2026-07-06T19:29:09.714725Z","submitted_at":"2024-10-07T17:47:50Z","title":"Navigating the Digital World as Humans Do: Universal Visual Grounding for GUI Agents","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.05243","snapshot_observed_at":"2026-07-14T06:25:23.264527Z","title":"Navigating the digital world as humans do: Universal visual grounding for gui agents","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.11185","last_updated":"2026-07-13T07:32:35Z","snapshot_observed_at":"2026-07-16T23:19:14.868685Z","submitted_at":"2026-07-13T07:32:35Z","title":"SCALECUA: Scaling Computer Use Agents with Verifiable Task Synthesis and Efficient Online RL","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-07-14T06:25:23.264527Z"},"links":{"cited_paper":"/paper/2410.05243","citing_paper":"/paper/2607.11185"},"observation_digest":"sha256:f36794ad9c5924c8353423939431846d0a6a19e2fb5ed919c0db7967c3ec7f78","observation_id":"33f473d4-e0e8-4825-8b6e-65aeb4b337d0","resolution":{"observed_at":"2026-07-14T06:25:23.264527Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.12948","last_updated":"2026-01-04T03:57:36Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-01-22T15:19:35Z","title":"DeepSeek-R1: Incentivizing Reasoning Capability in LLMs via Reinforcement Learning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.12948","snapshot_observed_at":"2026-07-14T06:25:23.264527Z","title":"Deepseek-r1: Incentivizing reasoning capability in llms via reinforcement learning.arXiv preprint arXiv:2501.12948,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.11185","last_updated":"2026-07-13T07:32:35Z","snapshot_observed_at":"2026-07-16T23:19:14.868685Z","submitted_at":"2026-07-13T07:32:35Z","title":"SCALECUA: Scaling Computer Use Agents with Verifiable Task Synthesis and Efficient Online RL","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-07-14T06:25:23.264527Z"},"links":{"cited_paper":"/paper/2501.12948","citing_paper":"/paper/2607.11185"},"observation_digest":"sha256:7e926a300fd9d337e70e6678813fcf6b5ef1c359231003edb98e1fb4d62df2f5","observation_id":"1a3f5382-04c4-449d-98eb-82326fcd579e","resolution":{"observed_at":"2026-07-14T06:25:23.264527Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2507.01006","last_updated":"2026-01-01T13:07:25Z","snapshot_observed_at":"2026-08-03T18:50:30.558321Z","submitted_at":"2025-07-01T17:55:04Z","title":"GLM-4.5V and GLM-4.1V-Thinking: Towards Versatile Multimodal Reasoning with Scalable Reinforcement Learning","version":6},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2507.01006","snapshot_observed_at":"2026-07-14T06:25:23.264527Z","title":"Glm-4.5 v and glm-4.1 v-thinking: Towards versatile multimodal reasoning with scalable reinforcement learning.arXiv preprint arXiv:2507.01006,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.11185","last_updated":"2026-07-13T07:32:35Z","snapshot_observed_at":"2026-07-16T23:19:14.868685Z","submitted_at":"2026-07-13T07:32:35Z","title":"SCALECUA: Scaling Computer Use Agents with Verifiable Task Synthesis and Efficient Online RL","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-07-14T06:25:23.264527Z"},"links":{"cited_paper":"/paper/2507.01006","citing_paper":"/paper/2607.11185"},"observation_digest":"sha256:8d656e17e3be090ef6eb8a3d1cf26e11be549e34c7498aa0c5f981b75bff545e","observation_id":"74475d5f-119d-4f3f-92b8-a00338a67808","resolution":{"observed_at":"2026-07-14T06:25:23.264527Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.13381","last_updated":"2024-06-19T09:23:53Z","snapshot_observed_at":"2026-08-04T03:51:46.195668Z","submitted_at":"2024-06-19T09:23:53Z","title":"CoAct: A Global-Local Hierarchy for Autonomous Agent Collaboration","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.13381","snapshot_observed_at":"2026-07-14T06:25:23.264527Z","title":"Coact: A global-local hierarchy for autonomous agent collaboration.arXiv preprint arXiv:2406.13381,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.11185","last_updated":"2026-07-13T07:32:35Z","snapshot_observed_at":"2026-07-16T23:19:14.868685Z","submitted_at":"2026-07-13T07:32:35Z","title":"SCALECUA: Scaling Computer Use Agents with Verifiable Task Synthesis and Efficient Online RL","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-07-14T06:25:23.264527Z"},"links":{"cited_paper":"/paper/2406.13381","citing_paper":"/paper/2607.11185"},"observation_digest":"sha256:b5e0e2b5ca4403ed435ea91bb7e18ea6f44c447fd85dfff1451f9d649da37d3f","observation_id":"239cc74f-50ec-42de-85ec-e2362093f294","resolution":{"observed_at":"2026-07-14T06:25:23.264527Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-14T06:25:23.264527Z","title":"Androidgen: Building an android language agent under data scarcity","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.11185","last_updated":"2026-07-13T07:32:35Z","snapshot_observed_at":"2026-07-16T23:19:14.868685Z","submitted_at":"2026-07-13T07:32:35Z","title":"SCALECUA: Scaling Computer Use Agents with Verifiable Task Synthesis and Efficient Online RL","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-07-14T06:25:23.264527Z"},"links":{"citing_paper":"/paper/2607.11185"},"observation_digest":"sha256:7e3bc542b4df7d74a394568c92c04f6e2a212879ac5311d5ad2515e767ff4e27","observation_id":"36f00fbb-47fe-489b-8ba3-734bb9eac714","resolution":{"observed_at":"2026-07-14T06:25:23.264527Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2408.06327","last_updated":"2024-08-12T17:44:17Z","snapshot_observed_at":"2026-08-07T04:07:56.293097Z","submitted_at":"2024-08-12T17:44:17Z","title":"VisualAgentBench: Towards Large Multimodal Models as Visual Foundation Agents","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2408.06327","snapshot_observed_at":"2026-07-14T06:25:23.264527Z","title":"Visualagentbench: Towards large multimodal models as visual foundation agents.arXiv preprint arXiv:2408.06327,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.11185","last_updated":"2026-07-13T07:32:35Z","snapshot_observed_at":"2026-07-16T23:19:14.868685Z","submitted_at":"2026-07-13T07:32:35Z","title":"SCALECUA: Scaling Computer Use Agents with Verifiable Task Synthesis and Efficient Online RL","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-07-14T06:25:23.264527Z"},"links":{"cited_paper":"/paper/2408.06327","citing_paper":"/paper/2607.11185"},"observation_digest":"sha256:0a7495e75abbd05928271c70d4a22f15c634996e96faeee014c37c61b1f0104d","observation_id":"01ac9fef-c562-4224-a789-9452760bec0e","resolution":{"observed_at":"2026-07-14T06:25:23.264527Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2411.02337","last_updated":"2025-01-27T11:56:15Z","snapshot_observed_at":"2026-07-06T19:45:00.034364Z","submitted_at":"2024-11-04T17:59:58Z","title":"WebRL: Training LLM Web Agents via Self-Evolving Online Curriculum Reinforcement Learning","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.02337","snapshot_observed_at":"2026-07-14T06:25:23.264527Z","title":"GPT-5.4 system card.https://openai.com/index/introducing-gpt-5-4/, 2025a","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.11185","last_updated":"2026-07-13T07:32:35Z","snapshot_observed_at":"2026-07-16T23:19:14.868685Z","submitted_at":"2026-07-13T07:32:35Z","title":"SCALECUA: Scaling Computer Use Agents with Verifiable Task Synthesis and Efficient Online RL","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-07-14T06:25:23.264527Z"},"links":{"cited_paper":"/paper/2411.02337","citing_paper":"/paper/2607.11185"},"observation_digest":"sha256:dd422eeeb6ca6b71f75da8262eca73f968d3f70d837417451e9be4a0f9cd4428","observation_id":"3ea66aef-dc4e-4898-8121-a7bf36a89eb6","resolution":{"observed_at":"2026-07-14T06:25:23.264527Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.12326","last_updated":"2025-01-21T17:48:10Z","snapshot_observed_at":"2026-07-06T20:23:58.426780Z","submitted_at":"2025-01-21T17:48:10Z","title":"UI-TARS: Pioneering Automated GUI Interaction with Native Agents","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.12326","snapshot_observed_at":"2026-07-14T06:25:23.264527Z","title":"Ui-tars: Pioneering automated gui interaction with native agents.arXiv preprint arXiv:2501.12326,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.11185","last_updated":"2026-07-13T07:32:35Z","snapshot_observed_at":"2026-07-16T23:19:14.868685Z","submitted_at":"2026-07-13T07:32:35Z","title":"SCALECUA: Scaling Computer Use Agents with Verifiable Task Synthesis and Efficient Online RL","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-07-14T06:25:23.264527Z"},"links":{"cited_paper":"/paper/2501.12326","citing_paper":"/paper/2607.11185"},"observation_digest":"sha256:5bef3fcd3e9336dd953356a7d0e3e5fcded0b75b9202cbf54c03516a7fecedbd","observation_id":"fb084d7f-f1d3-4193-812e-0e02e328cb22","resolution":{"observed_at":"2026-07-14T06:25:23.264527Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2603.20633","last_updated":"2026-04-17T08:57:17Z","snapshot_observed_at":"2026-07-06T22:49:57.103822Z","submitted_at":"2026-03-21T04:03:45Z","title":"Seed1.8 Model Card: Towards Generalized Real-World Agency","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2603.20633","snapshot_observed_at":"2026-07-14T06:25:23.264527Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.11185","last_updated":"2026-07-13T07:32:35Z","snapshot_observed_at":"2026-07-16T23:19:14.868685Z","submitted_at":"2026-07-13T07:32:35Z","title":"SCALECUA: Scaling Computer Use Agents with Verifiable Task Synthesis and Efficient Online RL","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-07-14T06:25:23.264527Z"},"links":{"cited_paper":"/paper/2603.20633","citing_paper":"/paper/2607.11185"},"observation_digest":"sha256:98670a4f47b0db8745300dc2d75cb40c96f93421113959b43ea72e4890321c56","observation_id":"6fb7a1f6-f563-46ed-a0cc-59fa0b91d7f3","resolution":{"observed_at":"2026-07-14T06:25:23.264527Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.03300","last_updated":"2024-04-27T15:25:53Z","snapshot_observed_at":"2026-08-06T14:58:42.911363Z","submitted_at":"2024-02-05T18:55:32Z","title":"DeepSeekMath: Pushing the Limits of Mathematical Reasoning in Open Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.03300","snapshot_observed_at":"2026-07-14T06:25:23.264527Z","title":"Deepseekmath: Pushing the limits of mathemat- ical reasoning in open language models.arXiv preprint arXiv:2402.03300,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.11185","last_updated":"2026-07-13T07:32:35Z","snapshot_observed_at":"2026-07-16T23:19:14.868685Z","submitted_at":"2026-07-13T07:32:35Z","title":"SCALECUA: Scaling Computer Use Agents with Verifiable Task Synthesis and Efficient Online RL","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-07-14T06:25:23.264527Z"},"links":{"cited_paper":"/paper/2402.03300","citing_paper":"/paper/2607.11185"},"observation_digest":"sha256:25b22617b41748fa969387c4e6e8f08c2e31b318f350732508f99dcb07032e90","observation_id":"502dc265-67b8-467e-bd31-97eb0af9def3","resolution":{"observed_at":"2026-07-14T06:25:23.264527Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2507.05720","last_updated":"2025-07-08T07:07:53Z","snapshot_observed_at":"2026-08-07T22:33:02.019770Z","submitted_at":"2025-07-08T07:07:53Z","title":"MobileGUI-RL: Advancing Mobile GUI Agent through Reinforcement Learning in Online Environment","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2507.05720","snapshot_observed_at":"2026-07-14T06:25:23.264527Z","title":"Mobilegui-rl: Advancing mobile gui agent through reinforcement learning in online environment.arXiv preprint arXiv:2507.05720,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.11185","last_updated":"2026-07-13T07:32:35Z","snapshot_observed_at":"2026-07-16T23:19:14.868685Z","submitted_at":"2026-07-13T07:32:35Z","title":"SCALECUA: Scaling Computer Use Agents with Verifiable Task Synthesis and Efficient Online RL","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-07-14T06:25:23.264527Z"},"links":{"cited_paper":"/paper/2507.05720","citing_paper":"/paper/2607.11185"},"observation_digest":"sha256:db819596c5c1f751f0928223b07dd03f1eb38ceabcd04011a9956b3edd1456db","observation_id":"f9893b93-b9d0-42f5-8f8a-aa23860c5100","resolution":{"observed_at":"2026-07-14T06:25:23.264527Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1909.08053","last_updated":"2020-03-13T23:45:18Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2019-09-17T19:42:54Z","title":"Megatron-LM: Training Multi-Billion Parameter Language Models Using Model Parallelism","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1909.08053","snapshot_observed_at":"2026-07-14T06:25:23.264527Z","title":"Megatron-lm: Training multi-billion parameter language models using model parallelism","venue":null,"work_id":null,"year":1909},"citing_paper":{"arxiv_id":"2607.11185","last_updated":"2026-07-13T07:32:35Z","snapshot_observed_at":"2026-07-16T23:19:14.868685Z","submitted_at":"2026-07-13T07:32:35Z","title":"SCALECUA: Scaling Computer Use Agents with Verifiable Task Synthesis and Efficient Online RL","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-07-14T06:25:23.264527Z"},"links":{"cited_paper":"/paper/1909.08053","citing_paper":"/paper/2607.11185"},"observation_digest":"sha256:09d3f148d78ec7bfe28cd6a816bd348a33695556c7be4fbe59f1762551db7e6d","observation_id":"1ad2afee-0ccc-4c4c-b390-d2c729b00246","resolution":{"observed_at":"2026-07-14T06:25:23.264527Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2505.19897","last_updated":"2026-04-20T16:29:04Z","snapshot_observed_at":"2026-07-06T21:30:37.657153Z","submitted_at":"2025-05-26T12:27:27Z","title":"ScienceBoard: Evaluating Multimodal Autonomous Agents in Realistic Scientific Workflows","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2505.19897","snapshot_observed_at":"2026-07-14T06:25:23.264527Z","title":"Os-genesis: Automating gui agent trajec- tory construction via reverse task synthesis","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.11185","last_updated":"2026-07-13T07:32:35Z","snapshot_observed_at":"2026-07-16T23:19:14.868685Z","submitted_at":"2026-07-13T07:32:35Z","title":"SCALECUA: Scaling Computer Use Agents with Verifiable Task Synthesis and Efficient Online RL","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-07-14T06:25:23.264527Z"},"links":{"cited_paper":"/paper/2505.19897","citing_paper":"/paper/2607.11185"},"observation_digest":"sha256:736625a59714088c3546fa1608e846f9bf6c90233f9f5c709a4d584f464f6392","observation_id":"4f051194-9f2c-4999-84e5-0e16fcbcb860","resolution":{"observed_at":"2026-07-14T06:25:23.264527Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2509.02544","last_updated":"2025-09-05T14:59:27Z","snapshot_observed_at":"2026-08-05T09:41:26.544360Z","submitted_at":"2025-09-02T17:44:45Z","title":"UI-TARS-2 Technical Report: Advancing GUI Agent with Multi-Turn Reinforcement Learning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2509.02544","snapshot_observed_at":"2026-07-14T06:25:23.264527Z","title":"Ui-tars-2 technical report: Advancing gui agent with multi-turn reinforcement learning.arXiv preprint arXiv:2509.02544, 2025a","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.11185","last_updated":"2026-07-13T07:32:35Z","snapshot_observed_at":"2026-07-16T23:19:14.868685Z","submitted_at":"2026-07-13T07:32:35Z","title":"SCALECUA: Scaling Computer Use Agents with Verifiable Task Synthesis and Efficient Online RL","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-07-14T06:25:23.264527Z"},"links":{"cited_paper":"/paper/2509.02544","citing_paper":"/paper/2607.11185"},"observation_digest":"sha256:fb27e3694fb03b00e9568a8e02b449f9046322c11de6474a6916bb0e01e0f18c","observation_id":"b74c6ba6-0d27-49b2-b84d-a9b00bd6fbf5","resolution":{"observed_at":"2026-07-14T06:25:23.264527Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-14T06:25:23.264527Z","title":"Agentsynth: Scalable task generation for generalist computer-use agents.arXiv preprint arXiv:2506.14205, 2025a","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.11185","last_updated":"2026-07-13T07:32:35Z","snapshot_observed_at":"2026-07-16T23:19:14.868685Z","submitted_at":"2026-07-13T07:32:35Z","title":"SCALECUA: Scaling Computer Use Agents with Verifiable Task Synthesis and Efficient Online RL","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-07-14T06:25:23.264527Z"},"links":{"citing_paper":"/paper/2607.11185"},"observation_digest":"sha256:a855a22677c8f4dc78cd2d27c22f97b73da6711aff2109352af129403492944b","observation_id":"94ed22cc-330f-412d-bdaa-bb9e2358b852","resolution":{"observed_at":"2026-07-14T06:25:23.264527Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-14T06:25:23.264527Z","title":"Scaling computer-use grounding via user interface decomposition and synthesis.arXiv preprint arXiv:2505.13227, 2025b","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.11185","last_updated":"2026-07-13T07:32:35Z","snapshot_observed_at":"2026-07-16T23:19:14.868685Z","submitted_at":"2026-07-13T07:32:35Z","title":"SCALECUA: Scaling Computer Use Agents with Verifiable Task Synthesis and Efficient Online RL","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-07-14T06:25:23.264527Z"},"links":{"citing_paper":"/paper/2607.11185"},"observation_digest":"sha256:71d69494906319c31a354ee226e33904f62ace086606644eac94e3fdde6c991d","observation_id":"95e88f73-9f88-4eae-83ec-d3a228595f0d","resolution":{"observed_at":"2026-07-14T06:25:23.264527Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-14T06:25:23.264527Z","title":"Mobilerl: Online agentic reinforcement learning for mobile gui agents.arXiv preprint arXiv:2509.18119,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.11185","last_updated":"2026-07-13T07:32:35Z","snapshot_observed_at":"2026-07-16T23:19:14.868685Z","submitted_at":"2026-07-13T07:32:35Z","title":"SCALECUA: Scaling Computer Use Agents with Verifiable Task Synthesis and Efficient Online RL","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-07-14T06:25:23.264527Z"},"links":{"citing_paper":"/paper/2607.11185"},"observation_digest":"sha256:d06aae9ad8766456e29b1d0bdc87b904b670d32ed8d04622b8e7580ced1ad839","observation_id":"d5546386-87e0-42bb-accb-77f2d703c242","resolution":{"observed_at":"2026-07-14T06:25:23.264527Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2412.09605","last_updated":"2025-03-03T18:59:36Z","snapshot_observed_at":"2026-08-06T04:29:42.591501Z","submitted_at":"2024-12-12T18:59:27Z","title":"AgentTrek: Agent Trajectory Synthesis via Guiding Replay with Web Tutorials","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.09605","snapshot_observed_at":"2026-07-14T06:25:23.264527Z","title":"Agenttrek: Agent trajectory synthesis via guiding replay with web tutorials.arXiv preprint arXiv:2412.09605,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.11185","last_updated":"2026-07-13T07:32:35Z","snapshot_observed_at":"2026-07-16T23:19:14.868685Z","submitted_at":"2026-07-13T07:32:35Z","title":"SCALECUA: Scaling Computer Use Agents with Verifiable Task Synthesis and Efficient Online RL","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-07-14T06:25:23.264527Z"},"links":{"cited_paper":"/paper/2412.09605","citing_paper":"/paper/2607.11185"},"observation_digest":"sha256:4cc65b4d83571f1bd5ceb8273cfb264ce81a2c6b02d61d0421d61ddc245333d5","observation_id":"4e4b3f1c-7045-4b0a-bfa0-17a11abee2a9","resolution":{"observed_at":"2026-07-14T06:25:23.264527Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-14T06:25:23.264527Z","title":"Evocua: Evolving computer use agents via learning from scalable synthetic experience.arXiv preprint arXiv:2601.15876,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.11185","last_updated":"2026-07-13T07:32:35Z","snapshot_observed_at":"2026-07-16T23:19:14.868685Z","submitted_at":"2026-07-13T07:32:35Z","title":"SCALECUA: Scaling Computer Use Agents with Verifiable Task Synthesis and Efficient Online RL","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-07-14T06:25:23.264527Z"},"links":{"citing_paper":"/paper/2607.11185"},"observation_digest":"sha256:cf65726a487839bb9edf6477eaccf27e3873e06ffba9a5117aefa1858336cbbf","observation_id":"6202e2ba-0fee-4776-b1c7-df3dc0e56d20","resolution":{"observed_at":"2026-07-14T06:25:23.264527Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.12122","last_updated":"2024-09-18T16:45:37Z","snapshot_observed_at":"2026-07-06T19:17:41.512834Z","submitted_at":"2024-09-18T16:45:37Z","title":"Qwen2.5-Math Technical Report: Toward Mathematical Expert Model via Self-Improvement","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.12122","snapshot_observed_at":"2026-07-14T06:25:23.264527Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.11185","last_updated":"2026-07-13T07:32:35Z","snapshot_observed_at":"2026-07-16T23:19:14.868685Z","submitted_at":"2026-07-13T07:32:35Z","title":"SCALECUA: Scaling Computer Use Agents with Verifiable Task Synthesis and Efficient Online RL","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-07-14T06:25:23.264527Z"},"links":{"cited_paper":"/paper/2409.12122","citing_paper":"/paper/2607.11185"},"observation_digest":"sha256:b8f6c90bca979334b18fb2f53e1c40b5a5747591f41e603ff46096d934641af6","observation_id":"b4c8fc65-81ae-4a55-ba50-6873b7dba976","resolution":{"observed_at":"2026-07-14T06:25:23.264527Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2505.09388","last_updated":"2025-05-14T13:41:34Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-05-14T13:41:34Z","title":"Qwen3 Technical Report","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2505.09388","snapshot_observed_at":"2026-07-14T06:25:23.264527Z","title":"Qwen3 technical report.arXiv preprint arXiv:2505.09388, 2025a","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.11185","last_updated":"2026-07-13T07:32:35Z","snapshot_observed_at":"2026-07-16T23:19:14.868685Z","submitted_at":"2026-07-13T07:32:35Z","title":"SCALECUA: Scaling Computer Use Agents with Verifiable Task Synthesis and Efficient Online RL","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-07-14T06:25:23.264527Z"},"links":{"cited_paper":"/paper/2505.09388","citing_paper":"/paper/2607.11185"},"observation_digest":"sha256:a0209529f86a1c90699ddec416d5ffee0dfc81721ebacdbf57f2ef1d0c0dc385","observation_id":"b93b78de-be2a-44df-b0d3-6f5e972cddca","resolution":{"observed_at":"2026-07-14T06:25:23.264527Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2503.14476","last_updated":"2025-05-20T01:37:34Z","snapshot_observed_at":"2026-08-02T01:40:54.187278Z","submitted_at":"2025-03-18T17:49:06Z","title":"DAPO: An Open-Source LLM Reinforcement Learning System at Scale","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2503.14476","snapshot_observed_at":"2026-07-14T06:25:23.264527Z","title":"Dapo: An open-source llm reinforcement learning system at scale.arXiv preprint arXiv:2503.14476,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.11185","last_updated":"2026-07-13T07:32:35Z","snapshot_observed_at":"2026-07-16T23:19:14.868685Z","submitted_at":"2026-07-13T07:32:35Z","title":"SCALECUA: Scaling Computer Use Agents with Verifiable Task Synthesis and Efficient Online RL","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-07-14T06:25:23.264527Z"},"links":{"cited_paper":"/paper/2503.14476","citing_paper":"/paper/2607.11185"},"observation_digest":"sha256:32c4db550a2d64389646e4bedf9242fbc8871d09d04be5de76a17eb761a5d02c","observation_id":"4a19d40a-4490-4ba9-934a-265c9c3a1bff","resolution":{"observed_at":"2026-07-14T06:25:23.264527Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-14T06:25:23.264527Z","title":"Agentrl: Scaling agentic reinforcement learning with a multi-turn, multi-task framework.arXiv preprint arXiv:2510.04206,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.11185","last_updated":"2026-07-13T07:32:35Z","snapshot_observed_at":"2026-07-16T23:19:14.868685Z","submitted_at":"2026-07-13T07:32:35Z","title":"SCALECUA: Scaling Computer Use Agents with Verifiable Task Synthesis and Efficient Online RL","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-07-14T06:25:23.264527Z"},"links":{"citing_paper":"/paper/2607.11185"},"observation_digest":"sha256:2ba37293fc1e3ae3c42b21436cc8bf8501c439587d04095100c30f8cdfbfaea7","observation_id":"3ae1de40-ba82-4c3f-8ce2-69c056ecca52","resolution":{"observed_at":"2026-07-14T06:25:23.264527Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2307.13854","last_updated":"2024-04-16T15:13:18Z","snapshot_observed_at":"2026-08-06T12:47:01.809185Z","submitted_at":"2023-07-25T22:59:32Z","title":"WebArena: A Realistic Web Environment for Building Autonomous Agents","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2307.13854","snapshot_observed_at":"2026-07-14T06:25:23.264527Z","title":"Webarena: A realistic web environment for building autonomous agents.arXiv preprint arXiv:2307.13854,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.11185","last_updated":"2026-07-13T07:32:35Z","snapshot_observed_at":"2026-07-16T23:19:14.868685Z","submitted_at":"2026-07-13T07:32:35Z","title":"SCALECUA: Scaling Computer Use Agents with Verifiable Task Synthesis and Efficient Online RL","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-07-14T06:25:23.264527Z"},"links":{"cited_paper":"/paper/2307.13854","citing_paper":"/paper/2607.11185"},"observation_digest":"sha256:7594338eb310b1358b3cfdf9685168486c2fba83b62ac49a990988b14b4a493c","observation_id":"6695e059-b3a6-4f8a-9f5e-94fc0249b227","resolution":{"observed_at":"2026-07-14T06:25:23.264527Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-14T06:25:23.264527Z","title":null,"venue":null,"work_id":null,"year":2000},"citing_paper":{"arxiv_id":"2607.11185","last_updated":"2026-07-13T07:32:35Z","snapshot_observed_at":"2026-07-16T23:19:14.868685Z","submitted_at":"2026-07-13T07:32:35Z","title":"SCALECUA: Scaling Computer Use Agents with Verifiable Task Synthesis and Efficient Online RL","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-07-14T06:25:23.264527Z"},"links":{"citing_paper":"/paper/2607.11185"},"observation_digest":"sha256:456fb135521c1d5fe692a4080478d75329ea233c046faf789073703bcdc6c555","observation_id":"3282f93e-ae12-4ea7-9bc7-acbf88a40564","resolution":{"observed_at":"2026-07-14T06:25:23.264527Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-14T06:25:23.264527Z","title":"Larger K values produce fewer but longer segments, requiring more rollout workers to keep the training engine fed","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.11185","last_updated":"2026-07-13T07:32:35Z","snapshot_observed_at":"2026-07-16T23:19:14.868685Z","submitted_at":"2026-07-13T07:32:35Z","title":"SCALECUA: Scaling Computer Use Agents with Verifiable Task Synthesis and Efficient Online RL","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-07-14T06:25:23.264527Z"},"links":{"citing_paper":"/paper/2607.11185"},"observation_digest":"sha256:8e4d2c2eb657798871d799dbad6e8899b3b88427f4576e6b13c6daa48429eb75","observation_id":"b21034aa-9e33-4325-985d-262f88f9dfad","resolution":{"observed_at":"2026-07-14T06:25:23.264527Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2607.11185","last_updated":"2026-07-13T07:32:35Z","latest_version":1,"primary_category":"cs.AI","snapshot_observed_at":"2026-07-16T23:19:14.868685Z","submitted_at":"2026-07-13T07:32:35Z","title":"SCALECUA: Scaling Computer Use Agents with Verifiable Task Synthesis and Efficient Online RL"},"reference_resolution":{"displayed":31,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":31,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":31},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"thesis":"As of 8 August 2026, this Paper Citation Record lists 31 of 31 outbound references and 0 inbound Pith citation observations for arXiv:2607.11185."}