{"as_of":"2026-08-12T04:15:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:6d8e1a5c23d6ba9e6f56051ec9dcaf764cb4f05f17428dc8be7fbabb479741c1","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":100,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":100,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-11T06:34:44.6726+00:00","state":"measured"},{"denominator":107,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":100,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-11T14:30:38.008772Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":1,"source":"pith","source_observed_at":"2026-08-05T02:28:24.338817Z","state":"measured"}],"external_citation_measurements":[{"count":8,"observed_at":"2026-08-05T02:28:24.338817Z","source":"pith"}],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-11T14:30:38.008772Z","title":"arXiv preprint arXiv:2501.04519","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2412.11934","last_updated":"2025-06-14T06:06:53Z","snapshot_observed_at":"2026-08-11T14:24:24.392794Z","submitted_at":"2024-12-16T16:20:41Z","title":"Stepwise Reasoning Error Disruption Attack of LLMs","version":5},"reference_index":2025,"source":"pdf_text","source_observed_at":"2026-08-11T14:30:38.008772Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2412.11934"},"observation_digest":"sha256:c1649bb9b789be0460c0d2159bb8abd2786878782524eb8c6405be71869de7f6","observation_id":"3b9c28fe-e1d3-42d1-966e-c279540ba545","resolution":{"observed_at":"2026-08-11T14:30:38.008772Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-10T18:36:54.504256Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2501.11223","last_updated":"2025-06-11T13:19:22Z","snapshot_observed_at":"2026-08-11T15:33:29.567283Z","submitted_at":"2025-01-20T02:16:19Z","title":"Reasoning Language Models: A Blueprint","version":4},"reference_index":57,"source":"pdf_text","source_observed_at":"2026-08-10T18:36:54.504256Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2501.11223"},"observation_digest":"sha256:de64d745b186ee02790c5c2c71484fabd44cee30e94286e1e15e0be4e53da52b","observation_id":"ae3cd6c2-b95c-4168-b2ff-fe81fead2e87","resolution":{"observed_at":"2026-08-10T18:36:54.504256Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-10T00:43:28.905055Z","title":"L., Liu, Y ., Shang, N., Sun, Y ., Zhu, Y ., Yang, F., and Yang, M","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2501.18107","last_updated":"2025-06-07T00:03:08Z","snapshot_observed_at":"2026-08-11T23:20:55.965204Z","submitted_at":"2025-01-30T03:16:44Z","title":"Scaling Inference-Efficient Language Models","version":2},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-10T00:43:28.905055Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2501.18107"},"observation_digest":"sha256:f8ea067252e66b201d248167833bf687a5cb90dbead8acdc3c602c1d99c80c2c","observation_id":"034e27fd-824d-48e9-bf03-ea75755be632","resolution":{"observed_at":"2026-08-10T00:43:28.905055Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":"2501.04519","doi":"10.48550/arxiv.2501.04519","metadata_source":"pith","pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","venue":"cs.CL","work_id":"49792b83-569e-4f5f-ae80-e96cbd3b7a43","year":2025},"citing_paper":{"arxiv_id":"2502.00955","last_updated":"2026-04-24T08:45:39Z","snapshot_observed_at":"2026-07-06T20:29:55.352526Z","submitted_at":"2025-02-02T23:20:16Z","title":"Efficient Multi-Agent System Training with Data Influence-Oriented Tree Search","version":2},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-05-23T04:06:23.521344Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2502.00955"},"observation_digest":"sha256:3e4bd9169408be595e4d7e1dc70c715a8ad645de8950335e3ff08629bbbc0482","observation_id":"586d5cd4-62e1-4a27-bd81-499fb27f53dd","resolution":{"observed_at":"2026-05-23T04:07:30.497780Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-09T13:25:52.006995Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2502.02095","last_updated":"2025-05-20T12:35:46Z","snapshot_observed_at":"2026-08-10T02:46:24.418647Z","submitted_at":"2025-02-04T08:25:17Z","title":"LongDPO: Unlock Better Long-form Generation Abilities for LLMs via Critique-augmented Stepwise Information","version":2},"reference_index":16,"source":"arxiv_source","source_observed_at":"2026-08-09T13:25:52.006995Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2502.02095"},"observation_digest":"sha256:35bf425922402a56fb155885fd5cadbbddbb031376173e2dca785a25d8000a98","observation_id":"1d2a5277-40df-4e3c-a160-c552a7fc7537","resolution":{"observed_at":"2026-08-09T13:25:52.006995Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-09T11:53:29.385817Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2502.02523","last_updated":"2025-02-07T15:02:21Z","snapshot_observed_at":"2026-08-10T20:10:31.948167Z","submitted_at":"2025-02-04T17:45:32Z","title":"Brief analysis of DeepSeek R1 and its implications for Generative AI","version":3},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-09T11:53:29.385817Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2502.02523"},"observation_digest":"sha256:52c79e92c6ed99c05a8791693f7036fb444cfa9936129b340ccaec68e6b1f731","observation_id":"4bf0e9b3-3ecd-4465-98cd-4db3051e59e0","resolution":{"observed_at":"2026-08-09T11:53:29.385817Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-08T23:50:35.741225Z","title":"L., Liu, Y., Shang, N., Sun, Y., Zhu, Y., Yang, F., and Yang, M","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2502.04040","last_updated":"2025-05-30T09:43:42Z","snapshot_observed_at":"2026-08-10T18:19:42.295295Z","submitted_at":"2025-02-06T13:01:44Z","title":"Safety Reasoning with Guidelines","version":2},"reference_index":21,"source":"arxiv_source","source_observed_at":"2026-08-08T23:50:35.741225Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2502.04040"},"observation_digest":"sha256:14bdc7d35eb3815cef4b94a0dc54da1d35ea7ff5bae85a2664048da86224493a","observation_id":"05159511-cc5c-4ca3-aad5-3350ce6dfb06","resolution":{"observed_at":"2026-08-08T23:50:35.741225Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-09T12:16:27.380681Z","title":"L., Liu, Y ., Shang, N., Sun, Y ., Zhu, Y ., Yang, F., and Yang, M","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2502.04350","last_updated":"2025-05-29T00:38:10Z","snapshot_observed_at":"2026-08-09T23:37:02.751429Z","submitted_at":"2025-02-04T15:53:59Z","title":"CodeSteer: Symbolic-Augmented Language Models via Code/Text Guidance","version":2},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-09T12:16:27.380681Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2502.04350"},"observation_digest":"sha256:054d4d6e34f83407455ff26fcfb9b09760842c7ec66b22ecc3cf2e6ce67de3bc","observation_id":"996155ce-2629-4548-a377-bcbfe0093546","resolution":{"observed_at":"2026-08-09T12:16:27.380681Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":"2501.04519","doi":"10.48550/arxiv.2501.04519","metadata_source":"pith","pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","venue":"cs.CL","work_id":"49792b83-569e-4f5f-ae80-e96cbd3b7a43","year":2025},"citing_paper":{"arxiv_id":"2502.04689","last_updated":"2026-04-16T17:03:43Z","snapshot_observed_at":"2026-07-06T20:32:39.321377Z","submitted_at":"2025-02-07T06:30:33Z","title":"Improving Language Models with Intentional Analysis","version":4},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-05-23T03:44:33.625252Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2502.04689"},"observation_digest":"sha256:9f34d3423c79a73a537ec221664df3cb37985676c4229db58e02881954e6f25f","observation_id":"2f97b4a9-2d70-4e09-8a11-3d0d9e19b7cc","resolution":{"observed_at":"2026-05-23T03:45:21.598416Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-08T19:23:26.305225Z","title":"L., Liu, Y ., Shang, N., Sun, Y ., Zhu, Y ., Yang, F., and Yang, M","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2502.05449","last_updated":"2025-06-01T20:14:16Z","snapshot_observed_at":"2026-08-10T20:57:26.493297Z","submitted_at":"2025-02-08T04:39:51Z","title":"Iterative Deepening Sampling as Efficient Test-Time Scaling","version":2},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-08T19:23:26.305225Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2502.05449"},"observation_digest":"sha256:aadd9179e212cfeb916c28bd7ac7493290befd674bb231483a04dd36d42f0160","observation_id":"d9fb27c5-3ba9-4b3c-a5b9-5fc998f8172d","resolution":{"observed_at":"2026-08-08T19:23:26.305225Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-08T18:10:53.373963Z","title":"L., Liu, Y ., Shang, N., Sun, Y ., Zhu, Y ., Yang, F., and Yang, M","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2502.05773","last_updated":"2025-07-24T22:59:42Z","snapshot_observed_at":"2026-08-10T01:49:25.494930Z","submitted_at":"2025-02-09T04:31:30Z","title":"PIPA: Preference Alignment as Prior-Informed Statistical Estimation","version":2},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-08T18:10:53.373963Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2502.05773"},"observation_digest":"sha256:544dc905e3a9b05792570a1c1dd0775501ca2c97df0a3f9989be25273aaeb277","observation_id":"bb67e40f-4cf7-42ae-8179-a6522be88ea6","resolution":{"observed_at":"2026-08-08T18:10:53.373963Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-08T14:40:35.941372Z","title":"rStar-Math : Small llms can master math reasoning with self-evolved deep thinking","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2502.06703","last_updated":"2025-02-10T17:30:23Z","snapshot_observed_at":"2026-08-08T18:40:06.157350Z","submitted_at":"2025-02-10T17:30:23Z","title":"Can 1B LLM Surpass 405B LLM? Rethinking Compute-Optimal Test-Time Scaling","version":1},"reference_index":13,"source":"arxiv_source","source_observed_at":"2026-08-08T14:40:35.941372Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2502.06703"},"observation_digest":"sha256:44980a5b061fcc032c986e021aac1b2b2749de2575991e0cb29d023b41eeed67","observation_id":"7e1abdd1-50fb-48bf-8a14-e2f02454c213","resolution":{"observed_at":"2026-08-08T14:40:35.941372Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-08T14:25:53.340082Z","title":"rstar-math: Small llms can master math reasoning with self-evolved deep thinking","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2502.06773","last_updated":"2025-02-10T18:52:04Z","snapshot_observed_at":"2026-08-10T11:31:53.867683Z","submitted_at":"2025-02-10T18:52:04Z","title":"On the Emergence of Thinking in LLMs I: Searching for the Right Intuition","version":1},"reference_index":19,"source":"arxiv_source","source_observed_at":"2026-08-08T14:25:53.340082Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2502.06773"},"observation_digest":"sha256:e30018c5d53c0c116c3eb830e5f766595700c47665e4c9a3b8d509cab0525cfa","observation_id":"c7e1636c-8b03-4263-91b1-ae5239e0bc6a","resolution":{"observed_at":"2026-08-08T14:25:53.340082Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-08T14:27:51.099046Z","title":"rstar-math: Small llms can master math reasoning with self-evolved deep thinking","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2502.06781","last_updated":"2025-02-10T18:57:29Z","snapshot_observed_at":"2026-08-09T17:45:05.773418Z","submitted_at":"2025-02-10T18:57:29Z","title":"Exploring the Limit of Outcome Reward for Learning Mathematical Reasoning","version":1},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-08T14:27:51.099046Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2502.06781"},"observation_digest":"sha256:22de3efde360b76da046b5fbb97afea0198c8b35ab8044108a044790dc8a6e38","observation_id":"1593a911-7d7b-4aa5-ba85-16fd0f66797b","resolution":{"observed_at":"2026-08-08T14:27:51.099046Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-08T06:00:32.234438Z","title":"L., Liu, Y., Shang, N., Sun, Y., Zhu, Y., Yang, F., and Yang, M","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2502.08235","last_updated":"2025-02-12T09:23:26Z","snapshot_observed_at":"2026-08-10T00:41:29.483613Z","submitted_at":"2025-02-12T09:23:26Z","title":"The Danger of Overthinking: Examining the Reasoning-Action Dilemma in Agentic Tasks","version":1},"reference_index":13,"source":"arxiv_source","source_observed_at":"2026-08-08T06:00:32.234438Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2502.08235"},"observation_digest":"sha256:cb966bf3d9af2272740e5ff65e9ed589e0a63f0d7dded969bfdd424de734c632","observation_id":"0db4016a-38d6-4cf6-ae2a-a510da52520a","resolution":{"observed_at":"2026-08-08T06:00:32.234438Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-07T22:52:54.318244Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking , 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2502.09042","last_updated":"2025-03-27T06:45:15Z","snapshot_observed_at":"2026-08-11T16:14:52.641432Z","submitted_at":"2025-02-13T07:55:54Z","title":"Typhoon T1: An Open Thai Reasoning Model","version":2},"reference_index":22,"source":"arxiv_source","source_observed_at":"2026-08-07T22:52:54.318244Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2502.09042"},"observation_digest":"sha256:c5a0a6b4955ad813b0beae5f2b0808f6c71a9b00b286b985c3f2b5e03db166d6","observation_id":"654b6651-bd04-4583-a1d9-2df3143df84d","resolution":{"observed_at":"2026-08-07T22:52:54.318244Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-07T20:57:42.034406Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2502.09601","last_updated":"2025-02-13T18:52:36Z","snapshot_observed_at":"2026-08-08T03:37:50.177168Z","submitted_at":"2025-02-13T18:52:36Z","title":"CoT-Valve: Length-Compressible Chain-of-Thought Tuning","version":1},"reference_index":13,"source":"arxiv_source","source_observed_at":"2026-08-07T20:57:42.034406Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2502.09601"},"observation_digest":"sha256:24901a674921f844ad3cd13c884512c0af9f4131626f58059072256f666884ed","observation_id":"a43164e7-7775-49b3-b0df-02a5a357d75e","resolution":{"observed_at":"2026-08-07T20:57:42.034406Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":"2501.04519","doi":"10.48550/arxiv.2501.04519","metadata_source":"pith","pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","venue":"cs.CL","work_id":"49792b83-569e-4f5f-ae80-e96cbd3b7a43","year":2025},"citing_paper":{"arxiv_id":"2502.17419","last_updated":"2025-06-25T02:24:46Z","snapshot_observed_at":"2026-08-10T09:42:17.185681Z","submitted_at":"2025-02-24T18:50:52Z","title":"From System 1 to System 2: A Survey of Reasoning Large Language Models","version":6},"reference_index":143,"source":"pdf_text","source_observed_at":"2026-05-13T01:36:23.845366Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2502.17419"},"observation_digest":"sha256:14a14d5c0acb2409b1cdcc6d38e89ba03e0fd2eaf4a8823396c729b0454700eb","observation_id":"67a98d22-7542-4dc7-9e50-1c9897cc54c2","resolution":{"observed_at":"2026-05-17T04:43:48.343981Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":"2501.04519","doi":"10.48550/arxiv.2501.04519","metadata_source":"pith","pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","venue":"cs.CL","work_id":"49792b83-569e-4f5f-ae80-e96cbd3b7a43","year":2025},"citing_paper":{"arxiv_id":"2502.19918","last_updated":"2026-05-07T16:58:35Z","snapshot_observed_at":"2026-08-02T05:40:27.766263Z","submitted_at":"2025-02-27T09:40:13Z","title":"Meta-Reasoner: Dynamic Guidance for Optimized Inference-time Reasoning in Large Language Models","version":6},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-05-23T02:46:20.578680Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2502.19918"},"observation_digest":"sha256:db9f284b3be98f81ef27c66f629012e52107e1c727d969af4a77707b5407ab54","observation_id":"14267516-3810-4312-bc97-390ab80baef5","resolution":{"observed_at":"2026-05-23T02:47:26.385462Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":"2501.04519","doi":"10.48550/arxiv.2501.04519","metadata_source":"pith","pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","venue":"cs.CL","work_id":"49792b83-569e-4f5f-ae80-e96cbd3b7a43","year":2025},"citing_paper":{"arxiv_id":"2503.09567","last_updated":"2025-07-18T15:57:54Z","snapshot_observed_at":"2026-08-08T22:33:20.124926Z","submitted_at":"2025-03-12T17:35:03Z","title":"Towards Reasoning Era: A Survey of Long Chain-of-Thought for Reasoning Large Language Models","version":5},"reference_index":226,"source":"pdf_text","source_observed_at":"2026-05-12T08:40:40.910461Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2503.09567"},"observation_digest":"sha256:27f6ced6da0f1eccff5d4359db55fd954dea9c786227c8e70f86fb22056553db","observation_id":"d9affb95-ebff-4890-bb90-3ffdcfd77928","resolution":{"observed_at":"2026-05-17T04:43:48.343981Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":"2501.04519","doi":"10.48550/arxiv.2501.04519","metadata_source":"pith","pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","venue":"cs.CL","work_id":"49792b83-569e-4f5f-ae80-e96cbd3b7a43","year":2025},"citing_paper":{"arxiv_id":"2504.21318","last_updated":"2025-04-30T05:05:09Z","snapshot_observed_at":"2026-08-08T08:43:53.657438Z","submitted_at":"2025-04-30T05:05:09Z","title":"Phi-4-reasoning Technical Report","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-05-17T03:40:25.706499Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2504.21318"},"observation_digest":"sha256:1354e74b798bea9757d73b2c00deeb6af97195d55e719b6e7e878359f3976fb9","observation_id":"14c552b0-c3f1-4f16-aa6b-e44d33d2395b","resolution":{"observed_at":"2026-05-17T04:43:48.343981Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":"2501.04519","doi":"10.48550/arxiv.2501.04519","metadata_source":"pith","pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","venue":"cs.CL","work_id":"49792b83-569e-4f5f-ae80-e96cbd3b7a43","year":2025},"citing_paper":{"arxiv_id":"2505.11737","last_updated":"2026-04-11T18:08:02Z","snapshot_observed_at":"2026-08-02T07:06:47.956538Z","submitted_at":"2025-05-16T22:47:32Z","title":"TokUR: Token-Level Uncertainty Estimation for Large Language Model Reasoning","version":4},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-05-22T13:58:07.913104Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2505.11737"},"observation_digest":"sha256:3d8890bbc6aab9b6d48ebdbcbc503c26590ca305a9185b2d9b023cf7a418da1d","observation_id":"0ead29b4-f399-46dc-a0f7-0837fd53f05c","resolution":{"observed_at":"2026-05-22T14:01:38.382405Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-07T15:42:21.412901Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.14107","last_updated":"2025-05-29T08:24:00Z","snapshot_observed_at":"2026-08-07T23:55:00.741921Z","submitted_at":"2025-05-20T09:14:53Z","title":"DiagnosisArena: Benchmarking Diagnostic Reasoning for Large Language Models","version":4},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-07T15:42:21.412901Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2505.14107"},"observation_digest":"sha256:9f48c8f36b2bcc4438879c7d253c92ca47df6b8b5e79201f05275413322b33f4","observation_id":"50cd3b33-c295-4122-838b-a37faeff1e13","resolution":{"observed_at":"2026-08-07T15:42:21.412901Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":"2501.04519","doi":"10.48550/arxiv.2501.04519","metadata_source":"pith","pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","venue":"cs.CL","work_id":"49792b83-569e-4f5f-ae80-e96cbd3b7a43","year":2025},"citing_paper":{"arxiv_id":"2505.15692","last_updated":"2026-05-15T09:38:51Z","snapshot_observed_at":"2026-08-02T00:54:12.697668Z","submitted_at":"2025-05-21T16:06:10Z","title":"TemplateRL: Structured Template-Guided Reinforcement Learning for LLM Reasoning","version":5},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-05-22T13:39:17.205077Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2505.15692"},"observation_digest":"sha256:82b8064d5b96a33cd7de38e40ff7f01e5f7b150767bfbfdc5a50aa9827fba036","observation_id":"bcd871aa-cf77-468b-981c-26d38e7d523a","resolution":{"observed_at":"2026-05-22T13:41:36.183624Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-07T15:15:42.134427Z","title":"rstar-math: Small llms can master math reasoning with self-evolved deep thinking.arXiv preprint arXiv:2501.04519, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.15817","last_updated":"2025-06-09T21:22:15Z","snapshot_observed_at":"2026-08-07T23:33:11.410147Z","submitted_at":"2025-05-21T17:59:54Z","title":"Learning to Reason via Mixture-of-Thought for Logical Reasoning","version":2},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-07T15:15:42.134427Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2505.15817"},"observation_digest":"sha256:773a0e10189450a5efed05329f0c3b18de48c7cb51c7caac84b518fa8c905e38","observation_id":"0b1ccf3e-7388-41e7-bc96-c6e6f01be5f8","resolution":{"observed_at":"2026-08-07T15:15:42.134427Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-07T14:46:17.111371Z","title":"rstar-math: Small llms can master math reasoning with self-evolved deep thinking.arXiv preprint arXiv:2501.04519, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.17697","last_updated":"2025-05-23T10:07:18Z","snapshot_observed_at":"2026-08-07T21:54:30.483014Z","submitted_at":"2025-05-23T10:07:18Z","title":"Activation Control for Efficiently Eliciting Long Chain-of-thought Ability of Language Models","version":1},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-07T14:46:17.111371Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2505.17697"},"observation_digest":"sha256:9962d4b410ea171809dbf153465b7d0bd7b3ee1cf1ef3f81f27a3373a0d53923","observation_id":"8694828c-2bb2-4f8b-acde-9d3ae93b6fb4","resolution":{"observed_at":"2026-08-07T14:46:17.111371Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-07T14:45:16.786796Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.17829","last_updated":"2025-05-23T12:42:50Z","snapshot_observed_at":"2026-08-10T01:27:23.018231Z","submitted_at":"2025-05-23T12:42:50Z","title":"Stepwise Reasoning Checkpoint Analysis: A Test Time Scaling Method to Enhance LLMs' Reasoning","version":1},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-08-07T14:45:16.786796Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2505.17829"},"observation_digest":"sha256:e43e6bc4afdfeff1f6ac710163020341549f452531ae3e8ded06cfe6b62d348b","observation_id":"38b3018e-0145-4e90-a9ac-f5ecea67b504","resolution":{"observed_at":"2026-08-07T14:45:16.786796Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-07T14:42:11.309063Z","title":"rstar-math: Small llms can master math reasoning with self-evolved deep thinking","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.17941","last_updated":"2025-05-23T14:17:56Z","snapshot_observed_at":"2026-08-09T06:05:38.396561Z","submitted_at":"2025-05-23T14:17:56Z","title":"VeriThinker: Learning to Verify Makes Reasoning Model Efficient","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-07T14:42:11.309063Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2505.17941"},"observation_digest":"sha256:fdfbdbe83b8fd33c17221bc5720a1da881bbcdbace1b510258bdf53d7f94e5c0","observation_id":"bedae641-e22c-4a95-97c6-a0b014872ec4","resolution":{"observed_at":"2026-08-07T14:42:11.309063Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-07T14:35:19.843349Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.18405","last_updated":"2025-05-27T04:20:31Z","snapshot_observed_at":"2026-08-09T09:27:17.241700Z","submitted_at":"2025-05-23T22:18:32Z","title":"RaDeR: Reasoning-aware Dense Retrieval Models","version":2},"reference_index":13,"source":"arxiv_source","source_observed_at":"2026-08-07T14:35:19.843349Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2505.18405"},"observation_digest":"sha256:c64c30f350d3793c9cfed70017cb654fe1fbcfbf3060d255d64394d2512aa782","observation_id":"6ac5c1b8-bbad-4caf-b8b5-86409613798d","resolution":{"observed_at":"2026-08-07T14:35:19.843349Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-07T14:23:07.990118Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.19126","last_updated":"2025-05-25T12:47:39Z","snapshot_observed_at":"2026-08-09T11:06:54.014327Z","submitted_at":"2025-05-25T12:47:39Z","title":"MMATH: A Multilingual Benchmark for Mathematical Reasoning","version":1},"reference_index":10,"source":"arxiv_source","source_observed_at":"2026-08-07T14:23:07.990118Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2505.19126"},"observation_digest":"sha256:340eaebe8a8aea06a5c0d1d3f3acb7bcc12575a458c03a59ca9a58c112927487","observation_id":"55471775-4c12-4441-a108-254b5ab1853e","resolution":{"observed_at":"2026-08-07T14:23:07.990118Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-07T14:14:21.569416Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.19634","last_updated":"2025-09-12T01:41:20Z","snapshot_observed_at":"2026-08-10T02:29:59.608724Z","submitted_at":"2025-05-26T07:51:30Z","title":"Faster and Better LLMs via Latency-Aware Test-Time Scaling","version":4},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-08-07T14:14:21.569416Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2505.19634"},"observation_digest":"sha256:0489b8272834429dfdaadc63c711f8fe2df7cb10777be1ad02f83bc89c8f5261","observation_id":"bf166085-f870-4f4d-bd71-c6c1dacb4a81","resolution":{"observed_at":"2026-08-07T14:14:21.569416Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-07T14:12:20.744113Z","title":"rstar-math: Small llms can master math reasoning with self-evolved deep thinking,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.19716","last_updated":"2025-05-26T09:04:44Z","snapshot_observed_at":"2026-08-09T02:14:56.474379Z","submitted_at":"2025-05-26T09:04:44Z","title":"Concise Reasoning, Big Gains: Pruning Long Reasoning Trace with Difficulty-Aware Prompting","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-07T14:12:20.744113Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2505.19716"},"observation_digest":"sha256:8b26bad9c43bb435f8c3b54629f4e84d6166f9f3be65e9747fef9114086efea2","observation_id":"3e200c4b-2622-4226-bfdb-4d622dc54603","resolution":{"observed_at":"2026-08-07T14:12:20.744113Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-07T13:53:56.145963Z","title":"rstar-math: Small llms can master math reasoning with self-evolved deep thinking.arXiv preprint arXiv:2501.04519,","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2505.20643","last_updated":"2025-05-27T02:44:00Z","snapshot_observed_at":"2026-08-07T21:57:04.990865Z","submitted_at":"2025-05-27T02:44:00Z","title":"Can Past Experience Accelerate LLM Reasoning?","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-07T13:53:56.145963Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2505.20643"},"observation_digest":"sha256:0d8060ed471e0e61d9068468e789add3a0f0149dc70b01ddeaca01fab1e31f54","observation_id":"1216ef42-4154-4d8a-b68e-e40f0f23babc","resolution":{"observed_at":"2026-08-07T13:53:56.145963Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-07T13:40:16.872558Z","title":"rstar-math: Small llms can master math reasoning with self-evolved deep thinking.arXiv preprint arXiv:2501.04519, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.21297","last_updated":"2025-05-27T15:00:57Z","snapshot_observed_at":"2026-08-12T00:40:56.022178Z","submitted_at":"2025-05-27T15:00:57Z","title":"rStar-Coder: Scaling Competitive Code Reasoning with a Large-Scale Verified Dataset","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-07T13:40:16.872558Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2505.21297"},"observation_digest":"sha256:79f33b9e24646d4e89d5be337e1858e66424a27cb686dd350767603fb52c5763","observation_id":"6db72fc4-c218-4d06-a80d-91872720b3d0","resolution":{"observed_at":"2026-08-07T13:40:16.872558Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-07T13:34:12.734899Z","title":"rstar-math: Small llms can master math reasoning with self-evolved deep thinking.arXiv preprint arXiv:2501.04519, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.21496","last_updated":"2025-05-27T17:58:06Z","snapshot_observed_at":"2026-08-08T15:12:37.593143Z","submitted_at":"2025-05-27T17:58:06Z","title":"UI-Genie: A Self-Improving Approach for Iteratively Boosting MLLM-based Mobile GUI Agents","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-07T13:34:12.734899Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2505.21496"},"observation_digest":"sha256:a055e120e270f4198a666f50aa0b2e78d9b233c2f8ad0d8b0fa485895ad5dbf9","observation_id":"f779bc17-e800-4b2d-bdc9-34e5539ba275","resolution":{"observed_at":"2026-08-07T13:34:12.734899Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-07T13:14:29.991659Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.22430","last_updated":"2025-05-28T14:55:33Z","snapshot_observed_at":"2026-08-11T15:37:34.461537Z","submitted_at":"2025-05-28T14:55:33Z","title":"RAG-Zeval: Towards Robust and Interpretable Evaluation on RAG Responses through End-to-End Rule-Guided Reasoning","version":1},"reference_index":9,"source":"arxiv_source","source_observed_at":"2026-08-07T13:14:29.991659Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2505.22430"},"observation_digest":"sha256:1c1559e74ed930faf7634b31c09c5aded4cd40d705680ad931f7e4b470b0d188","observation_id":"2d721ffe-0a6c-461f-9e63-ff41c627315e","resolution":{"observed_at":"2026-08-07T13:14:29.991659Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-07T12:32:21.738872Z","title":"rstar-math: Small llms can master math reasoning with self-evolved deep thinking, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2505.24273","last_updated":"2026-08-09T16:14:48Z","snapshot_observed_at":"2026-08-12T03:16:09.415159Z","submitted_at":"2025-05-30T06:49:00Z","title":"How Much Backtracking is Enough? Exploring the Interplay of SFT and RL in Enhancing LLM Reasoning","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-07T12:32:21.738872Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2505.24273"},"observation_digest":"sha256:4af0216ed1f771a89cc431f9f581102579a3ee79a9a6306652f29741690d72d2","observation_id":"5f7f1c3c-cd7a-45be-ad00-ebd75ad3a8b4","resolution":{"observed_at":"2026-08-07T12:32:21.738872Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":"2501.04519","doi":"10.48550/arxiv.2501.04519","metadata_source":"pith","pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","venue":"cs.CL","work_id":"49792b83-569e-4f5f-ae80-e96cbd3b7a43","year":2025},"citing_paper":{"arxiv_id":"2506.01939","last_updated":"2025-11-13T10:08:29Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-06-02T17:54:39Z","title":"Beyond the 80/20 Rule: High-Entropy Minority Tokens Drive Effective Reinforcement Learning for LLM Reasoning","version":2},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-05-12T12:12:08.724844Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2506.01939"},"observation_digest":"sha256:f2892b815fedacb3fdfa8a6c316789c9f434c9c45fe214f07211e89d4e220d76","observation_id":"317aff50-0850-44ce-bba9-846ccbf27ed4","resolution":{"observed_at":"2026-05-17T04:43:48.343981Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-07T11:30:40.139932Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.02338","last_updated":"2025-06-03T00:29:15Z","snapshot_observed_at":"2026-08-08T15:09:47.963101Z","submitted_at":"2025-06-03T00:29:15Z","title":"One Missing Piece for Open-Source Reasoning Models: A Dataset to Mitigate Cold-Starting Short CoT LLMs in RL","version":1},"reference_index":5,"source":"arxiv_source","source_observed_at":"2026-08-07T11:30:40.139932Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2506.02338"},"observation_digest":"sha256:577376f562cf7a415bd5317c0bbe88d87b08f962ace77c24e9d99c0930b74ffa","observation_id":"abc0b61a-31c9-4586-8bf7-70f503b7e8fe","resolution":{"observed_at":"2026-08-07T11:30:40.139932Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-07T11:02:39.226807Z","title":"L., Liu, Y., Shang, N., Sun, Y., Zhu, Y., Yang, F., and Yang, M","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.03700","last_updated":"2025-06-04T08:32:30Z","snapshot_observed_at":"2026-08-08T01:16:24.422333Z","submitted_at":"2025-06-04T08:32:30Z","title":"AdaDecode: Accelerating LLM Decoding with Adaptive Layer Parallelism","version":1},"reference_index":26,"source":"arxiv_source","source_observed_at":"2026-08-07T11:02:39.226807Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2506.03700"},"observation_digest":"sha256:b3df986cbf6c65bafd8f984dfa7621d9e51d263b16c1085fffe15061e02d4bea","observation_id":"94f56415-7df1-433c-b465-91a8350e4ce4","resolution":{"observed_at":"2026-08-07T11:02:39.226807Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-07T10:56:23.571793Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.03978","last_updated":"2025-06-09T04:15:05Z","snapshot_observed_at":"2026-08-12T02:45:06.429220Z","submitted_at":"2025-06-04T14:08:44Z","title":"Structured Pruning for Diverse Best-of-N Reasoning Optimization","version":2},"reference_index":16,"source":"arxiv_source","source_observed_at":"2026-08-07T10:56:23.571793Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2506.03978"},"observation_digest":"sha256:9becb489e9bf0904134e62f7db2b8153994d544e78c97de975f99f100185ea87","observation_id":"82381dd6-b3cb-4230-a172-84bd0c8bfe11","resolution":{"observed_at":"2026-08-07T10:56:23.571793Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-07T10:21:27.136837Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.05936","last_updated":"2025-06-06T10:02:13Z","snapshot_observed_at":"2026-08-07T10:44:40.702753Z","submitted_at":"2025-06-06T10:02:13Z","title":"DynamicMind: A Tri-Mode Thinking System for Large Language Models","version":1},"reference_index":14,"source":"arxiv_source","source_observed_at":"2026-08-07T10:21:27.136837Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2506.05936"},"observation_digest":"sha256:3d14a7e17bca86617f16f75fdd4f7ec0c72c6695b19cc00c90277f26b512d3b7","observation_id":"8e199341-7a67-43c8-a4fa-e3400918a3eb","resolution":{"observed_at":"2026-08-07T10:21:27.136837Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-07T06:07:02.005898Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.06009","last_updated":"2025-06-06T11:54:06Z","snapshot_observed_at":"2026-08-10T14:21:07.495292Z","submitted_at":"2025-06-06T11:54:06Z","title":"Unlocking Recursive Thinking of LLMs: Alignment via Refinement","version":1},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-08-07T06:07:02.005898Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2506.06009"},"observation_digest":"sha256:d77984ed9d26f65fdc173482841bc72d6b937fc2b56b641210398876201f4437","observation_id":"ffc07461-ffdc-44a3-8845-da654ca7584f","resolution":{"observed_at":"2026-08-07T06:07:02.005898Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-07T05:14:45.965584Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2506.08446","last_updated":"2025-06-10T04:44:28Z","snapshot_observed_at":"2026-08-07T09:15:20.058485Z","submitted_at":"2025-06-10T04:44:28Z","title":"A Survey on Large Language Models for Mathematical Reasoning","version":1},"reference_index":2014,"source":"pdf_text","source_observed_at":"2026-08-07T05:14:45.965584Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2506.08446"},"observation_digest":"sha256:878c8638312c61154dd52a2467f3a2a13711fde69adcf46395b0ee98f59da5bf","observation_id":"bfdd0ab2-dde6-46e3-a06c-933d250e130b","resolution":{"observed_at":"2026-08-07T05:14:45.965584Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-07T05:01:18.720845Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.08927","last_updated":"2025-06-10T15:51:16Z","snapshot_observed_at":"2026-08-09T20:17:32.286068Z","submitted_at":"2025-06-10T15:51:16Z","title":"Socratic-MCTS: Test-Time Visual Reasoning by Asking the Right Questions","version":1},"reference_index":10,"source":"arxiv_source","source_observed_at":"2026-08-07T05:01:18.720845Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2506.08927"},"observation_digest":"sha256:f654a72045412cf1c0e221130df5cc8b25b099876fb8492ca1d42c8cd9472999","observation_id":"37401fbc-e87f-498d-9c1b-9519b094ab57","resolution":{"observed_at":"2026-08-07T05:01:18.720845Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-07T05:01:24.112886Z","title":"rstar-math: Small llms can master math reasoning with self-evolved deep thinking.arXiv preprint arXiv:2501.04519, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.08989","last_updated":"2025-06-10T17:02:00Z","snapshot_observed_at":"2026-08-07T04:55:03.980740Z","submitted_at":"2025-06-10T17:02:00Z","title":"SwS: Self-aware Weakness-driven Problem Synthesis in Reinforcement Learning for LLM Reasoning","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-07T05:01:24.112886Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2506.08989"},"observation_digest":"sha256:ef1241b8786bc25c83e73fec02bb428b8c673f802226b7e1e7d3eb0cbf0aabd2","observation_id":"06e02db3-49d4-4303-b2bb-2dae77510a5a","resolution":{"observed_at":"2026-08-07T05:01:24.112886Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-07T04:38:41.146393Z","title":"Chaoqun He, Renjie Luo, Yuzhuo Bai, Shengding Hu, Zhen Thai, Junhao Shen, Jinyi Hu, Xu Han, Yujie Huang, Yuxiang Zhang, Jie Liu, Lei Qi, Zhiyuan Liu, and Maosong Sun","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.10209","last_updated":"2025-06-11T22:03:19Z","snapshot_observed_at":"2026-08-07T04:29:32.419935Z","submitted_at":"2025-06-11T22:03:19Z","title":"TTT-Bench: A Benchmark for Evaluating Reasoning Ability with Simple and Novel Tic-Tac-Toe-style Games","version":1},"reference_index":2025,"source":"pdf_text","source_observed_at":"2026-08-07T04:38:41.146393Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2506.10209"},"observation_digest":"sha256:63a6c295af182b0a33599a342f3f60a75c00c45b528aa3be2b732a386bc5f2ba","observation_id":"301449ee-3612-4f6f-a783-45d2679d11f0","resolution":{"observed_at":"2026-08-07T04:38:41.146393Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-07T00:49:34.039982Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking[J]","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.12728","last_updated":"2025-06-15T05:42:01Z","snapshot_observed_at":"2026-08-10T14:35:40.671292Z","submitted_at":"2025-06-15T05:42:01Z","title":"MCTS-Refined CoT: High-Quality Fine-Tuning Data for LLM-Based Repository Issue Resolution","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-07T00:49:34.039982Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2506.12728"},"observation_digest":"sha256:4927d6430769e21599fcb8b9331f37cb085cf1de368c9fce3b5c7ed7ccce1b83","observation_id":"59d93812-8bd4-4b1c-8421-7c8c5550685c","resolution":{"observed_at":"2026-08-07T00:49:34.039982Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-07T00:24:14.735972Z","title":"rstar-math: Small llms can master math reasoning with self-evolved deep thinking, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.14234","last_updated":"2025-06-17T06:47:19Z","snapshot_observed_at":"2026-08-10T00:27:45.764398Z","submitted_at":"2025-06-17T06:47:19Z","title":"Xolver: Multi-Agent Reasoning with Holistic Experience Learning Just Like an Olympiad Team","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-07T00:24:14.735972Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2506.14234"},"observation_digest":"sha256:c620f7b2db0d230d43a908bdf855f047739a7dbda5cbb86d49e43dec49c14106","observation_id":"05224726-4f85-4934-aa07-8d6c13110327","resolution":{"observed_at":"2026-08-07T00:24:14.735972Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-06T23:35:35.759763Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.17533","last_updated":"2025-06-21T01:11:01Z","snapshot_observed_at":"2026-08-09T12:43:12.249421Z","submitted_at":"2025-06-21T01:11:01Z","title":"DuaShepherd: Integrating Stepwise Correctness and Potential Rewards for Mathematical Reasoning","version":1},"reference_index":10,"source":"arxiv_source","source_observed_at":"2026-08-06T23:35:35.759763Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2506.17533"},"observation_digest":"sha256:29838ea8044c171777c97a088f54e939f5548346bc22aa6da0309465bd8881b9","observation_id":"2656d7cb-e0c5-402f-bb80-5de2452c58f7","resolution":{"observed_at":"2026-08-06T23:35:35.759763Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-06T23:11:58.354030Z","title":"rstar- math: Small llms can master math reasoning with self- evolved deep thinking","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.19171","last_updated":"2025-06-23T22:10:38Z","snapshot_observed_at":"2026-08-08T10:48:33.185126Z","submitted_at":"2025-06-23T22:10:38Z","title":"Distilling Tool Knowledge into Language Models via Back-Translated Traces","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-06T23:11:58.354030Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2506.19171"},"observation_digest":"sha256:a686e4fe1f3c4cca59fa2836ecc55e18eace9162a596a120420f52daa0426d89","observation_id":"e85fecc2-2d41-4bdd-ad63-eef3dcabfd24","resolution":{"observed_at":"2026-08-06T23:11:58.354030Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-07T04:07:56.268149Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.21571","last_updated":"2026-07-28T12:38:29Z","snapshot_observed_at":"2026-08-11T21:42:03.576031Z","submitted_at":"2025-06-13T05:40:56Z","title":"Towards Understanding the Cognitive Habits of Large Reasoning Models","version":3},"reference_index":15,"source":"arxiv_source","source_observed_at":"2026-08-07T04:07:56.268149Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2506.21571"},"observation_digest":"sha256:e858c2925d2dba07346aaae7fb0cce10316c1728e2fbebdd3b41fce27e23377d","observation_id":"d26efe52-e749-4d62-9209-0bb1d6082afe","resolution":{"observed_at":"2026-08-07T04:07:56.268149Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-06T20:46:44.759235Z","title":null,"venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2507.01951","last_updated":"2025-07-09T12:28:31Z","snapshot_observed_at":"2026-08-10T04:39:01.641296Z","submitted_at":"2025-07-02T17:58:01Z","title":"Test-Time Scaling with Reflective Generative Model","version":2},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-06T20:46:44.759235Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2507.01951"},"observation_digest":"sha256:4363ea15e63a5963d8397d90e8a5eb05acf11746de00f1eb6ed996dfa5dd0522","observation_id":"8226f95e-f046-4cbf-919f-72c19f342603","resolution":{"observed_at":"2026-08-06T20:46:44.759235Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-06T19:28:12.313777Z","title":"arXiv preprint arXiv:2501.04519","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.05557","last_updated":"2025-07-08T00:41:12Z","snapshot_observed_at":"2026-08-06T19:21:14.781528Z","submitted_at":"2025-07-08T00:41:12Z","title":"Enhancing Test-Time Scaling of Large Language Models with Hierarchical Retrieval-Augmented MCTS","version":1},"reference_index":2025,"source":"pdf_text","source_observed_at":"2026-08-06T19:28:12.313777Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2507.05557"},"observation_digest":"sha256:417d354dba07924f9d708615dee1f14bd2f3bc37cc1e6d14adaea74cd072a771","observation_id":"6c2624e4-c766-48db-adb9-7fdc2332769a","resolution":{"observed_at":"2026-08-06T19:28:12.313777Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-06T18:24:50.462671Z","title":"Bilevel programming for hyperparameter optimization and meta-learning","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2507.08501","last_updated":"2025-07-11T11:24:09Z","snapshot_observed_at":"2026-08-07T23:20:34.881688Z","submitted_at":"2025-07-11T11:24:09Z","title":"From Language to Logic: A Bi-Level Framework for Structured Reasoning","version":1},"reference_index":906,"source":"pdf_text","source_observed_at":"2026-08-06T18:24:50.462671Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2507.08501"},"observation_digest":"sha256:36f4768e1b6ef7580a10fcdbaec4485d0bb758d5d59b11862e6c0d8274a636cf","observation_id":"3113f3ab-8d7c-4976-bfb2-2ace325cef0b","resolution":{"observed_at":"2026-08-06T18:24:50.462671Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-06T18:04:53.340378Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.09374","last_updated":"2025-07-12T18:44:32Z","snapshot_observed_at":"2026-08-10T00:10:49.037395Z","submitted_at":"2025-07-12T18:44:32Z","title":"EduFlow: Advancing MLLMs' Problem-Solving Proficiency through Multi-Stage, Multi-Perspective Critique","version":1},"reference_index":10,"source":"arxiv_source","source_observed_at":"2026-08-06T18:04:53.340378Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2507.09374"},"observation_digest":"sha256:04c5956c58a5c19a6bff90459f96e29c3dfff955b5ddb83d329051be1fdb6d58","observation_id":"474e7e1f-ebec-4769-90c1-086e16e92013","resolution":{"observed_at":"2026-08-06T18:04:53.340378Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-06T17:29:26.331533Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.12484","last_updated":"2025-07-14T20:35:16Z","snapshot_observed_at":"2026-08-08T07:11:37.298140Z","submitted_at":"2025-07-14T20:35:16Z","title":"AI-Powered Math Tutoring: Platform for Personalized and Adaptive Education","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-06T17:29:26.331533Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2507.12484"},"observation_digest":"sha256:b95965bc073764b8a636c1c4f6dc56ecd8f68cd1fd112c2d953332e832f340e7","observation_id":"6dfdf8e0-b155-410e-8274-540a77f81581","resolution":{"observed_at":"2026-08-06T17:29:26.331533Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-06T14:44:35.823791Z","title":"Preprint, arXiv:2501.04519","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2507.17849","last_updated":"2025-07-23T18:17:22Z","snapshot_observed_at":"2026-08-10T07:28:39.609575Z","submitted_at":"2025-07-23T18:17:22Z","title":"Dynamic and Generalizable Process Reward Modeling","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-06T14:44:35.823791Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2507.17849"},"observation_digest":"sha256:064a29ec5640fa96945ded842a95e5d206be98934dc0307e0c58ae9f9e944355","observation_id":"d753e3c7-eb92-43de-a765-31c634f38b70","resolution":{"observed_at":"2026-08-06T14:44:35.823791Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-06T14:37:28.746721Z","title":"Proceedings of the National Academy of Sciences, 120(30):e2305016120","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2507.18584","last_updated":"2025-07-24T17:03:27Z","snapshot_observed_at":"2026-08-10T02:44:19.084180Z","submitted_at":"2025-07-24T17:03:27Z","title":"AQuilt: Weaving Logic and Self-Inspection into Low-Cost, High-Relevance Data Synthesis for Specialist LLMs","version":1},"reference_index":2023,"source":"pdf_text","source_observed_at":"2026-08-06T14:37:28.746721Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2507.18584"},"observation_digest":"sha256:ffe2493a5173ca9f107e5bc38bcda9f96da2e0dc438d2560f3c4588da8f278c8","observation_id":"6dfeb8e6-c345-49ff-9bc1-20b33a167609","resolution":{"observed_at":"2026-08-06T14:37:28.746721Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":"2501.04519","doi":"10.48550/arxiv.2501.04519","metadata_source":"pith","pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","venue":"cs.CL","work_id":"49792b83-569e-4f5f-ae80-e96cbd3b7a43","year":2025},"citing_paper":{"arxiv_id":"2507.21046","last_updated":"2026-01-16T20:59:08Z","snapshot_observed_at":"2026-08-01T06:32:44.461162Z","submitted_at":"2025-07-28T17:59:05Z","title":"A Survey of Self-Evolving Agents: What, When, How, and Where to Evolve on the Path to Artificial Super Intelligence","version":4},"reference_index":234,"source":"arxiv_source","source_observed_at":"2026-05-14T22:23:14.621091Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2507.21046"},"observation_digest":"sha256:748b5e62e74c094b99d965eda641ded79e758efed8f2b84e80d3276a7e14cf27","observation_id":"fdb1f98b-6f30-4fc9-aa4c-de8503fd383a","resolution":{"observed_at":"2026-05-17T04:43:48.343981Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-05T23:29:17.082217Z","title":"rstar-math: Small llms can master math reasoning with self-evolved deep thinking.arXiv preprint arXiv:2501.04519, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2508.05383","last_updated":"2025-08-07T13:31:21Z","snapshot_observed_at":"2026-08-08T12:45:30.983934Z","submitted_at":"2025-08-07T13:31:21Z","title":"StructVRM: Aligning Multimodal Reasoning with Structured and Verifiable Reward Models","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-05T23:29:17.082217Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2508.05383"},"observation_digest":"sha256:e348ff3cd9632f26c877bd2891c6ac3df4b0e8e6778497b6b8f0a73fa5fba25a","observation_id":"9f30538b-8588-4cba-add5-7cc3256747de","resolution":{"observed_at":"2026-08-05T23:29:17.082217Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-05T23:25:50.433469Z","title":"rstar-math: Small llms can master math reasoning with self-evolved deep thinking.arXiv preprint arXiv:2501.04519, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2508.05388","last_updated":"2025-08-07T13:38:49Z","snapshot_observed_at":"2026-08-06T12:42:41.700036Z","submitted_at":"2025-08-07T13:38:49Z","title":"An Explainable Machine Learning Framework for Railway Predictive Maintenance using Data Streams from the Metro Operator of Portugal","version":1},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-05T23:25:50.433469Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2508.05388"},"observation_digest":"sha256:43149273a2bbedab8d6184bad39daa7e44e622ed4ae4c2b5f67b0e0e0b6177ec","observation_id":"d9fac228-abbe-4aac-8de2-91b4c37dae0f","resolution":{"observed_at":"2026-08-05T23:25:50.433469Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":"2501.04519","doi":"10.48550/arxiv.2501.04519","metadata_source":"pith","pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","venue":"cs.CL","work_id":"49792b83-569e-4f5f-ae80-e96cbd3b7a43","year":2025},"citing_paper":{"arxiv_id":"2508.06412","last_updated":"2026-05-07T07:25:18Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2025-08-08T15:56:49Z","title":"Sample-efficient LLM Optimization with Reset Replay","version":3},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-05-18T23:38:08.044450Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2508.06412"},"observation_digest":"sha256:9ca42ddcdd343e3531dd6879605fbd2f8c31a5888d8dae486e5f86fa4ecad800","observation_id":"62909ce8-c75d-4330-8c4d-8d2d6a3e3c6e","resolution":{"observed_at":"2026-05-18T23:41:54.694793Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-05T21:23:23.591867Z","title":"Guan, X.; Zhang, L","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2508.15791","last_updated":"2025-08-12T11:53:57Z","snapshot_observed_at":"2026-08-09T01:20:59.427158Z","submitted_at":"2025-08-12T11:53:57Z","title":"InteChar: A Unified Oracle Bone Character List for Ancient Chinese Language Modeling","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-05T21:23:23.591867Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2508.15791"},"observation_digest":"sha256:26a6087ae16afb6b767d90670fc4a34e80c72b80ab89fc839b28e30104206541","observation_id":"1ab3c225-c272-492f-8113-19d1233e2401","resolution":{"observed_at":"2026-08-05T21:23:23.591867Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-05T14:57:39.405467Z","title":"Glaive function calling v2 dataset","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2508.20722","last_updated":"2025-08-28T12:45:25Z","snapshot_observed_at":"2026-08-09T05:32:48.085714Z","submitted_at":"2025-08-28T12:45:25Z","title":"rStar2-Agent: Agentic Reasoning Technical Report","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-05T14:57:39.405467Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2508.20722"},"observation_digest":"sha256:8a9229aceb008958c4566a047aabd3e943075fe9393b7bf76e60764a57dc9d32","observation_id":"1e6a9ace-d75a-451e-b5c1-cc2d3f741e21","resolution":{"observed_at":"2026-08-05T14:57:39.405467Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":"2501.04519","doi":"10.48550/arxiv.2501.04519","metadata_source":"pith","pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","venue":"cs.CL","work_id":"49792b83-569e-4f5f-ae80-e96cbd3b7a43","year":2025},"citing_paper":{"arxiv_id":"2509.02547","last_updated":"2026-04-17T18:09:08Z","snapshot_observed_at":"2026-08-03T09:07:42.489237Z","submitted_at":"2025-09-02T17:46:26Z","title":"The Landscape of Agentic Reinforcement Learning for LLMs: A Survey","version":5},"reference_index":202,"source":"pdf_text","source_observed_at":"2026-05-18T19:19:36.427337Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2509.02547"},"observation_digest":"sha256:2ced7b89b326feddb9e199482b1666e0e5c06aed5d6b3f9efe160d986e5a4444","observation_id":"e887d306-6f75-4c59-9230-544deff51a20","resolution":{"observed_at":"2026-05-18T19:21:48.008345Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-05T13:52:04.430116Z","title":"rstar-math: Small llms can master math reasoning with self-evolved deep thinking, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2509.04475","last_updated":"2025-08-30T03:09:07Z","snapshot_observed_at":"2026-08-08T21:44:17.416488Z","submitted_at":"2025-08-30T03:09:07Z","title":"ParaThinker: Native Parallel Thinking as a New Paradigm to Scale LLM Test-time Compute","version":1},"reference_index":15,"source":"arxiv_source","source_observed_at":"2026-08-05T13:52:04.430116Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2509.04475"},"observation_digest":"sha256:36940940705e5aad9c7fa572372d62a47f9d45bf4b5c5a92bca586611cadf9f2","observation_id":"255eb9be-145d-4f37-a01e-831da3972f87","resolution":{"observed_at":"2026-08-05T13:52:04.430116Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":"2501.04519","doi":"10.48550/arxiv.2501.04519","metadata_source":"pith","pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","venue":"cs.CL","work_id":"49792b83-569e-4f5f-ae80-e96cbd3b7a43","year":2025},"citing_paper":{"arxiv_id":"2509.25454","last_updated":"2026-04-06T19:16:24Z","snapshot_observed_at":"2026-07-06T22:31:10.099674Z","submitted_at":"2025-09-29T20:00:29Z","title":"DeepSearch: Overcome the Bottleneck of Reinforcement Learning with Verifiable Rewards via Monte Carlo Tree Search","version":4},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-05-18T12:12:25.437344Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2509.25454"},"observation_digest":"sha256:b8abf16646e0a2e1b916d5a9bddcb4dc7c62db031b29775769115dcb7ae36c48","observation_id":"c04d4468-5067-4343-8006-c47994ef6c65","resolution":{"observed_at":"2026-05-18T12:12:35.851659Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":"2501.04519","doi":"10.48550/arxiv.2501.04519","metadata_source":"pith","pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","venue":"cs.CL","work_id":"49792b83-569e-4f5f-ae80-e96cbd3b7a43","year":2025},"citing_paper":{"arxiv_id":"2510.18245","last_updated":"2026-05-13T04:16:31Z","snapshot_observed_at":"2026-08-08T00:42:38.787054Z","submitted_at":"2025-10-21T03:08:48Z","title":"Scaling Laws Meet Model Architecture: Toward Inference-Efficient LLMs","version":3},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-05-18T05:30:11.389756Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2510.18245"},"observation_digest":"sha256:2319f2a7560b82a085dc49f1e8b10b3d781cc416681c0701d8e78bae7d6b2a82","observation_id":"2b12d6ce-7fe1-40c5-baf4-25f423b5945b","resolution":{"observed_at":"2026-05-18T05:30:55.039488Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-04T07:56:46.376327Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2510.24803","last_updated":"2026-07-15T00:10:20Z","snapshot_observed_at":"2026-08-08T21:42:31.396671Z","submitted_at":"2025-10-28T00:48:20Z","title":"MASPRM: Multi-Agent System Process Reward Model","version":3},"reference_index":9,"source":"arxiv_source","source_observed_at":"2026-08-04T07:56:46.376327Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2510.24803"},"observation_digest":"sha256:8a30cae1644ab7d302547d5825fe99808f41ccc9716644e7ea13d215086868f6","observation_id":"0c1c197b-2910-421a-b4f6-af84b593fac8","resolution":{"observed_at":"2026-08-04T07:56:46.376327Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":"2501.04519","doi":"10.48550/arxiv.2501.04519","metadata_source":"pith","pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","venue":"cs.CL","work_id":"49792b83-569e-4f5f-ae80-e96cbd3b7a43","year":2025},"citing_paper":{"arxiv_id":"2511.00066","last_updated":"2026-05-13T07:57:37Z","snapshot_observed_at":"2026-08-06T07:33:01.385285Z","submitted_at":"2025-10-29T08:07:47Z","title":"Sharpness-Guided Group Relative Policy Optimization via Probability Shaping","version":4},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-05-18T03:10:57.839146Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2511.00066"},"observation_digest":"sha256:eb8fbffc7acd4090cbfd7db6a7d09622bfd9a37e8426fe66024c49405be153de","observation_id":"a7b835e5-3650-4486-af52-465b278ec30d","resolution":{"observed_at":"2026-05-18T03:12:22.378053Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":"2501.04519","doi":"10.48550/arxiv.2501.04519","metadata_source":"pith","pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","venue":"cs.CL","work_id":"49792b83-569e-4f5f-ae80-e96cbd3b7a43","year":2025},"citing_paper":{"arxiv_id":"2511.22277","last_updated":"2026-04-24T08:52:17Z","snapshot_observed_at":"2026-08-11T02:15:49.098646Z","submitted_at":"2025-11-27T09:59:39Z","title":"TreeCoder: Systematic Exploration and Optimisation of Decoding and Constraints for LLM Code Generation","version":2},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-05-17T04:19:18.805697Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2511.22277"},"observation_digest":"sha256:1deb3cad709ed48b1ed55afd0f428820a01bc1c0e78cd5d5f87558ea8bccd76e","observation_id":"4e046aa7-c60f-4465-9b24-55412e400030","resolution":{"observed_at":"2026-05-17T04:43:48.343981Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":"2501.04519","doi":"10.48550/arxiv.2501.04519","metadata_source":"pith","pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","venue":"cs.CL","work_id":"49792b83-569e-4f5f-ae80-e96cbd3b7a43","year":2025},"citing_paper":{"arxiv_id":"2512.14735","last_updated":"2026-04-08T06:53:15Z","snapshot_observed_at":"2026-08-11T12:01:10.700411Z","submitted_at":"2025-12-11T06:04:33Z","title":"PyFi: Toward Pyramid-like Financial Image Understanding for VLMs via Adversarial Agents","version":2},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-05-16T23:38:28.197932Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2512.14735"},"observation_digest":"sha256:f5bc51ec32856deef41bbee825da9b973ce9d47270080e510adaeb2a5b2afe40","observation_id":"152560ad-ddcc-4670-a512-74c44d9220e2","resolution":{"observed_at":"2026-05-17T04:43:48.343981Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":"2501.04519","doi":"10.48550/arxiv.2501.04519","metadata_source":"pith","pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","venue":"cs.CL","work_id":"49792b83-569e-4f5f-ae80-e96cbd3b7a43","year":2025},"citing_paper":{"arxiv_id":"2601.21619","last_updated":"2026-05-09T01:57:30Z","snapshot_observed_at":"2026-08-11T01:08:48.247252Z","submitted_at":"2026-01-29T12:22:45Z","title":"On the Overscaling Curse of Parallel Thinking: System Efficacy Contradicts Sample Efficiency","version":2},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-05-16T10:38:09.875786Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2601.21619"},"observation_digest":"sha256:63ec3ca9078dc0299cbfc6ab62cdf2c9b1efad73707c8af0f430c57a69616444","observation_id":"4ca6de60-51a9-4c36-ab8b-9c0b293768c0","resolution":{"observed_at":"2026-05-17T04:43:48.343981Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-03T02:43:34.738415Z","title":"rstar-math: Small llms can master math reasoning with self-evolved deep thinking.arXiv preprint arXiv:2501.04519, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2602.10014","last_updated":"2026-05-31T19:13:22Z","snapshot_observed_at":"2026-08-06T02:34:35.293978Z","submitted_at":"2026-02-10T17:36:41Z","title":"A Task-Centric Theory for Iterative Self-Improvement with Easy-to-Hard Curricula","version":3},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-03T02:43:34.738415Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2602.10014"},"observation_digest":"sha256:9eb1f6b20e29784d611f8dad98a46c8b81b5f0ddecb58a4c0a26b8abaeea5d41","observation_id":"2a65ce3f-459f-4877-9f61-32c6531bc8d9","resolution":{"observed_at":"2026-08-03T02:43:34.738415Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-02T23:50:34.097649Z","title":"L., Liu, Y., Shang, N., Sun, Y., Zhu, Y., Yang, F., and Yang, M","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2602.12586","last_updated":"2026-05-27T16:01:59Z","snapshot_observed_at":"2026-08-08T22:46:16.097170Z","submitted_at":"2026-02-13T03:56:22Z","title":"Can I Have Your Order? Monte-Carlo Tree Search for Slot Filling Ordering in Diffusion Language Models","version":2},"reference_index":17,"source":"arxiv_source","source_observed_at":"2026-08-02T23:50:34.097649Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2602.12586"},"observation_digest":"sha256:48e5c89606f3f735a0076a727c2f722c145a09c57f0a0307db1e7b681152f9d3","observation_id":"45829dcc-6a81-4821-b00a-25494ece48f0","resolution":{"observed_at":"2026-08-02T23:50:34.097649Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":"2501.04519","doi":"10.48550/arxiv.2501.04519","metadata_source":"pith","pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","venue":"cs.CL","work_id":"49792b83-569e-4f5f-ae80-e96cbd3b7a43","year":2025},"citing_paper":{"arxiv_id":"2604.10228","last_updated":"2026-05-28T16:26:54Z","snapshot_observed_at":"2026-08-11T07:43:38.834025Z","submitted_at":"2026-04-11T14:25:17Z","title":"SVSR: A Self-Verification and Self-Rectification Paradigm for Multimodal Reasoning","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-05-10T15:59:04.802241Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2604.10228"},"observation_digest":"sha256:9242e58de109698edb77d877b238a91ac8c8e448674b5c295295da6291e5cfb7","observation_id":"d82cbb9a-454d-484b-972c-986e1b46981a","resolution":{"observed_at":"2026-05-17T04:43:48.343981Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-07-12T22:47:33.917420Z","title":"rstar-math: Small llms can master math reasoning with self-evolved deep thinking.arXiv preprint arXiv:2501.04519, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2604.10228","last_updated":"2026-05-28T16:26:54Z","snapshot_observed_at":"2026-08-11T07:43:38.834025Z","submitted_at":"2026-04-11T14:25:17Z","title":"SVSR: A Self-Verification and Self-Rectification Paradigm for Multimodal Reasoning","version":2},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-07-12T22:47:33.917420Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2604.10228"},"observation_digest":"sha256:de0c489bb681632c91dfe67c7fe62a68fda49cca3c56907e1c0c954d4c9011f5","observation_id":"bda38c06-9ec6-46ea-8562-c5027c81e588","resolution":{"observed_at":"2026-07-12T22:47:33.917420Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":"2501.04519","doi":"10.48550/arxiv.2501.04519","metadata_source":"pith","pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","venue":"cs.CL","work_id":"49792b83-569e-4f5f-ae80-e96cbd3b7a43","year":2025},"citing_paper":{"arxiv_id":"2604.14768","last_updated":"2026-04-16T08:29:22Z","snapshot_observed_at":"2026-08-02T15:11:59.771593Z","submitted_at":"2026-04-16T08:29:22Z","title":"CoTEvol: Self-Evolving Chain-of-Thoughts for Data Synthesis in Mathematical Reasoning","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-05-10T10:58:27.736448Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2604.14768"},"observation_digest":"sha256:04a75d79262a50348111f881951ec3fff7007a180799719d6e5146443189e91e","observation_id":"76a04c12-74a7-4501-a022-b6575dd8aab0","resolution":{"observed_at":"2026-05-17T04:43:48.343981Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":"2501.04519","doi":"10.48550/arxiv.2501.04519","metadata_source":"pith","pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","venue":"cs.CL","work_id":"49792b83-569e-4f5f-ae80-e96cbd3b7a43","year":2025},"citing_paper":{"arxiv_id":"2604.18936","last_updated":"2026-04-21T00:21:05Z","snapshot_observed_at":"2026-07-06T23:05:39.724237Z","submitted_at":"2026-04-21T00:21:05Z","title":"Fine-Tuning Small Reasoning Models for Quantum Field Theory","version":1},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-05-10T03:23:18.770963Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2604.18936"},"observation_digest":"sha256:2c3461c3276e94855b3252b71d587f3149e438c441c83199f9b9125e09a53dd1","observation_id":"bf86c04b-90ac-4460-a3bf-f6cb036765e5","resolution":{"observed_at":"2026-05-17T04:43:48.343981Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":"2501.04519","doi":"10.48550/arxiv.2501.04519","metadata_source":"pith","pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","venue":"cs.CL","work_id":"49792b83-569e-4f5f-ae80-e96cbd3b7a43","year":2025},"citing_paper":{"arxiv_id":"2604.24037","last_updated":"2026-06-22T01:13:01Z","snapshot_observed_at":"2026-07-06T23:10:12.016186Z","submitted_at":"2026-04-27T04:43:42Z","title":"A Limit Theory of Foundation Models: A Mathematical Approach to Understanding Emergent Intelligence and Scaling Laws","version":3},"reference_index":75,"source":"pdf_text","source_observed_at":"2026-05-13T07:27:21.118156Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2604.24037"},"observation_digest":"sha256:adaed50a0fa5b87e619fc9de536294ad42d9acdc88517bdc63f743ca822ad179","observation_id":"84324a20-d793-4820-9f8a-78fecad56a1f","resolution":{"observed_at":"2026-05-17T04:43:48.343981Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":"2501.04519","doi":"10.48550/arxiv.2501.04519","metadata_source":"pith","pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","venue":"cs.CL","work_id":"49792b83-569e-4f5f-ae80-e96cbd3b7a43","year":2025},"citing_paper":{"arxiv_id":"2604.24037","last_updated":"2026-06-22T01:13:01Z","snapshot_observed_at":"2026-07-06T23:10:12.016186Z","submitted_at":"2026-04-27T04:43:42Z","title":"A Limit Theory of Foundation Models: A Mathematical Approach to Understanding Emergent Intelligence and Scaling Laws","version":4},"reference_index":76,"source":"pdf_text","source_observed_at":"2026-07-01T09:03:59.522516Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2604.24037"},"observation_digest":"sha256:91810e648c970fc3010181a3090958f8331205c2c701dfdd1a6e445391cd2396","observation_id":"48368654-04ff-49e0-84d2-81553fe4a869","resolution":{"observed_at":"2026-07-01T09:05:36.328782Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":"2501.04519","doi":"10.48550/arxiv.2501.04519","metadata_source":"pith","pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","venue":"cs.CL","work_id":"49792b83-569e-4f5f-ae80-e96cbd3b7a43","year":2025},"citing_paper":{"arxiv_id":"2604.24114","last_updated":"2026-04-27T07:10:35Z","snapshot_observed_at":"2026-08-03T01:17:11.673708Z","submitted_at":"2026-04-27T07:10:35Z","title":"IRIS: Interleaved Reinforcement with Incremental Staged Curriculum for Cross-Lingual Mathematical Reasoning","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-05-08T03:44:56.098128Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2604.24114"},"observation_digest":"sha256:18b00544fa0ae7954830af53fc8cd4e090bf03ba5e9e4018dd1c0df931cf42ca","observation_id":"431bed34-8c12-4e7b-b2f4-f4f571024482","resolution":{"observed_at":"2026-05-17T04:43:48.343981Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":"2501.04519","doi":"10.48550/arxiv.2501.04519","metadata_source":"pith","pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","venue":"cs.CL","work_id":"49792b83-569e-4f5f-ae80-e96cbd3b7a43","year":2025},"citing_paper":{"arxiv_id":"2605.07353","last_updated":"2026-05-08T07:08:25Z","snapshot_observed_at":"2026-08-11T04:46:32.902221Z","submitted_at":"2026-05-08T07:08:25Z","title":"Confidence-Aware Alignment Makes Reasoning LLMs More Reliable","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-05-11T02:10:40.020460Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2605.07353"},"observation_digest":"sha256:e0bd2dbaa1fd37c7b3f148f73923a7fc4b1aa92bad20e7948a73b454cbc60e17","observation_id":"980b8362-b7da-4dd0-943b-77bf37fa5c77","resolution":{"observed_at":"2026-05-17T04:43:48.343981Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":"2501.04519","doi":"10.48550/arxiv.2501.04519","metadata_source":"pith","pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","venue":"cs.CL","work_id":"49792b83-569e-4f5f-ae80-e96cbd3b7a43","year":2025},"citing_paper":{"arxiv_id":"2605.07600","last_updated":"2026-05-08T11:19:01Z","snapshot_observed_at":"2026-07-06T23:19:56.415523Z","submitted_at":"2026-05-08T11:19:01Z","title":"Mathematical Reasoning via Intervention-Based Time-Series Causal Discovery Using LLMs as Concept Mastery Simulators","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-05-11T02:05:48.483640Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2605.07600"},"observation_digest":"sha256:ac11ab5a78a52062248c3f945d9512e65babd9c42bcb05b0c9b7fd1c741ded4a","observation_id":"f820cb88-4dff-4174-a2d3-7ab13abeaacc","resolution":{"observed_at":"2026-05-17T04:43:48.343981Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":"2501.04519","doi":"10.48550/arxiv.2501.04519","metadata_source":"pith","pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","venue":"cs.CL","work_id":"49792b83-569e-4f5f-ae80-e96cbd3b7a43","year":2025},"citing_paper":{"arxiv_id":"2605.12289","last_updated":"2026-05-12T15:47:18Z","snapshot_observed_at":"2026-07-06T23:23:59.123377Z","submitted_at":"2026-05-12T15:47:18Z","title":"PriorZero: Bridging Language Priors and World Models for Decision Making","version":1},"reference_index":17,"source":"arxiv_source","source_observed_at":"2026-05-13T05:25:05.907123Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2605.12289"},"observation_digest":"sha256:af2c0fcee41d4d2ae1c1c15bd23cf2660d2fa8abb9c672db1d683dbbcb762516","observation_id":"68428fb5-62ff-4d18-b2d9-6d86d8def2e4","resolution":{"observed_at":"2026-05-17T04:43:48.343981Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":"2501.04519","doi":"10.48550/arxiv.2501.04519","metadata_source":"pith","pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","venue":"cs.CL","work_id":"49792b83-569e-4f5f-ae80-e96cbd3b7a43","year":2025},"citing_paper":{"arxiv_id":"2605.13511","last_updated":"2026-05-31T14:41:44Z","snapshot_observed_at":"2026-07-06T23:25:07.022065Z","submitted_at":"2026-05-13T13:30:12Z","title":"Many-Shot CoT-ICL: Making In-Context Learning Truly Learn","version":1},"reference_index":28,"source":"arxiv_source","source_observed_at":"2026-05-14T19:15:44.379686Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2605.13511"},"observation_digest":"sha256:01e18279c1def8b6112c5c0c18eca5ddd0ca813ec00afe69011406f95d4ff298","observation_id":"90b203d4-4837-4398-ba2a-202f27da154a","resolution":{"observed_at":"2026-05-17T04:43:48.343981Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":"2501.04519","doi":"10.48550/arxiv.2501.04519","metadata_source":"pith","pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","venue":"cs.CL","work_id":"49792b83-569e-4f5f-ae80-e96cbd3b7a43","year":2025},"citing_paper":{"arxiv_id":"2605.14790","last_updated":"2026-05-14T12:57:56Z","snapshot_observed_at":"2026-07-06T23:26:09.066863Z","submitted_at":"2026-05-14T12:57:56Z","title":"Graphs of Research: Citation Evolution Graphs as Supervision for Research Idea Generation","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-06-30T20:52:32.567107Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2605.14790"},"observation_digest":"sha256:29a14214c6f36f8b666c863f32822ca74095704d4d8f7c27a6a38d0611230af8","observation_id":"a8d5e88f-8379-4512-ab09-629fe4024d65","resolution":{"observed_at":"2026-06-30T20:55:04.014261Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":"2501.04519","doi":"10.48550/arxiv.2501.04519","metadata_source":"pith","pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","venue":"cs.CL","work_id":"49792b83-569e-4f5f-ae80-e96cbd3b7a43","year":2025},"citing_paper":{"arxiv_id":"2605.16727","last_updated":"2026-05-16T00:29:35Z","snapshot_observed_at":"2026-08-09T13:41:57.056244Z","submitted_at":"2026-05-16T00:29:35Z","title":"PopuLoRA: Co-Evolving LLM Populations for Reasoning Self-Play","version":1},"reference_index":64,"source":"arxiv_source","source_observed_at":"2026-05-19T21:37:56.570173Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2605.16727"},"observation_digest":"sha256:b7445528792d69ba651e31056c6a61f333ee0171c2a2201adeb2a60eb2b192a8","observation_id":"335db1e3-ef65-49d7-a9e9-c6de0432b6f6","resolution":{"observed_at":"2026-05-19T21:42:48.129101Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":"2501.04519","doi":"10.48550/arxiv.2501.04519","metadata_source":"pith","pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","venue":"cs.CL","work_id":"49792b83-569e-4f5f-ae80-e96cbd3b7a43","year":2025},"citing_paper":{"arxiv_id":"2605.18851","last_updated":"2026-05-13T11:04:31Z","snapshot_observed_at":"2026-07-06T23:29:38.069295Z","submitted_at":"2026-05-13T11:04:31Z","title":"STRIDE: Learnable Stepwise Language Feedback for LLM Reasoning","version":1},"reference_index":53,"source":"pdf_text","source_observed_at":"2026-05-20T20:47:16.236629Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2605.18851"},"observation_digest":"sha256:90590e88b9dbce4a8e65b033b3b1f5bbb7b7ce9bb562bd345bdcd7b2ffe7b922","observation_id":"d890c629-6899-4748-9604-6f67267dcb76","resolution":{"observed_at":"2026-05-20T20:49:00.759777Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-07-12T16:29:56.962393Z","title":"URL:https://arxiv.org/abs/2501.04519,arXiv:2501.04519","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2605.19723","last_updated":"2026-07-08T07:17:25Z","snapshot_observed_at":"2026-08-05T22:38:33.840848Z","submitted_at":"2026-05-19T11:56:03Z","title":"Mathematical Reasoning in Large Language Models: Benchmarks, Architectures, Evaluation, and Open Challenges","version":3},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-07-12T16:29:56.962393Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2605.19723"},"observation_digest":"sha256:302686c72ef521edb6bc3b805e56648ae638d69f71ace0069e11e8506cf662d7","observation_id":"89f1e9a3-e938-492c-bd54-04a02c7bbea4","resolution":{"observed_at":"2026-07-12T16:29:56.962393Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":"2501.04519","doi":"10.48550/arxiv.2501.04519","metadata_source":"pith","pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","venue":"cs.CL","work_id":"49792b83-569e-4f5f-ae80-e96cbd3b7a43","year":2025},"citing_paper":{"arxiv_id":"2605.25864","last_updated":"2026-05-25T13:55:12Z","snapshot_observed_at":"2026-07-06T23:35:46.157653Z","submitted_at":"2026-05-25T13:55:12Z","title":"When Self-Belief Misleads: Active Label Acquisition for Reinforcement Learning with Verifiable Rewards","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-06-29T22:58:46.313028Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2605.25864"},"observation_digest":"sha256:473ce4958af2298dc62d071ad46eb03f709c83febbaf7df53d2c89a328ed220f","observation_id":"cdd2fe36-2bab-4888-9b2f-9ea87f2b85d7","resolution":{"observed_at":"2026-06-29T23:04:01.277826Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":"2501.04519","doi":"10.48550/arxiv.2501.04519","metadata_source":"pith","pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","venue":"cs.CL","work_id":"49792b83-569e-4f5f-ae80-e96cbd3b7a43","year":2025},"citing_paper":{"arxiv_id":"2605.26924","last_updated":"2026-05-26T12:20:53Z","snapshot_observed_at":"2026-08-07T10:43:53.331655Z","submitted_at":"2026-05-26T12:20:53Z","title":"Learning to Adapt SFT Data for Better Reasoning Generalization","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-06-29T18:14:22.896656Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2605.26924"},"observation_digest":"sha256:6e19f107e26266ffc611f448aeed823be294421b46ccce3f67ca849bae566a80","observation_id":"0eaae3cb-75a9-4dd7-8adf-e5e48f807740","resolution":{"observed_at":"2026-06-29T18:23:51.093488Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":"2501.04519","doi":"10.48550/arxiv.2501.04519","metadata_source":"pith","pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","venue":"cs.CL","work_id":"49792b83-569e-4f5f-ae80-e96cbd3b7a43","year":2025},"citing_paper":{"arxiv_id":"2606.00618","last_updated":"2026-06-23T11:36:28Z","snapshot_observed_at":"2026-07-06T23:41:19.955034Z","submitted_at":"2026-05-30T08:46:44Z","title":"Efficient Test-time Inference for Generative Planning Models with OCL Search","version":2},"reference_index":14,"source":"arxiv_source","source_observed_at":"2026-06-28T18:53:08.784890Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2606.00618"},"observation_digest":"sha256:0b392da81b450db7a0409738a1914814edab4acbebb7d5e5d9296c8162152da0","observation_id":"1ea2f4e8-f7a8-4aec-aab2-c024f249b556","resolution":{"observed_at":"2026-06-28T19:52:35.142296Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":"2501.04519","doi":"10.48550/arxiv.2501.04519","metadata_source":"pith","pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","venue":"cs.CL","work_id":"49792b83-569e-4f5f-ae80-e96cbd3b7a43","year":2025},"citing_paper":{"arxiv_id":"2606.03660","last_updated":"2026-06-03T14:05:21Z","snapshot_observed_at":"2026-07-06T23:43:55.632508Z","submitted_at":"2026-06-02T13:47:19Z","title":"From Answers to States: Verifiable Process-Level Evaluation of Chemical Reasoning in Large Language Models","version":2},"reference_index":50,"source":"arxiv_source","source_observed_at":"2026-06-28T09:42:02.298729Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2606.03660"},"observation_digest":"sha256:7a1410a7e622f5d0a0292de2a7c9aa09e8cd52082d817b085584852e59449305","observation_id":"59d40125-828c-4873-971f-b962ae58cf41","resolution":{"observed_at":"2026-07-02T03:46:32.853097Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":"2501.04519","doi":"10.48550/arxiv.2501.04519","metadata_source":"pith","pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","venue":"cs.CL","work_id":"49792b83-569e-4f5f-ae80-e96cbd3b7a43","year":2025},"citing_paper":{"arxiv_id":"2606.05464","last_updated":"2026-06-03T21:43:38Z","snapshot_observed_at":"2026-08-03T10:19:08.677126Z","submitted_at":"2026-06-03T21:43:38Z","title":"Step-by-Step Optimization-like Reasoning in LLMs over Expanding Search Spaces","version":1},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-06-28T05:46:26.938277Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2606.05464"},"observation_digest":"sha256:5ca385df5ceb2c8a91e29cac3c80e50c9272021289f8ba52ea344b2db4bbde78","observation_id":"87c65115-8ac5-42d5-b79b-9fea48415870","resolution":{"observed_at":"2026-07-02T08:46:48.975501Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":"2501.04519","doi":"10.48550/arxiv.2501.04519","metadata_source":"pith","pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","venue":"cs.CL","work_id":"49792b83-569e-4f5f-ae80-e96cbd3b7a43","year":2025},"citing_paper":{"arxiv_id":"2606.06840","last_updated":"2026-06-05T02:32:24Z","snapshot_observed_at":"2026-08-05T15:03:33.373071Z","submitted_at":"2026-06-05T02:32:24Z","title":"Characterize Then Distill: Mechanistic Reasoning in Large Output Spaces","version":1},"reference_index":48,"source":"arxiv_source","source_observed_at":"2026-06-27T22:22:52.690010Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2606.06840"},"observation_digest":"sha256:65cdcac08518401016961b802b414d61d2f858680014a66935d97b3c16ea5e9d","observation_id":"9e8d504b-0622-434f-a03e-94d45c695e4a","resolution":{"observed_at":"2026-06-27T22:31:21.447916Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":"2501.04519","doi":"10.48550/arxiv.2501.04519","metadata_source":"pith","pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","venue":"cs.CL","work_id":"49792b83-569e-4f5f-ae80-e96cbd3b7a43","year":2025},"citing_paper":{"arxiv_id":"2606.07801","last_updated":"2026-06-05T19:32:23Z","snapshot_observed_at":"2026-07-06T23:47:24.424525Z","submitted_at":"2026-06-05T19:32:23Z","title":"Improving Multimodal Reasoning via Worst Dimension Optimization","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-06-27T21:55:57.691185Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2606.07801"},"observation_digest":"sha256:ecfff5b6ab96211606d7f0c432446cecc18ebaa1e4de30ad1d5ee5313eccd43e","observation_id":"90d15e96-d409-4d4a-9b50-c4a9c5f5f5ca","resolution":{"observed_at":"2026-07-02T17:37:14.796624Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":"2501.04519","doi":"10.48550/arxiv.2501.04519","metadata_source":"pith","pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","venue":"cs.CL","work_id":"49792b83-569e-4f5f-ae80-e96cbd3b7a43","year":2025},"citing_paper":{"arxiv_id":"2606.08346","last_updated":"2026-06-06T21:29:01Z","snapshot_observed_at":"2026-08-07T19:43:22.618216Z","submitted_at":"2026-06-06T21:29:01Z","title":"CATPO: Critique-Augmented Tree Policy Optimization","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-07-11T11:50:26.030339Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2606.08346"},"observation_digest":"sha256:99ddb597f7cf44020cee4bcdabf560c81064c8295c109bf34e8fe222b5e40e4a","observation_id":"d65ea05a-f974-467a-a5e5-e3cd492beba8","resolution":{"observed_at":"2026-06-27T19:31:09.966790Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","version":1},"cited_work":{"arxiv_id":"2501.04519","doi":"10.48550/arxiv.2501.04519","metadata_source":"pith","pith_arxiv_id":"2501.04519","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking","venue":"cs.CL","work_id":"49792b83-569e-4f5f-ae80-e96cbd3b7a43","year":2025},"citing_paper":{"arxiv_id":"2606.11470","last_updated":"2026-08-10T19:13:02Z","snapshot_observed_at":"2026-08-12T03:19:39.808979Z","submitted_at":"2026-06-09T21:59:37Z","title":"The Periodic Table of LLM Reasoning: A Structured Survey of Reasoning Paradigms, Methods, and Failure Modes","version":1},"reference_index":82,"source":"arxiv_source","source_observed_at":"2026-06-27T12:59:51.091008Z"},"links":{"cited_paper":"/paper/2501.04519","citing_paper":"/paper/2606.11470"},"observation_digest":"sha256:6da98ed2fdb3bab40fec49ca4335640074cc7aa9d230c5806f90c58f7aa24a1e","observation_id":"ed7f0660-4b5f-487c-bb4d-736115358585","resolution":{"observed_at":"2026-06-27T13:00:56.139993Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-11T06:34:44.6726+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-23T21:53:07.471573+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2501.04519/citation-record","integrity":"/paper/2501.04519/integrity","json":"/paper/2501.04519/citation-record.json","paper":"/paper/2501.04519"},"outbound":[],"paper":{"arxiv_id":"2501.04519","last_updated":"2025-01-08T14:12:57Z","latest_version":1,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-11T09:20:49.313562Z","submitted_at":"2025-01-08T14:12:57Z","title":"rStar-Math: Small LLMs Can Master Math Reasoning with Self-Evolved Deep Thinking"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-11T06:34:44.6726+00:00","source":"crossref"},{"observed_at":"2026-08-11T06:34:36.301508+00:00","source":"retraction_watch"}],"thesis":"As of 12 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 100 inbound Pith citation observations for arXiv:2501.04519."}