{"as_of":"2026-08-08T05:09:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:3d2dfce3f04daae19408bb13709354ef7357fef3e2b97b73d903298d8ae65ec3","coverage":[{"denominator":69,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":69,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-07T10:44:59.510685Z","state":"measured"},{"denominator":71,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":71,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-07T06:34:17.273281+00:00","state":"measured"},{"denominator":2,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":2,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-06T22:17:28.616993Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"pith","source_observed_at":"2026-08-06T15:20:18.946831Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2506.04592","last_updated":"2025-06-05T03:16:08Z","snapshot_observed_at":"2026-08-07T23:35:52.088117Z","submitted_at":"2025-06-05T03:16:08Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2506.04592","snapshot_observed_at":"2026-08-06T22:17:28.616993Z","title":"Safe: Enhancing mathematical reasoning in large language models via retrospective step-aware formal verification, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.22005","last_updated":"2025-06-27T08:17:18Z","snapshot_observed_at":"2026-08-06T22:10:53.633909Z","submitted_at":"2025-06-27T08:17:18Z","title":"LeanConjecturer: Automatic Generation of Mathematical Conjectures for Theorem Proving","version":1},"reference_index":4,"source":"arxiv_source","source_observed_at":"2026-08-06T22:17:28.616993Z"},"links":{"cited_paper":"/paper/2506.04592","citing_paper":"/paper/2506.22005"},"observation_digest":"sha256:7ccbe6f0c6641f6b9049637b50d6c3ba088efc0c53846b5e7ab55ab77ea7475a","observation_id":"2b38c5b5-ca51-4071-9234-64a65a46db49","resolution":{"observed_at":"2026-08-06T22:17:28.616993Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2506.04592","last_updated":"2025-06-05T03:16:08Z","snapshot_observed_at":"2026-08-07T23:35:52.088117Z","submitted_at":"2025-06-05T03:16:08Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification","version":1},"cited_work":{"arxiv_id":"2506.04592","doi":null,"metadata_source":"pith","pith_arxiv_id":"2506.04592","snapshot_observed_at":"2026-08-06T15:20:18.946831Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification","venue":"cs.CL","work_id":"63acbaa4-b53c-4a19-b5aa-ad99e4cde664","year":2025},"citing_paper":{"arxiv_id":"2507.16331","last_updated":"2026-07-03T09:17:03Z","snapshot_observed_at":"2026-08-07T04:15:41.914807Z","submitted_at":"2025-07-22T08:13:01Z","title":"Re:Form -- Reducing Human Annotations in Scalable Formal Software Verification with RL in LLMs: A Preliminary Study on Dafny","version":4},"reference_index":48,"source":"arxiv_source","source_observed_at":"2026-08-06T15:20:14.514001Z"},"links":{"cited_paper":"/paper/2506.04592","citing_paper":"/paper/2507.16331"},"observation_digest":"sha256:7ce7ae2d9ef05046140b2f839f0e3867deca04f2111add03bc184efe131b166b","observation_id":"44aa9e5a-8ce3-4ece-b7bd-f98cf55671ea","resolution":{"observed_at":"2026-08-06T15:20:18.953748Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2506.04592/citation-record","integrity":"/paper/2506.04592/integrity","json":"/paper/2506.04592/citation-record.json","paper":"/paper/2506.04592"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2402.00157","last_updated":"2024-09-16T19:20:59Z","snapshot_observed_at":"2026-08-03T23:38:38.676420Z","submitted_at":"2024-01-31T20:26:32Z","title":"Large Language Models for Mathematical Reasoning: Progresses and Challenges","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.00157","snapshot_observed_at":"2026-08-07T10:44:59.314638Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.04592","last_updated":"2025-06-05T03:16:08Z","snapshot_observed_at":"2026-08-07T23:35:52.088117Z","submitted_at":"2025-06-05T03:16:08Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification","version":1},"reference_index":1,"source":"arxiv_source","source_observed_at":"2026-08-07T10:44:59.314638Z"},"links":{"cited_paper":"/paper/2402.00157","citing_paper":"/paper/2506.04592"},"observation_digest":"sha256:30e8a3cf84df924caec3eaf3ed867f3d92ddfe278eaf3183494e7da95721be34","observation_id":"0a8b4406-5832-45e1-8897-10e849757da2","resolution":{"observed_at":"2026-08-07T10:44:59.314638Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:44:59.964261Z","title":null,"venue":null,"work_id":"23f8e401-d361-4fe7-a35d-1f17a3107b83","year":2021},"citing_paper":{"arxiv_id":"2506.04592","last_updated":"2025-06-05T03:16:08Z","snapshot_observed_at":"2026-08-07T23:35:52.088117Z","submitted_at":"2025-06-05T03:16:08Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification","version":1},"reference_index":2,"source":"arxiv_source","source_observed_at":"2026-08-07T10:44:59.318138Z"},"links":{"citing_paper":"/paper/2506.04592"},"observation_digest":"sha256:8bcd0b46d87f9fc82b4f8b2f7fceaeebf57a766f093bdbed6277b655be865ddf","observation_id":"4ff635bf-b878-4082-823d-39da13267101","resolution":{"observed_at":"2026-08-07T10:44:59.967125Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:44:59.956082Z","title":null,"venue":null,"work_id":"e1640f89-c0f5-4e23-b440-9873fa3dc43c","year":2021},"citing_paper":{"arxiv_id":"2506.04592","last_updated":"2025-06-05T03:16:08Z","snapshot_observed_at":"2026-08-07T23:35:52.088117Z","submitted_at":"2025-06-05T03:16:08Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification","version":1},"reference_index":3,"source":"arxiv_source","source_observed_at":"2026-08-07T10:44:59.322532Z"},"links":{"citing_paper":"/paper/2506.04592"},"observation_digest":"sha256:cb525fd10e192000b1330130618613ff89fcd7d6e23beda0938bbf748c22ab60","observation_id":"c26706cc-b685-4334-8756-2889c066eeac","resolution":{"observed_at":"2026-08-07T10:44:59.958644Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:44:59.325494Z","title":null,"venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2506.04592","last_updated":"2025-06-05T03:16:08Z","snapshot_observed_at":"2026-08-07T23:35:52.088117Z","submitted_at":"2025-06-05T03:16:08Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification","version":1},"reference_index":4,"source":"arxiv_source","source_observed_at":"2026-08-07T10:44:59.325494Z"},"links":{"citing_paper":"/paper/2506.04592"},"observation_digest":"sha256:3f11d9e4a35c6c895eca1476a51180c527dc771e43197af5776d2677fd39b330","observation_id":"497770ac-70fd-4c6f-b22b-c258db0c496f","resolution":{"observed_at":"2026-08-07T10:44:59.325494Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2110.14168","last_updated":"2021-11-18T00:23:45Z","snapshot_observed_at":"2026-08-07T01:45:38.840969Z","submitted_at":"2021-10-27T04:49:45Z","title":"Training Verifiers to Solve Math Word Problems","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2110.14168","snapshot_observed_at":"2026-08-07T10:44:59.328997Z","title":null,"venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2506.04592","last_updated":"2025-06-05T03:16:08Z","snapshot_observed_at":"2026-08-07T23:35:52.088117Z","submitted_at":"2025-06-05T03:16:08Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification","version":1},"reference_index":5,"source":"arxiv_source","source_observed_at":"2026-08-07T10:44:59.328997Z"},"links":{"cited_paper":"/paper/2110.14168","citing_paper":"/paper/2506.04592"},"observation_digest":"sha256:365c1bacce43abd6184299c534241cfd17557b21f121d81b165ecc14638b1d29","observation_id":"b3665208-52fa-4820-92d2-8b40814d69eb","resolution":{"observed_at":"2026-08-07T10:44:59.328997Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2407.21783","last_updated":"2024-11-23T23:27:33Z","snapshot_observed_at":"2026-07-06T18:55:11.576666Z","submitted_at":"2024-07-31T17:54:27Z","title":"The Llama 3 Herd of Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2407.21783","snapshot_observed_at":"2026-08-07T10:44:59.332576Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.04592","last_updated":"2025-06-05T03:16:08Z","snapshot_observed_at":"2026-08-07T23:35:52.088117Z","submitted_at":"2025-06-05T03:16:08Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification","version":1},"reference_index":6,"source":"arxiv_source","source_observed_at":"2026-08-07T10:44:59.332576Z"},"links":{"cited_paper":"/paper/2407.21783","citing_paper":"/paper/2506.04592"},"observation_digest":"sha256:6596b42cfc8c0f195cfbd5aae462b8d8f7fef0d4eba1ae78f751ba545e9003de","observation_id":"7bc8e8dd-c668-4b57-87ca-4c9f2983465a","resolution":{"observed_at":"2026-08-07T10:44:59.332576Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:44:59.944847Z","title":null,"venue":null,"work_id":"cd5acbc0-1feb-40f3-af51-4b602a3d7ed4","year":2024},"citing_paper":{"arxiv_id":"2506.04592","last_updated":"2025-06-05T03:16:08Z","snapshot_observed_at":"2026-08-07T23:35:52.088117Z","submitted_at":"2025-06-05T03:16:08Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification","version":1},"reference_index":7,"source":"arxiv_source","source_observed_at":"2026-08-07T10:44:59.335833Z"},"links":{"citing_paper":"/paper/2506.04592"},"observation_digest":"sha256:fdb55995cf168356162b1d29ba310880d7fd1cef6d469e1d82dd1f72a6a9ddd3","observation_id":"f38c60a5-3741-49ad-a78c-0c371af6bbcd","resolution":{"observed_at":"2026-08-07T10:44:59.947476Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:44:59.936710Z","title":null,"venue":null,"work_id":"1f4e1250-b8a8-4418-912f-6e7ae9c156a9","year":2022},"citing_paper":{"arxiv_id":"2506.04592","last_updated":"2025-06-05T03:16:08Z","snapshot_observed_at":"2026-08-07T23:35:52.088117Z","submitted_at":"2025-06-05T03:16:08Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification","version":1},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-08-07T10:44:59.338852Z"},"links":{"citing_paper":"/paper/2506.04592"},"observation_digest":"sha256:f9c66283e934f73f14958345adf2ae067ba2cb926423534521aa478f99c36d68","observation_id":"4df0577f-4001-468a-b95d-8cb8457a93d4","resolution":{"observed_at":"2026-08-07T10:44:59.939689Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:44:59.341710Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.04592","last_updated":"2025-06-05T03:16:08Z","snapshot_observed_at":"2026-08-07T23:35:52.088117Z","submitted_at":"2025-06-05T03:16:08Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification","version":1},"reference_index":9,"source":"arxiv_source","source_observed_at":"2026-08-07T10:44:59.341710Z"},"links":{"citing_paper":"/paper/2506.04592"},"observation_digest":"sha256:4840a1f40c7aeeee5e2138b85b6f475d38f4660ee2d7b79a0c2c26632452ed4b","observation_id":"ed325208-c5ff-458d-96c5-1494a38d43a2","resolution":{"observed_at":"2026-08-07T10:44:59.341710Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:44:59.924860Z","title":null,"venue":null,"work_id":"9fea3f99-c064-4d1f-9d1a-3f22b60353f4","year":2024},"citing_paper":{"arxiv_id":"2506.04592","last_updated":"2025-06-05T03:16:08Z","snapshot_observed_at":"2026-08-07T23:35:52.088117Z","submitted_at":"2025-06-05T03:16:08Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification","version":1},"reference_index":10,"source":"arxiv_source","source_observed_at":"2026-08-07T10:44:59.344582Z"},"links":{"citing_paper":"/paper/2506.04592"},"observation_digest":"sha256:a9e17da70aed7ceed82bbf8acf22244475cffb5924a0930ab28f8b8b30110b12","observation_id":"87723558-d3f0-4d3a-aa07-54c58eae4073","resolution":{"observed_at":"2026-08-07T10:44:59.928164Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:44:59.347562Z","title":null,"venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2506.04592","last_updated":"2025-06-05T03:16:08Z","snapshot_observed_at":"2026-08-07T23:35:52.088117Z","submitted_at":"2025-06-05T03:16:08Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification","version":1},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-08-07T10:44:59.347562Z"},"links":{"citing_paper":"/paper/2506.04592"},"observation_digest":"sha256:789f790c1b9dede10a7f9b2fad70d0299d3981356a81ac08c47e661af497d376","observation_id":"39ff0bcf-742b-43bb-91b2-f4294e0a6195","resolution":{"observed_at":"2026-08-07T10:44:59.347562Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:44:59.913242Z","title":null,"venue":null,"work_id":"ac379cf7-3e02-42eb-9fe8-040276671cb8","year":1997},"citing_paper":{"arxiv_id":"2506.04592","last_updated":"2025-06-05T03:16:08Z","snapshot_observed_at":"2026-08-07T23:35:52.088117Z","submitted_at":"2025-06-05T03:16:08Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification","version":1},"reference_index":12,"source":"arxiv_source","source_observed_at":"2026-08-07T10:44:59.350611Z"},"links":{"citing_paper":"/paper/2506.04592"},"observation_digest":"sha256:f0c96b1541b7b17cedec3282ca56de717dbea8c03c80459fe070ded8b269ba04","observation_id":"ca38ed18-8ed7-46b5-ac4c-99993534495d","resolution":{"observed_at":"2026-08-07T10:44:59.915939Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:44:59.353232Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.04592","last_updated":"2025-06-05T03:16:08Z","snapshot_observed_at":"2026-08-07T23:35:52.088117Z","submitted_at":"2025-06-05T03:16:08Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification","version":1},"reference_index":13,"source":"arxiv_source","source_observed_at":"2026-08-07T10:44:59.353232Z"},"links":{"citing_paper":"/paper/2506.04592"},"observation_digest":"sha256:f466e0c608ea37eb3945befd34a7d8b4fc73dcc376ca73c5e5a48d684e11df09","observation_id":"80b99761-2363-40cc-8748-1e7f8a120661","resolution":{"observed_at":"2026-08-07T10:44:59.353232Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:44:59.902510Z","title":null,"venue":null,"work_id":"6ea2c965-17a5-4b3a-b356-befecbb09a16","year":2024},"citing_paper":{"arxiv_id":"2506.04592","last_updated":"2025-06-05T03:16:08Z","snapshot_observed_at":"2026-08-07T23:35:52.088117Z","submitted_at":"2025-06-05T03:16:08Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification","version":1},"reference_index":14,"source":"arxiv_source","source_observed_at":"2026-08-07T10:44:59.355900Z"},"links":{"citing_paper":"/paper/2506.04592"},"observation_digest":"sha256:13dd07c0903c9910ddf1b0a21a39bcf6061f3221764a3246d52ce8fc16fb258e","observation_id":"ed0ae6dc-8e83-49e2-9431-de8b1e60725f","resolution":{"observed_at":"2026-08-07T10:44:59.905116Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.21276","last_updated":"2024-10-25T17:43:01Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2024-10-25T17:43:01Z","title":"GPT-4o System Card","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.21276","snapshot_observed_at":"2026-08-07T10:44:59.358492Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.04592","last_updated":"2025-06-05T03:16:08Z","snapshot_observed_at":"2026-08-07T23:35:52.088117Z","submitted_at":"2025-06-05T03:16:08Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification","version":1},"reference_index":15,"source":"arxiv_source","source_observed_at":"2026-08-07T10:44:59.358492Z"},"links":{"cited_paper":"/paper/2410.21276","citing_paper":"/paper/2506.04592"},"observation_digest":"sha256:0e9fa54b1ce796fedea87ffbcbd7be39a6cb1e5e04dcb9d8ed21d15dcff72a59","observation_id":"f4700c45-62f0-4c33-a081-5873e1b4ab45","resolution":{"observed_at":"2026-08-07T10:44:59.358492Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:44:59.894378Z","title":null,"venue":null,"work_id":"3ecf9e2b-a9bc-4824-8619-019c4e3fbfb3","year":2022},"citing_paper":{"arxiv_id":"2506.04592","last_updated":"2025-06-05T03:16:08Z","snapshot_observed_at":"2026-08-07T23:35:52.088117Z","submitted_at":"2025-06-05T03:16:08Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification","version":1},"reference_index":16,"source":"arxiv_source","source_observed_at":"2026-08-07T10:44:59.361274Z"},"links":{"citing_paper":"/paper/2506.04592"},"observation_digest":"sha256:7fba4cd476c332a90d26f12edef71c636dd7b94f24c40eb2eb7d1b564af240ed","observation_id":"81485b53-4e8b-4c58-9bd6-58a5cc4fa8a2","resolution":{"observed_at":"2026-08-07T10:44:59.897276Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:44:59.886892Z","title":null,"venue":null,"work_id":"5cee0f2d-8738-4642-a2cd-256dd01c9bbb","year":2024},"citing_paper":{"arxiv_id":"2506.04592","last_updated":"2025-06-05T03:16:08Z","snapshot_observed_at":"2026-08-07T23:35:52.088117Z","submitted_at":"2025-06-05T03:16:08Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification","version":1},"reference_index":17,"source":"arxiv_source","source_observed_at":"2026-08-07T10:44:59.363852Z"},"links":{"citing_paper":"/paper/2506.04592"},"observation_digest":"sha256:aa2a56f2fc529f737261cfcdff79967317cc904a24e59c25c8b74ea8a2d44c95","observation_id":"59386955-86c9-4bc4-bbc1-45b74d453371","resolution":{"observed_at":"2026-08-07T10:44:59.890093Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:44:59.366523Z","title":null,"venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2506.04592","last_updated":"2025-06-05T03:16:08Z","snapshot_observed_at":"2026-08-07T23:35:52.088117Z","submitted_at":"2025-06-05T03:16:08Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification","version":1},"reference_index":18,"source":"arxiv_source","source_observed_at":"2026-08-07T10:44:59.366523Z"},"links":{"citing_paper":"/paper/2506.04592"},"observation_digest":"sha256:fb8f7a60a31eb518e6667a4c9ed8582e2f8688183a02f63d7170eca2bbf0f62d","observation_id":"4a9d1346-d436-4a62-86b2-424fcf092b1f","resolution":{"observed_at":"2026-08-07T10:44:59.366523Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:44:59.874219Z","title":null,"venue":null,"work_id":"1b48250f-1ae1-4c3e-871a-5c421aa97781","year":2020},"citing_paper":{"arxiv_id":"2506.04592","last_updated":"2025-06-05T03:16:08Z","snapshot_observed_at":"2026-08-07T23:35:52.088117Z","submitted_at":"2025-06-05T03:16:08Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification","version":1},"reference_index":19,"source":"arxiv_source","source_observed_at":"2026-08-07T10:44:59.369350Z"},"links":{"citing_paper":"/paper/2506.04592"},"observation_digest":"sha256:889d1abc835096df1bd6e05634cf5f53c2db677acd6e41e124eb1edde67a5e6d","observation_id":"2cb400b4-8036-43ca-a2ea-2611bf5d6388","resolution":{"observed_at":"2026-08-07T10:44:59.877589Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:44:59.372876Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.04592","last_updated":"2025-06-05T03:16:08Z","snapshot_observed_at":"2026-08-07T23:35:52.088117Z","submitted_at":"2025-06-05T03:16:08Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification","version":1},"reference_index":20,"source":"arxiv_source","source_observed_at":"2026-08-07T10:44:59.372876Z"},"links":{"citing_paper":"/paper/2506.04592"},"observation_digest":"sha256:49a645acd9569447dae909b32462eb3a9f2dcd963b73bf66b42a900522349d9e","observation_id":"5717639a-b226-448f-af52-38c1bcc5855b","resolution":{"observed_at":"2026-08-07T10:44:59.372876Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:44:59.862698Z","title":null,"venue":null,"work_id":"a5de2e3b-c8d5-4009-bbf3-6893f0bb3091","year":2024},"citing_paper":{"arxiv_id":"2506.04592","last_updated":"2025-06-05T03:16:08Z","snapshot_observed_at":"2026-08-07T23:35:52.088117Z","submitted_at":"2025-06-05T03:16:08Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification","version":1},"reference_index":21,"source":"arxiv_source","source_observed_at":"2026-08-07T10:44:59.375796Z"},"links":{"citing_paper":"/paper/2506.04592"},"observation_digest":"sha256:45da539fb42475f562ae3809e04c77aec740ee424c8ec6233d40ef7d908a8cb6","observation_id":"0435f491-9514-4b92-bdcd-52633e0415b5","resolution":{"observed_at":"2026-08-07T10:44:59.865476Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2403.13787","last_updated":"2024-06-08T16:40:12Z","snapshot_observed_at":"2026-08-02T18:11:57.036767Z","submitted_at":"2024-03-20T17:49:54Z","title":"RewardBench: Evaluating Reward Models for Language Modeling","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.13787","snapshot_observed_at":"2026-08-07T10:44:59.378386Z","title":"Smith, and Hannaneh Hajishirzi","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.04592","last_updated":"2025-06-05T03:16:08Z","snapshot_observed_at":"2026-08-07T23:35:52.088117Z","submitted_at":"2025-06-05T03:16:08Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification","version":1},"reference_index":22,"source":"arxiv_source","source_observed_at":"2026-08-07T10:44:59.378386Z"},"links":{"cited_paper":"/paper/2403.13787","citing_paper":"/paper/2506.04592"},"observation_digest":"sha256:f53e155a7d7995578d9f00663c2097e73154a3ff1f285b3dafba315cef7eb021","observation_id":"96f48841-4829-40d3-bdba-7726951f6802","resolution":{"observed_at":"2026-08-07T10:44:59.378386Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:44:59.854593Z","title":null,"venue":null,"work_id":"ef4c44c3-64e6-4c82-b078-4bd3e45745ae","year":2022},"citing_paper":{"arxiv_id":"2506.04592","last_updated":"2025-06-05T03:16:08Z","snapshot_observed_at":"2026-08-07T23:35:52.088117Z","submitted_at":"2025-06-05T03:16:08Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification","version":1},"reference_index":23,"source":"arxiv_source","source_observed_at":"2026-08-07T10:44:59.381405Z"},"links":{"citing_paper":"/paper/2506.04592"},"observation_digest":"sha256:43b8b405757cb14d81abace30a236aba49c458899c42d33d75521ec3cf84b9ca","observation_id":"3ea1986f-e7d7-4531-b341-cbf7ca3fdef4","resolution":{"observed_at":"2026-08-07T10:44:59.858080Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2404.00599","last_updated":"2024-03-31T08:10:50Z","snapshot_observed_at":"2026-08-07T19:41:09.683401Z","submitted_at":"2024-03-31T08:10:50Z","title":"EvoCodeBench: An Evolving Code Generation Benchmark Aligned with Real-World Code Repositories","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.00599","snapshot_observed_at":"2026-08-07T10:44:59.384125Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.04592","last_updated":"2025-06-05T03:16:08Z","snapshot_observed_at":"2026-08-07T23:35:52.088117Z","submitted_at":"2025-06-05T03:16:08Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification","version":1},"reference_index":24,"source":"arxiv_source","source_observed_at":"2026-08-07T10:44:59.384125Z"},"links":{"cited_paper":"/paper/2404.00599","citing_paper":"/paper/2506.04592"},"observation_digest":"sha256:00292d33b138946bf7558176bf2479acc822e2cd241e9eefbce14e41d1d681d4","observation_id":"38b8e995-9922-44d7-9064-ba4401ed6059","resolution":{"observed_at":"2026-08-07T10:44:59.384125Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:44:59.387846Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.04592","last_updated":"2025-06-05T03:16:08Z","snapshot_observed_at":"2026-08-07T23:35:52.088117Z","submitted_at":"2025-06-05T03:16:08Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification","version":1},"reference_index":25,"source":"arxiv_source","source_observed_at":"2026-08-07T10:44:59.387846Z"},"links":{"citing_paper":"/paper/2506.04592"},"observation_digest":"sha256:2441378fce496744714c5e32dcb61848bbd1db9d7ea0153a2ce2ebec420a78ed","observation_id":"fb43529e-c24f-49f8-8dd5-9f0a1e47bab7","resolution":{"observed_at":"2026-08-07T10:44:59.387846Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2309.04295","last_updated":"2023-12-05T08:38:01Z","snapshot_observed_at":"2026-08-07T23:35:18.737859Z","submitted_at":"2023-09-08T12:34:28Z","title":"FIMO: A Challenge Formal Dataset for Automated Theorem Proving","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2309.04295","snapshot_observed_at":"2026-08-07T10:44:59.390584Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.04592","last_updated":"2025-06-05T03:16:08Z","snapshot_observed_at":"2026-08-07T23:35:52.088117Z","submitted_at":"2025-06-05T03:16:08Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification","version":1},"reference_index":26,"source":"arxiv_source","source_observed_at":"2026-08-07T10:44:59.390584Z"},"links":{"cited_paper":"/paper/2309.04295","citing_paper":"/paper/2506.04592"},"observation_digest":"sha256:78b7a5ec02bee8883570328ce83a3c4140159af23f02671b12077b3b153d8035","observation_id":"026ff056-7a23-492a-836d-a5406c650fee","resolution":{"observed_at":"2026-08-07T10:44:59.390584Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2410.18451","last_updated":"2024-10-24T06:06:26Z","snapshot_observed_at":"2026-08-07T05:53:12.643515Z","submitted_at":"2024-10-24T06:06:26Z","title":"Skywork-Reward: Bag of Tricks for Reward Modeling in LLMs","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.18451","snapshot_observed_at":"2026-08-07T10:44:59.393457Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.04592","last_updated":"2025-06-05T03:16:08Z","snapshot_observed_at":"2026-08-07T23:35:52.088117Z","submitted_at":"2025-06-05T03:16:08Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification","version":1},"reference_index":27,"source":"arxiv_source","source_observed_at":"2026-08-07T10:44:59.393457Z"},"links":{"cited_paper":"/paper/2410.18451","citing_paper":"/paper/2506.04592"},"observation_digest":"sha256:1ba674d0440d8771b501480462a6455fcf8b30fc4b20eb15c879fe1dc82347e1","observation_id":"6193df7e-5dc7-4567-90a4-ff73a47eb869","resolution":{"observed_at":"2026-08-07T10:44:59.393457Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.01940","last_updated":"2024-10-14T03:46:00Z","snapshot_observed_at":"2026-08-05T18:55:10.658942Z","submitted_at":"2024-06-04T03:48:08Z","title":"Process-Driven Autoformalization in Lean 4","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.01940","snapshot_observed_at":"2026-08-07T10:44:59.396198Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.04592","last_updated":"2025-06-05T03:16:08Z","snapshot_observed_at":"2026-08-07T23:35:52.088117Z","submitted_at":"2025-06-05T03:16:08Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification","version":1},"reference_index":28,"source":"arxiv_source","source_observed_at":"2026-08-07T10:44:59.396198Z"},"links":{"cited_paper":"/paper/2406.01940","citing_paper":"/paper/2506.04592"},"observation_digest":"sha256:a33463141f103dd573ef40232b73a00d39fe5d2ae45e4f8af47f021ef8b2b093","observation_id":"168df49c-4963-4de0-9ca2-46f2bacc2e9d","resolution":{"observed_at":"2026-08-07T10:44:59.396198Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:44:59.399512Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.04592","last_updated":"2025-06-05T03:16:08Z","snapshot_observed_at":"2026-08-07T23:35:52.088117Z","submitted_at":"2025-06-05T03:16:08Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification","version":1},"reference_index":29,"source":"arxiv_source","source_observed_at":"2026-08-07T10:44:59.399512Z"},"links":{"citing_paper":"/paper/2506.04592"},"observation_digest":"sha256:dcc88d07689ca1a222cd1408ec496e8d6a3cde31e621f3538c571da599846c3b","observation_id":"961f0443-416f-4425-92d8-b51cf3dd2027","resolution":{"observed_at":"2026-08-07T10:44:59.399512Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1910.09336","last_updated":"2019-12-16T20:37:52Z","snapshot_observed_at":"2026-07-06T08:31:00.030545Z","submitted_at":"2019-10-21T13:08:24Z","title":"The Lean mathematical library","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1910.09336","snapshot_observed_at":"2026-08-07T10:44:59.402345Z","title":null,"venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2506.04592","last_updated":"2025-06-05T03:16:08Z","snapshot_observed_at":"2026-08-07T23:35:52.088117Z","submitted_at":"2025-06-05T03:16:08Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification","version":1},"reference_index":30,"source":"arxiv_source","source_observed_at":"2026-08-07T10:44:59.402345Z"},"links":{"cited_paper":"/paper/1910.09336","citing_paper":"/paper/2506.04592"},"observation_digest":"sha256:9fc273c6dc378c7c0896c72a95b3c711b15c5cb3fa48acf53814caad59bc5474","observation_id":"df9bded4-17a9-478c-969b-f7b2d77c37bb","resolution":{"observed_at":"2026-08-07T10:44:59.402345Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:44:59.405311Z","title":null,"venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2506.04592","last_updated":"2025-06-05T03:16:08Z","snapshot_observed_at":"2026-08-07T23:35:52.088117Z","submitted_at":"2025-06-05T03:16:08Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification","version":1},"reference_index":31,"source":"arxiv_source","source_observed_at":"2026-08-07T10:44:59.405311Z"},"links":{"citing_paper":"/paper/2506.04592"},"observation_digest":"sha256:63e477eee8aa37a5b003cbec4f5145131cd3e2ba2b78b544b2ba3a0bfc3a2293","observation_id":"0101a69b-0713-480e-b43d-0ae6dd539cf6","resolution":{"observed_at":"2026-08-07T10:44:59.405311Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:44:59.839837Z","title":null,"venue":null,"work_id":"a0fe4d0f-b16f-42f2-90bd-9ff87ed98fde","year":2023},"citing_paper":{"arxiv_id":"2506.04592","last_updated":"2025-06-05T03:16:08Z","snapshot_observed_at":"2026-08-07T23:35:52.088117Z","submitted_at":"2025-06-05T03:16:08Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification","version":1},"reference_index":32,"source":"arxiv_source","source_observed_at":"2026-08-07T10:44:59.407975Z"},"links":{"citing_paper":"/paper/2506.04592"},"observation_digest":"sha256:9199d9833c29a4909852a596da430b4cccd6c1ba33edb2e6ee5a20591fd4598f","observation_id":"a5511e95-4098-49ec-90d0-271f5e5c273b","resolution":{"observed_at":"2026-08-07T10:44:59.842585Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:44:59.832300Z","title":null,"venue":null,"work_id":"4c849a84-1f85-4051-9c8c-f25be14c51ad","year":2023},"citing_paper":{"arxiv_id":"2506.04592","last_updated":"2025-06-05T03:16:08Z","snapshot_observed_at":"2026-08-07T23:35:52.088117Z","submitted_at":"2025-06-05T03:16:08Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification","version":1},"reference_index":33,"source":"arxiv_source","source_observed_at":"2026-08-07T10:44:59.410896Z"},"links":{"citing_paper":"/paper/2506.04592"},"observation_digest":"sha256:9f4d0e5ca7e878733bdf3d100f1d599b37b045ee5740a28f83e551ed2c6fbd65","observation_id":"ddb8de31-193a-4def-9166-1e14bbfd540c","resolution":{"observed_at":"2026-08-07T10:44:59.835296Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:44:59.824654Z","title":null,"venue":null,"work_id":"37ae90fc-ebce-4f6b-9244-6ea6ebb4fe52","year":2022},"citing_paper":{"arxiv_id":"2506.04592","last_updated":"2025-06-05T03:16:08Z","snapshot_observed_at":"2026-08-07T23:35:52.088117Z","submitted_at":"2025-06-05T03:16:08Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification","version":1},"reference_index":34,"source":"arxiv_source","source_observed_at":"2026-08-07T10:44:59.413763Z"},"links":{"citing_paper":"/paper/2506.04592"},"observation_digest":"sha256:f3b6fc70fe58a6cc970e580aaa2806f3ebe683313f18e0ac3c690ada970d2b61","observation_id":"938a2745-6525-4ab7-be01-3f64e617620f","resolution":{"observed_at":"2026-08-07T10:44:59.827195Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2009.03393","last_updated":"2020-09-07T19:50:10Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2020-09-07T19:50:10Z","title":"Generative Language Modeling for Automated Theorem Proving","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2009.03393","snapshot_observed_at":"2026-08-07T10:44:59.417804Z","title":null,"venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2506.04592","last_updated":"2025-06-05T03:16:08Z","snapshot_observed_at":"2026-08-07T23:35:52.088117Z","submitted_at":"2025-06-05T03:16:08Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification","version":1},"reference_index":35,"source":"arxiv_source","source_observed_at":"2026-08-07T10:44:59.417804Z"},"links":{"cited_paper":"/paper/2009.03393","citing_paper":"/paper/2506.04592"},"observation_digest":"sha256:75303c4d0e5a61c164d524a83bcb32e65137d7e1f7c14118d83ebcd69f3a810e","observation_id":"813c056d-6353-49bf-8115-256699197efa","resolution":{"observed_at":"2026-08-07T10:44:59.417804Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2311.12022","last_updated":"2023-11-20T18:57:34Z","snapshot_observed_at":"2026-08-04T22:55:15.345443Z","submitted_at":"2023-11-20T18:57:34Z","title":"GPQA: A Graduate-Level Google-Proof Q&A Benchmark","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2311.12022","snapshot_observed_at":"2026-08-07T10:44:59.420494Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.04592","last_updated":"2025-06-05T03:16:08Z","snapshot_observed_at":"2026-08-07T23:35:52.088117Z","submitted_at":"2025-06-05T03:16:08Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification","version":1},"reference_index":36,"source":"arxiv_source","source_observed_at":"2026-08-07T10:44:59.420494Z"},"links":{"cited_paper":"/paper/2311.12022","citing_paper":"/paper/2506.04592"},"observation_digest":"sha256:f5a835c85cd242035adc7ca633003269bdb7675877ab3f6b9ad1132234018eab","observation_id":"f6e26455-bf57-4a9c-baf8-9fe4999d157d","resolution":{"observed_at":"2026-08-07T10:44:59.420494Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:44:59.817241Z","title":null,"venue":null,"work_id":"a78ac7f6-df22-428b-af07-cbcecc036438","year":2022},"citing_paper":{"arxiv_id":"2506.04592","last_updated":"2025-06-05T03:16:08Z","snapshot_observed_at":"2026-08-07T23:35:52.088117Z","submitted_at":"2025-06-05T03:16:08Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification","version":1},"reference_index":37,"source":"arxiv_source","source_observed_at":"2026-08-07T10:44:59.423051Z"},"links":{"citing_paper":"/paper/2506.04592"},"observation_digest":"sha256:8303eb4513ba92732f852bbe2bf8746ed410640e78d8623c86b435fdb90f7554","observation_id":"4e3de15b-d59b-4b99-83ba-ca9bc7ad5a00","resolution":{"observed_at":"2026-08-07T10:44:59.820122Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:44:59.810033Z","title":null,"venue":null,"work_id":"30f7fbf3-1773-4444-9188-a4b46ff4125c","year":2024},"citing_paper":{"arxiv_id":"2506.04592","last_updated":"2025-06-05T03:16:08Z","snapshot_observed_at":"2026-08-07T23:35:52.088117Z","submitted_at":"2025-06-05T03:16:08Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification","version":1},"reference_index":38,"source":"arxiv_source","source_observed_at":"2026-08-07T10:44:59.425505Z"},"links":{"citing_paper":"/paper/2506.04592"},"observation_digest":"sha256:f7280ab76135be510dd5ef381d12306d65f2ca251e3c6beb1aab5baf265cc213","observation_id":"fbb76c3b-acac-4d4a-b03c-18e5a547bc59","resolution":{"observed_at":"2026-08-07T10:44:59.812590Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2402.03300","last_updated":"2024-04-27T15:25:53Z","snapshot_observed_at":"2026-08-06T14:58:42.911363Z","submitted_at":"2024-02-05T18:55:32Z","title":"DeepSeekMath: Pushing the Limits of Mathematical Reasoning in Open Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.03300","snapshot_observed_at":"2026-08-07T10:44:59.428476Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.04592","last_updated":"2025-06-05T03:16:08Z","snapshot_observed_at":"2026-08-07T23:35:52.088117Z","submitted_at":"2025-06-05T03:16:08Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification","version":1},"reference_index":39,"source":"arxiv_source","source_observed_at":"2026-08-07T10:44:59.428476Z"},"links":{"cited_paper":"/paper/2402.03300","citing_paper":"/paper/2506.04592"},"observation_digest":"sha256:8d34e0c44d1df505007c91f200566ca637a2804659ffbcf73d5e5fa92162b519","observation_id":"215def12-0c2f-4eae-a8b0-ec0e9144abf2","resolution":{"observed_at":"2026-08-07T10:44:59.428476Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.18653/v1/2021.findings-emnlp.195","metadata_source":"doi_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:44:59.543280Z","title":null,"venue":null,"work_id":"3992c0b7-f124-46f4-bf51-146a48e8438c","year":2021},"citing_paper":{"arxiv_id":"2506.04592","last_updated":"2025-06-05T03:16:08Z","snapshot_observed_at":"2026-08-07T23:35:52.088117Z","submitted_at":"2025-06-05T03:16:08Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification","version":1},"reference_index":40,"source":"arxiv_source","source_observed_at":"2026-08-07T10:44:59.431213Z"},"links":{"citing_paper":"/paper/2506.04592"},"observation_digest":"sha256:239297c507ba8891f203d619940aa9fb5e60536160d32811aa99c40e9e07c142","observation_id":"ce194bc3-da9b-4379-8ce0-ed78990b6f72","resolution":{"observed_at":"2026-08-07T10:44:59.547381Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:44:59.802681Z","title":null,"venue":null,"work_id":"255808b3-0d63-47c1-af62-274a632c0460","year":2024},"citing_paper":{"arxiv_id":"2506.04592","last_updated":"2025-06-05T03:16:08Z","snapshot_observed_at":"2026-08-07T23:35:52.088117Z","submitted_at":"2025-06-05T03:16:08Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification","version":1},"reference_index":41,"source":"arxiv_source","source_observed_at":"2026-08-07T10:44:59.433959Z"},"links":{"citing_paper":"/paper/2506.04592"},"observation_digest":"sha256:9e2d8be55faaaae43e1c9690ae640852cdf505f58fdf6772f3f93515b193e769","observation_id":"08a113e9-518e-444d-8ae3-9e1817b1ae8b","resolution":{"observed_at":"2026-08-07T10:44:59.805639Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:44:59.795790Z","title":null,"venue":null,"work_id":"eeadb995-58c5-4257-849a-3f4317f3a8a6","year":2024},"citing_paper":{"arxiv_id":"2506.04592","last_updated":"2025-06-05T03:16:08Z","snapshot_observed_at":"2026-08-07T23:35:52.088117Z","submitted_at":"2025-06-05T03:16:08Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification","version":1},"reference_index":42,"source":"arxiv_source","source_observed_at":"2026-08-07T10:44:59.436651Z"},"links":{"citing_paper":"/paper/2506.04592"},"observation_digest":"sha256:f5b020a350b4ce97c953677f696b08fa1d49c498adf9eeca393bb00846e9f902","observation_id":"c40e29e1-26e5-47b1-8ab9-34158c2c48ef","resolution":{"observed_at":"2026-08-07T10:44:59.798289Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:44:59.788030Z","title":null,"venue":null,"work_id":"77d178c9-165e-4e75-807e-968fee47bdee","year":2023},"citing_paper":{"arxiv_id":"2506.04592","last_updated":"2025-06-05T03:16:08Z","snapshot_observed_at":"2026-08-07T23:35:52.088117Z","submitted_at":"2025-06-05T03:16:08Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification","version":1},"reference_index":43,"source":"arxiv_source","source_observed_at":"2026-08-07T10:44:59.440156Z"},"links":{"citing_paper":"/paper/2506.04592"},"observation_digest":"sha256:de5a5488efbeff57b72278250cb033cb559e76e9486f1408c83653405624e8b6","observation_id":"ee564e06-07ce-43a8-a95c-c3870aad75cb","resolution":{"observed_at":"2026-08-07T10:44:59.790571Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:44:59.442685Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.04592","last_updated":"2025-06-05T03:16:08Z","snapshot_observed_at":"2026-08-07T23:35:52.088117Z","submitted_at":"2025-06-05T03:16:08Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification","version":1},"reference_index":44,"source":"arxiv_source","source_observed_at":"2026-08-07T10:44:59.442685Z"},"links":{"citing_paper":"/paper/2506.04592"},"observation_digest":"sha256:78ff0cdf7dfe6c15a73e7a897b1a342266cf00a224132eaa2f5e0732526eb8e9","observation_id":"1065ed31-53de-4fb0-8668-4b06dc655174","resolution":{"observed_at":"2026-08-07T10:44:59.442685Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:44:59.780190Z","title":null,"venue":null,"work_id":"0e6ef8cd-10e0-4bfd-b1e1-fa452e620eda","year":2024},"citing_paper":{"arxiv_id":"2506.04592","last_updated":"2025-06-05T03:16:08Z","snapshot_observed_at":"2026-08-07T23:35:52.088117Z","submitted_at":"2025-06-05T03:16:08Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification","version":1},"reference_index":45,"source":"arxiv_source","source_observed_at":"2026-08-07T10:44:59.445486Z"},"links":{"citing_paper":"/paper/2506.04592"},"observation_digest":"sha256:b395791849167888a6771ff98a91237f521b6b4402a7aaca6b1ef604f5711eba","observation_id":"a87bf1d8-4e0b-48c5-8295-fcd5a172eec4","resolution":{"observed_at":"2026-08-07T10:44:59.782908Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:44:59.448037Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.04592","last_updated":"2025-06-05T03:16:08Z","snapshot_observed_at":"2026-08-07T23:35:52.088117Z","submitted_at":"2025-06-05T03:16:08Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification","version":1},"reference_index":46,"source":"arxiv_source","source_observed_at":"2026-08-07T10:44:59.448037Z"},"links":{"citing_paper":"/paper/2506.04592"},"observation_digest":"sha256:11ad02f6ac8d2e9fa7c8a453e914c587400c050d17bfcb18cd4808330cd4549f","observation_id":"87674b87-b5dd-42ef-b8dc-140112621262","resolution":{"observed_at":"2026-08-07T10:44:59.448037Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2307.10635","last_updated":"2024-06-28T08:24:13Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-07-20T07:01:57Z","title":"SciBench: Evaluating College-Level Scientific Problem-Solving Abilities of Large Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2307.10635","snapshot_observed_at":"2026-08-07T10:44:59.450839Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.04592","last_updated":"2025-06-05T03:16:08Z","snapshot_observed_at":"2026-08-07T23:35:52.088117Z","submitted_at":"2025-06-05T03:16:08Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification","version":1},"reference_index":47,"source":"arxiv_source","source_observed_at":"2026-08-07T10:44:59.450839Z"},"links":{"cited_paper":"/paper/2307.10635","citing_paper":"/paper/2506.04592"},"observation_digest":"sha256:40189dc4d5f02947ebd784db9a9d90d7fcf26617bee2fbe2abb798b6d040e0fc","observation_id":"bf21d658-604d-4752-af9e-01bf920e9d7d","resolution":{"observed_at":"2026-08-07T10:44:59.450839Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:44:59.453738Z","title":"Chi, Quoc V Le, and Denny Zhou","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.04592","last_updated":"2025-06-05T03:16:08Z","snapshot_observed_at":"2026-08-07T23:35:52.088117Z","submitted_at":"2025-06-05T03:16:08Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification","version":1},"reference_index":48,"source":"arxiv_source","source_observed_at":"2026-08-07T10:44:59.453738Z"},"links":{"citing_paper":"/paper/2506.04592"},"observation_digest":"sha256:df4347c5ee271f2ece250be0eda7b1cde364e72223b0a2f1121d2641989b9239","observation_id":"f64dfd52-1a97-4a55-a69d-89c30c73eec2","resolution":{"observed_at":"2026-08-07T10:44:59.453738Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:44:59.768460Z","title":null,"venue":null,"work_id":"f2da9838-ff42-4483-80b2-cd0c9c9d97b7","year":2022},"citing_paper":{"arxiv_id":"2506.04592","last_updated":"2025-06-05T03:16:08Z","snapshot_observed_at":"2026-08-07T23:35:52.088117Z","submitted_at":"2025-06-05T03:16:08Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification","version":1},"reference_index":49,"source":"arxiv_source","source_observed_at":"2026-08-07T10:44:59.456378Z"},"links":{"citing_paper":"/paper/2506.04592"},"observation_digest":"sha256:c0f65482c9ca9fba29658dc63847471b1cb1128b301e96897240e6a36c488bbe","observation_id":"c6d941dd-0837-4bcb-8a3d-a9f192f22956","resolution":{"observed_at":"2026-08-07T10:44:59.771006Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2405.14333","last_updated":"2024-05-23T09:03:42Z","snapshot_observed_at":"2026-07-06T18:18:25.746022Z","submitted_at":"2024-05-23T09:03:42Z","title":"DeepSeek-Prover: Advancing Theorem Proving in LLMs through Large-Scale Synthetic Data","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2405.14333","snapshot_observed_at":"2026-08-07T10:44:59.459139Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.04592","last_updated":"2025-06-05T03:16:08Z","snapshot_observed_at":"2026-08-07T23:35:52.088117Z","submitted_at":"2025-06-05T03:16:08Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification","version":1},"reference_index":50,"source":"arxiv_source","source_observed_at":"2026-08-07T10:44:59.459139Z"},"links":{"cited_paper":"/paper/2405.14333","citing_paper":"/paper/2506.04592"},"observation_digest":"sha256:f5d092d756779988c97d1ec631597472398aa3a6921b4a769cafe521ca59d211","observation_id":"684c83d8-729a-41f3-89c1-e44ebaebea77","resolution":{"observed_at":"2026-08-07T10:44:59.459139Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:44:59.761410Z","title":null,"venue":null,"work_id":"9a24defd-9e41-4193-972f-c9a8c4112f74","year":2024},"citing_paper":{"arxiv_id":"2506.04592","last_updated":"2025-06-05T03:16:08Z","snapshot_observed_at":"2026-08-07T23:35:52.088117Z","submitted_at":"2025-06-05T03:16:08Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification","version":1},"reference_index":51,"source":"arxiv_source","source_observed_at":"2026-08-07T10:44:59.463401Z"},"links":{"citing_paper":"/paper/2506.04592"},"observation_digest":"sha256:53be80ab4dd2063b31e0a95df50ba17a7fb1a45d257df5964df5c3ddc40f9f11","observation_id":"caa98479-512a-418f-8c4e-74924e71f9cf","resolution":{"observed_at":"2026-08-07T10:44:59.764145Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:44:59.753960Z","title":null,"venue":null,"work_id":"8f5fd378-b336-48a8-99a0-f24c9968aff6","year":2023},"citing_paper":{"arxiv_id":"2506.04592","last_updated":"2025-06-05T03:16:08Z","snapshot_observed_at":"2026-08-07T23:35:52.088117Z","submitted_at":"2025-06-05T03:16:08Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification","version":1},"reference_index":52,"source":"arxiv_source","source_observed_at":"2026-08-07T10:44:59.466077Z"},"links":{"citing_paper":"/paper/2506.04592"},"observation_digest":"sha256:c9dcb20f527854873cdafdb9388ec08436c4480e51b58a295266b0632bfa971e","observation_id":"c3280eb4-68f2-4eed-a229-0980d1d7d322","resolution":{"observed_at":"2026-08-07T10:44:59.757003Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:44:59.468639Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.04592","last_updated":"2025-06-05T03:16:08Z","snapshot_observed_at":"2026-08-07T23:35:52.088117Z","submitted_at":"2025-06-05T03:16:08Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification","version":1},"reference_index":53,"source":"arxiv_source","source_observed_at":"2026-08-07T10:44:59.468639Z"},"links":{"citing_paper":"/paper/2506.04592"},"observation_digest":"sha256:2d9eebfd91670ed5efe190aca0af2b30cebdc06ff96b0c7d1419da3ee7bfc9d5","observation_id":"3c5875da-165a-47b0-b879-fd5c1cd983ad","resolution":{"observed_at":"2026-08-07T10:44:59.468639Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.10528","last_updated":"2025-05-19T03:59:29Z","snapshot_observed_at":"2026-08-05T20:37:15.940003Z","submitted_at":"2024-02-16T09:29:50Z","title":"Can We Verify Step by Step for Incorrect Answer Detection?","version":4},"cited_work":{"arxiv_id":"2402.10528","doi":null,"metadata_source":"pith","pith_arxiv_id":"2402.10528","snapshot_observed_at":"2026-08-07T10:44:59.617022Z","title":"Can We Verify Step by Step for Incorrect Answer Detection?","venue":"cs.CL","work_id":"029c0037-3cbb-455c-a8cf-332f07c2aa42","year":2024},"citing_paper":{"arxiv_id":"2506.04592","last_updated":"2025-06-05T03:16:08Z","snapshot_observed_at":"2026-08-07T23:35:52.088117Z","submitted_at":"2025-06-05T03:16:08Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification","version":1},"reference_index":54,"source":"arxiv_source","source_observed_at":"2026-08-07T10:44:59.471059Z"},"links":{"cited_paper":"/paper/2402.10528","citing_paper":"/paper/2506.04592"},"observation_digest":"sha256:164d1a53bfdd448ce279ae362ca97f02b5e06483c5d21b8ff528350fcaba946d","observation_id":"71f9c518-aba2-4af0-9c69-70e41899d8f0","resolution":{"observed_at":"2026-08-07T10:44:59.620047Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2405.14804","last_updated":"2025-02-26T02:21:40Z","snapshot_observed_at":"2026-07-06T18:18:47.640933Z","submitted_at":"2024-05-23T17:13:50Z","title":"Can LLMs Solve longer Math Word Problems Better?","version":4},"cited_work":{"arxiv_id":"2405.14804","doi":null,"metadata_source":"pith","pith_arxiv_id":"2405.14804","snapshot_observed_at":"2026-08-07T10:44:59.607033Z","title":"Can LLMs Solve longer Math Word Problems Better?","venue":"cs.CL","work_id":"987f6cf0-dde3-4e57-9680-1631b0b347ff","year":2024},"citing_paper":{"arxiv_id":"2506.04592","last_updated":"2025-06-05T03:16:08Z","snapshot_observed_at":"2026-08-07T23:35:52.088117Z","submitted_at":"2025-06-05T03:16:08Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification","version":1},"reference_index":55,"source":"arxiv_source","source_observed_at":"2026-08-07T10:44:59.473658Z"},"links":{"cited_paper":"/paper/2405.14804","citing_paper":"/paper/2506.04592"},"observation_digest":"sha256:4fa9322b3ed848b0a8d73c2da6258cd7216d915eaff591d485e005b6371c8e17","observation_id":"1ee74c80-f7ec-4123-a69f-08d92a139756","resolution":{"observed_at":"2026-08-07T10:44:59.609963Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2501.13766","last_updated":"2025-02-25T08:15:43Z","snapshot_observed_at":"2026-07-06T20:25:03.950783Z","submitted_at":"2025-01-23T15:46:43Z","title":"UGMathBench: A Diverse and Dynamic Benchmark for Undergraduate-Level Mathematical Reasoning with Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.13766","snapshot_observed_at":"2026-08-07T10:44:59.476120Z","title":null,"venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.04592","last_updated":"2025-06-05T03:16:08Z","snapshot_observed_at":"2026-08-07T23:35:52.088117Z","submitted_at":"2025-06-05T03:16:08Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification","version":1},"reference_index":56,"source":"arxiv_source","source_observed_at":"2026-08-07T10:44:59.476120Z"},"links":{"cited_paper":"/paper/2501.13766","citing_paper":"/paper/2506.04592"},"observation_digest":"sha256:c49e3b0ac8edbf74dab2395c626b61fbc2a01dff30e098691c810fc572d90bfd","observation_id":"2d14513f-748d-4736-b5d6-42ef87186ab8","resolution":{"observed_at":"2026-08-07T10:44:59.476120Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.01524","last_updated":"2025-06-28T04:34:18Z","snapshot_observed_at":"2026-07-06T19:09:27.758765Z","submitted_at":"2024-09-03T01:40:21Z","title":"S^3cMath: Spontaneous Step-level Self-correction Makes Large Language Models Better Mathematical Reasoners","version":3},"cited_work":{"arxiv_id":"2409.01524","doi":null,"metadata_source":"pith","pith_arxiv_id":"2409.01524","snapshot_observed_at":"2026-08-07T10:44:59.589269Z","title":"S^3cMath: Spontaneous Step-level Self-correction Makes Large Language Models Better Mathematical Reasoners","venue":"cs.CL","work_id":"1a76dcd5-6670-44da-aced-f305b79a6507","year":2024},"citing_paper":{"arxiv_id":"2506.04592","last_updated":"2025-06-05T03:16:08Z","snapshot_observed_at":"2026-08-07T23:35:52.088117Z","submitted_at":"2025-06-05T03:16:08Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification","version":1},"reference_index":57,"source":"arxiv_source","source_observed_at":"2026-08-07T10:44:59.478740Z"},"links":{"cited_paper":"/paper/2409.01524","citing_paper":"/paper/2506.04592"},"observation_digest":"sha256:8c0c8e9af2d27db1b3320df6fb79d5efa10a4f327ea6f89f726955f6574ab2e3","observation_id":"a0f3c176-fe1f-4350-a9e6-7ca808dd995d","resolution":{"observed_at":"2026-08-07T10:44:59.593313Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:44:59.742895Z","title":null,"venue":null,"work_id":"7568cea3-7e87-4566-b3ad-c6a87b2119c8","year":2024},"citing_paper":{"arxiv_id":"2506.04592","last_updated":"2025-06-05T03:16:08Z","snapshot_observed_at":"2026-08-07T23:35:52.088117Z","submitted_at":"2025-06-05T03:16:08Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification","version":1},"reference_index":58,"source":"arxiv_source","source_observed_at":"2026-08-07T10:44:59.481389Z"},"links":{"citing_paper":"/paper/2506.04592"},"observation_digest":"sha256:3eb5ef011296678dd2589f5de3458e44e7a4828f5c9681374d1fe74a44d8f053","observation_id":"47bce112-3f26-45ee-9d5b-9e50b4c4754d","resolution":{"observed_at":"2026-08-07T10:44:59.745624Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2412.16075","last_updated":"2024-12-20T17:19:24Z","snapshot_observed_at":"2026-08-05T03:18:11.955456Z","submitted_at":"2024-12-20T17:19:24Z","title":"Formal Mathematical Reasoning: A New Frontier in AI","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2412.16075","snapshot_observed_at":"2026-08-07T10:44:59.483870Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.04592","last_updated":"2025-06-05T03:16:08Z","snapshot_observed_at":"2026-08-07T23:35:52.088117Z","submitted_at":"2025-06-05T03:16:08Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification","version":1},"reference_index":59,"source":"arxiv_source","source_observed_at":"2026-08-07T10:44:59.483870Z"},"links":{"cited_paper":"/paper/2412.16075","citing_paper":"/paper/2506.04592"},"observation_digest":"sha256:4c16cae735bcb61c845d0ba811811ceb621210f0828b31ad38c7aa2bc5a61348","observation_id":"868daa88-127e-4bf0-bd83-83800ec58daa","resolution":{"observed_at":"2026-08-07T10:44:59.483870Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:44:59.735974Z","title":null,"venue":null,"work_id":"0cd362ce-15d0-4fcc-8068-16e6f5ce1fd9","year":2023},"citing_paper":{"arxiv_id":"2506.04592","last_updated":"2025-06-05T03:16:08Z","snapshot_observed_at":"2026-08-07T23:35:52.088117Z","submitted_at":"2025-06-05T03:16:08Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification","version":1},"reference_index":60,"source":"arxiv_source","source_observed_at":"2026-08-07T10:44:59.486479Z"},"links":{"citing_paper":"/paper/2506.04592"},"observation_digest":"sha256:0f9f13ef300af3ebd1d96e62755e9e634e20a1dbd3a42cc26084f1f8ed72f400","observation_id":"0fcec193-ce17-428e-8d42-99ac63171109","resolution":{"observed_at":"2026-08-07T10:44:59.738544Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.03847","last_updated":"2025-06-18T16:07:47Z","snapshot_observed_at":"2026-08-07T23:35:20.471377Z","submitted_at":"2024-06-06T08:25:43Z","title":"Lean Workbook: A large-scale Lean problem set formalized from natural language math problems","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.03847","snapshot_observed_at":"2026-08-07T10:44:59.489003Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.04592","last_updated":"2025-06-05T03:16:08Z","snapshot_observed_at":"2026-08-07T23:35:52.088117Z","submitted_at":"2025-06-05T03:16:08Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification","version":1},"reference_index":61,"source":"arxiv_source","source_observed_at":"2026-08-07T10:44:59.489003Z"},"links":{"cited_paper":"/paper/2406.03847","citing_paper":"/paper/2506.04592"},"observation_digest":"sha256:1f3fabc3d0e3f9419d33b85741f763bc675f2c63b1f124695fc9342f970f6a37","observation_id":"b5b0d599-69e5-47f2-9495-7836b841ec13","resolution":{"observed_at":"2026-08-07T10:44:59.489003Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.06332","last_updated":"2024-05-24T07:09:21Z","snapshot_observed_at":"2026-08-06T07:58:05.717753Z","submitted_at":"2024-02-09T11:22:08Z","title":"InternLM-Math: Open Math Large Language Models Toward Verifiable Reasoning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.06332","snapshot_observed_at":"2026-08-07T10:44:59.491464Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.04592","last_updated":"2025-06-05T03:16:08Z","snapshot_observed_at":"2026-08-07T23:35:52.088117Z","submitted_at":"2025-06-05T03:16:08Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification","version":1},"reference_index":62,"source":"arxiv_source","source_observed_at":"2026-08-07T10:44:59.491464Z"},"links":{"cited_paper":"/paper/2402.06332","citing_paper":"/paper/2506.04592"},"observation_digest":"sha256:8d481b1ef18b44afa2f3eed6feb96b93d2008c3c5e2929ffc6984d361bf8bb88","observation_id":"7320371f-1ce9-4e52-bbd3-073523d26039","resolution":{"observed_at":"2026-08-07T10:44:59.491464Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2309.12284","last_updated":"2024-05-03T17:36:07Z","snapshot_observed_at":"2026-08-02T15:00:50.388422Z","submitted_at":"2023-09-21T17:45:42Z","title":"MetaMath: Bootstrap Your Own Mathematical Questions for Large Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2309.12284","snapshot_observed_at":"2026-08-07T10:44:59.494004Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.04592","last_updated":"2025-06-05T03:16:08Z","snapshot_observed_at":"2026-08-07T23:35:52.088117Z","submitted_at":"2025-06-05T03:16:08Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification","version":1},"reference_index":63,"source":"arxiv_source","source_observed_at":"2026-08-07T10:44:59.494004Z"},"links":{"cited_paper":"/paper/2309.12284","citing_paper":"/paper/2506.04592"},"observation_digest":"sha256:1571507c2305e8bd39e4cfbe6a9c7d89d242d1bd4c010534705fc86edc599016","observation_id":"a07f3824-ea3e-4f7b-abbc-eab215c1c427","resolution":{"observed_at":"2026-08-07T10:44:59.494004Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:44:59.496588Z","title":null,"venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.04592","last_updated":"2025-06-05T03:16:08Z","snapshot_observed_at":"2026-08-07T23:35:52.088117Z","submitted_at":"2025-06-05T03:16:08Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification","version":1},"reference_index":64,"source":"arxiv_source","source_observed_at":"2026-08-07T10:44:59.496588Z"},"links":{"citing_paper":"/paper/2506.04592"},"observation_digest":"sha256:75df65d2799df39bcab2c39d846c6e67a75cd2c3db85963544c92c7605917afb","observation_id":"2867b088-41da-4f0e-a4a4-4dc062b4c585","resolution":{"observed_at":"2026-08-07T10:44:59.496588Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:44:59.725351Z","title":null,"venue":null,"work_id":"c198bf61-a7d0-4d3f-ad0f-aa80c611f97b","year":2021},"citing_paper":{"arxiv_id":"2506.04592","last_updated":"2025-06-05T03:16:08Z","snapshot_observed_at":"2026-08-07T23:35:52.088117Z","submitted_at":"2025-06-05T03:16:08Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification","version":1},"reference_index":65,"source":"arxiv_source","source_observed_at":"2026-08-07T10:44:59.499292Z"},"links":{"citing_paper":"/paper/2506.04592"},"observation_digest":"sha256:02042fd35e6314d7b57d5f6d720fb4a377492355edd88f5c581e9c679945eef2","observation_id":"1fd2d262-7974-486b-9173-4e194f3bee24","resolution":{"observed_at":"2026-08-07T10:44:59.727780Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:44:59.718176Z","title":null,"venue":null,"work_id":"20da57e4-65ce-42bc-9fc6-367b6966abad","year":2024},"citing_paper":{"arxiv_id":"2506.04592","last_updated":"2025-06-05T03:16:08Z","snapshot_observed_at":"2026-08-07T23:35:52.088117Z","submitted_at":"2025-06-05T03:16:08Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification","version":1},"reference_index":66,"source":"arxiv_source","source_observed_at":"2026-08-07T10:44:59.501836Z"},"links":{"citing_paper":"/paper/2506.04592"},"observation_digest":"sha256:f928b59edc4c74c441167289f3e1d3dc79364709fe1a0411e51600dc46e3ae2f","observation_id":"ddfadd51-50ab-474c-8081-eb27fbe4a990","resolution":{"observed_at":"2026-08-07T10:44:59.720850Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2406.15877","last_updated":"2025-04-01T08:36:44Z","snapshot_observed_at":"2026-07-31T19:00:59.311189Z","submitted_at":"2024-06-22T15:52:04Z","title":"BigCodeBench: Benchmarking Code Generation with Diverse Function Calls and Complex Instructions","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.15877","snapshot_observed_at":"2026-08-07T10:44:59.504652Z","title":null,"venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.04592","last_updated":"2025-06-05T03:16:08Z","snapshot_observed_at":"2026-08-07T23:35:52.088117Z","submitted_at":"2025-06-05T03:16:08Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification","version":1},"reference_index":67,"source":"arxiv_source","source_observed_at":"2026-08-07T10:44:59.504652Z"},"links":{"cited_paper":"/paper/2406.15877","citing_paper":"/paper/2506.04592"},"observation_digest":"sha256:070b61cd597cd2a3a9374fa2c5292acc6f89df35b919f0e4e84e7c9848f73191","observation_id":"cafb388e-d148-480e-88de-e0e563730abd","resolution":{"observed_at":"2026-08-07T10:44:59.504652Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:44:59.507565Z","title":"online\" 'onlinestring :=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2506.04592","last_updated":"2025-06-05T03:16:08Z","snapshot_observed_at":"2026-08-07T23:35:52.088117Z","submitted_at":"2025-06-05T03:16:08Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification","version":1},"reference_index":68,"source":"arxiv_source","source_observed_at":"2026-08-07T10:44:59.507565Z"},"links":{"citing_paper":"/paper/2506.04592"},"observation_digest":"sha256:3be9ef9830df25514dfd5dd0258f2a4b56e5c9bfd08ce9230d86ae1673a65d58","observation_id":"337bf8a3-ef04-45c3-b269-90747aa3e08c","resolution":{"observed_at":"2026-08-07T10:44:59.507565Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-07T10:44:59.510685Z","title":"write newline","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2506.04592","last_updated":"2025-06-05T03:16:08Z","snapshot_observed_at":"2026-08-07T23:35:52.088117Z","submitted_at":"2025-06-05T03:16:08Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification","version":1},"reference_index":69,"source":"arxiv_source","source_observed_at":"2026-08-07T10:44:59.510685Z"},"links":{"citing_paper":"/paper/2506.04592"},"observation_digest":"sha256:f960d8b689d3e3e11a9ef9dc82dce386efb96470a185b65e9cde34e06e824c27","observation_id":"6279587b-880c-4ec6-aa59-5de5c9663cad","resolution":{"observed_at":"2026-08-07T10:44:59.510685Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2506.04592","last_updated":"2025-06-05T03:16:08Z","latest_version":1,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-07T23:35:52.088117Z","submitted_at":"2025-06-05T03:16:08Z","title":"Safe: Enhancing Mathematical Reasoning in Large Language Models via Retrospective Step-aware Formal Verification"},"reference_resolution":{"displayed":69,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":65,"verified_exact":4,"verified_fuzzy":0},"total_outbound_references":69},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"thesis":"As of 8 August 2026, this Paper Citation Record lists 69 of 69 outbound references and 2 inbound Pith citation observations for arXiv:2506.04592."}