{"as_of":"2026-08-08T20:21:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:cb5ccfe340f67b4e4c0a9f800cd9330d8612551cd5651501f716e679967281ca","coverage":[{"denominator":25,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":25,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-06T22:10:20.375361Z","state":"measured"},{"denominator":25,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":25,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-08T06:32:00.761636+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2506.22370/citation-record","integrity":"/paper/2506.22370/integrity","json":"/paper/2506.22370/citation-record.json","paper":"/paper/2506.22370"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T22:10:23.978395Z","title":null,"venue":null,"work_id":"1affe567-0f29-4dec-ac12-67f2ea6e4491","year":2024},"citing_paper":{"arxiv_id":"2506.22370","last_updated":"2025-09-07T11:26:53Z","snapshot_observed_at":"2026-08-06T22:03:10.107310Z","submitted_at":"2025-06-27T16:34:13Z","title":"Can Large Language Models Help Students Prove Software Correctness? An Experimental Study with Dafny","version":4},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-06T22:10:19.122363Z"},"links":{"citing_paper":"/paper/2506.22370"},"observation_digest":"sha256:5c025b05982f3a7816c989501d81922b3fba7b5a2be59e9fd2eab4831788b6eb","observation_id":"d4c4811e-c837-481f-9b17-92d39f489c0f","resolution":{"observed_at":"2026-08-06T22:10:24.100340Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.07789","last_updated":"2025-01-20T13:00:29Z","snapshot_observed_at":"2026-08-05T11:12:22.697766Z","submitted_at":"2025-01-20T13:00:29Z","title":"Do AI assistants help students write formal specifications? A study with ChatGPT and the B-Method","version":1},"cited_work":{"arxiv_id":"2502.07789","doi":null,"metadata_source":"pith","pith_arxiv_id":"2502.07789","snapshot_observed_at":"2026-08-06T22:10:20.458751Z","title":"Do AI assistants help students write formal specifications? A study with ChatGPT and the B-Method","venue":"cs.CY","work_id":"942eb4ca-59eb-4fa7-ae05-5885e07a631e","year":2025},"citing_paper":{"arxiv_id":"2506.22370","last_updated":"2025-09-07T11:26:53Z","snapshot_observed_at":"2026-08-06T22:03:10.107310Z","submitted_at":"2025-06-27T16:34:13Z","title":"Can Large Language Models Help Students Prove Software Correctness? An Experimental Study with Dafny","version":4},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-06T22:10:19.146590Z"},"links":{"cited_paper":"/paper/2502.07789","citing_paper":"/paper/2506.22370"},"observation_digest":"sha256:08f2cefd1118aa28d4fcaaa4ddf3a55096aec7418d62ab2cfcfe5ef581bf52fc","observation_id":"61c8ef18-3838-4605-b2eb-af93dbae9874","resolution":{"observed_at":"2026-08-06T22:10:20.557101Z","resolver_source":"local_arxiv","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T22:10:23.839263Z","title":"In: Proceedings of the 23rd ACM International Workshop on Formal Techniques for Java-like Programs","venue":null,"work_id":"7f9d4fb7-d890-4913-97a3-b073957b2a2f","year":2021},"citing_paper":{"arxiv_id":"2506.22370","last_updated":"2025-09-07T11:26:53Z","snapshot_observed_at":"2026-08-06T22:03:10.107310Z","submitted_at":"2025-06-27T16:34:13Z","title":"Can Large Language Models Help Students Prove Software Correctness? An Experimental Study with Dafny","version":4},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-06T22:10:19.210768Z"},"links":{"citing_paper":"/paper/2506.22370"},"observation_digest":"sha256:ed6659415e47146660a1353f138eb7c7bc4069d17dc980aab7fb4b9acafe80c8","observation_id":"2b587b74-2640-4695-97a5-38c078019567","resolution":{"observed_at":"2026-08-06T22:10:23.900773Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T22:10:23.678818Z","title":"Computers and Education: Artificial Intelligence7, 100290 (2024)","venue":null,"work_id":"e924f305-78f2-4423-9644-1c2078a28fa7","year":2024},"citing_paper":{"arxiv_id":"2506.22370","last_updated":"2025-09-07T11:26:53Z","snapshot_observed_at":"2026-08-06T22:03:10.107310Z","submitted_at":"2025-06-27T16:34:13Z","title":"Can Large Language Models Help Students Prove Software Correctness? An Experimental Study with Dafny","version":4},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-06T22:10:19.243840Z"},"links":{"citing_paper":"/paper/2506.22370"},"observation_digest":"sha256:9713430ca65200a4163d0853dd5e3607965b2628ff16370d00ef23a632b828ef","observation_id":"54a65751-e22c-4207-afe0-22fc08e0a5ed","resolution":{"observed_at":"2026-08-06T22:10:23.768504Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T22:10:23.506028Z","title":null,"venue":null,"work_id":"c6f5ae17-1b4f-4401-8d35-2c07d87003d6","year":2023},"citing_paper":{"arxiv_id":"2506.22370","last_updated":"2025-09-07T11:26:53Z","snapshot_observed_at":"2026-08-06T22:03:10.107310Z","submitted_at":"2025-06-27T16:34:13Z","title":"Can Large Language Models Help Students Prove Software Correctness? An Experimental Study with Dafny","version":4},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-06T22:10:19.290210Z"},"links":{"citing_paper":"/paper/2506.22370"},"observation_digest":"sha256:29b9a1aebf6a048ab3c68550e64fff473e85c291104ba6858637b8bcb8408320","observation_id":"b76a109a-0889-4aa4-9422-14de82073c15","resolution":{"observed_at":"2026-08-06T22:10:23.588382Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T22:10:23.301009Z","title":"Applied Sciences14(10), 4115 (2024)","venue":null,"work_id":"5c91896d-7377-416f-9906-532d1fca5212","year":2024},"citing_paper":{"arxiv_id":"2506.22370","last_updated":"2025-09-07T11:26:53Z","snapshot_observed_at":"2026-08-06T22:03:10.107310Z","submitted_at":"2025-06-27T16:34:13Z","title":"Can Large Language Models Help Students Prove Software Correctness? An Experimental Study with Dafny","version":4},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-06T22:10:19.378703Z"},"links":{"citing_paper":"/paper/2506.22370"},"observation_digest":"sha256:418299250ff50a3b34ac51e9721faa75b0067a2dceabe075dc5e3d05fc2939f1","observation_id":"4370b81c-38ba-4fd1-9bf5-a84598731bb6","resolution":{"observed_at":"2026-08-06T22:10:23.410110Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T22:10:19.451424Z","title":"In: International conference on logic for programming artificial intelligence and reasoning","venue":null,"work_id":null,"year":2010},"citing_paper":{"arxiv_id":"2506.22370","last_updated":"2025-09-07T11:26:53Z","snapshot_observed_at":"2026-08-06T22:03:10.107310Z","submitted_at":"2025-06-27T16:34:13Z","title":"Can Large Language Models Help Students Prove Software Correctness? An Experimental Study with Dafny","version":4},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-06T22:10:19.451424Z"},"links":{"citing_paper":"/paper/2506.22370"},"observation_digest":"sha256:6b8865a84f227ff8578f1ec0b167164fa78f048fedb403639f35f7a67c3c7ff9","observation_id":"c0889314-3feb-4250-8d20-554ca277e91f","resolution":{"observed_at":"2026-08-06T22:10:19.451424Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2406.08467","last_updated":"2024-06-12T17:53:31Z","snapshot_observed_at":"2026-07-06T18:29:49.284504Z","submitted_at":"2024-06-12T17:53:31Z","title":"DafnyBench: A Benchmark for Formal Software Verification","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2406.08467","snapshot_observed_at":"2026-08-06T22:10:19.488681Z","title":"arXiv preprint arXiv:2406.08467 (2024)","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.22370","last_updated":"2025-09-07T11:26:53Z","snapshot_observed_at":"2026-08-06T22:03:10.107310Z","submitted_at":"2025-06-27T16:34:13Z","title":"Can Large Language Models Help Students Prove Software Correctness? An Experimental Study with Dafny","version":4},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-06T22:10:19.488681Z"},"links":{"cited_paper":"/paper/2406.08467","citing_paper":"/paper/2506.22370"},"observation_digest":"sha256:62735b31bff0fec39ec7da5c507004c31dd43b7f670ab5bbe0ef41f8e7424c2a","observation_id":"56e87d08-1a2d-41db-ba05-793a0798aa78","resolution":{"observed_at":"2026-08-06T22:10:19.488681Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T22:10:23.071449Z","title":"In: Proceedings of the 56th ACM Technical Symposium on Computer Science Education V","venue":null,"work_id":"01b9b5a3-8d81-4bac-839a-7d4b146bacc2","year":2025},"citing_paper":{"arxiv_id":"2506.22370","last_updated":"2025-09-07T11:26:53Z","snapshot_observed_at":"2026-08-06T22:03:10.107310Z","submitted_at":"2025-06-27T16:34:13Z","title":"Can Large Language Models Help Students Prove Software Correctness? An Experimental Study with Dafny","version":4},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-06T22:10:19.493600Z"},"links":{"citing_paper":"/paper/2506.22370"},"observation_digest":"sha256:94af206c1708477dea23482fac8383e762aa60b4845e19c2284eb21439997650","observation_id":"75f1cf6d-cdbf-4b22-aade-a14bb249f638","resolution":{"observed_at":"2026-08-06T22:10:23.156920Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.18494","last_updated":"2024-11-05T03:38:39Z","snapshot_observed_at":"2026-07-06T19:38:52.603055Z","submitted_at":"2024-10-24T07:29:15Z","title":"Assured Automatic Programming via Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.18494","snapshot_observed_at":"2026-08-06T22:10:19.502790Z","title":"arXiv preprint arXiv:2410.18494 (2024)","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.22370","last_updated":"2025-09-07T11:26:53Z","snapshot_observed_at":"2026-08-06T22:03:10.107310Z","submitted_at":"2025-06-27T16:34:13Z","title":"Can Large Language Models Help Students Prove Software Correctness? An Experimental Study with Dafny","version":4},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-06T22:10:19.502790Z"},"links":{"cited_paper":"/paper/2410.18494","citing_paper":"/paper/2506.22370"},"observation_digest":"sha256:6c049ff76b9e88242b01203da07239d44f17583c7e9ee1e53a96156c95eb21da","observation_id":"6073b738-31f2-4d1c-9042-a9549d87fde2","resolution":{"observed_at":"2026-08-06T22:10:19.502790Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T22:10:22.905256Z","title":"Proceedings of the ACM on Software Engineering1(FSE), 812–835 (2024)","venue":null,"work_id":"120642af-deb1-4aaa-a851-17328244efba","year":2024},"citing_paper":{"arxiv_id":"2506.22370","last_updated":"2025-09-07T11:26:53Z","snapshot_observed_at":"2026-08-06T22:03:10.107310Z","submitted_at":"2025-06-27T16:34:13Z","title":"Can Large Language Models Help Students Prove Software Correctness? An Experimental Study with Dafny","version":4},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-06T22:10:19.547781Z"},"links":{"citing_paper":"/paper/2506.22370"},"observation_digest":"sha256:046daa719c80e901658f047f399cde3dc16921e190869d36474e9fff422bf759","observation_id":"b3f59be7-b167-40a1-a408-850c86a0c474","resolution":{"observed_at":"2026-08-06T22:10:22.970377Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T22:10:22.725214Z","title":"Proceedings of the ACM on Programming Languages9(OOPSLA1), 1519–1545 (2025)","venue":null,"work_id":"4eecf444-7e16-4e24-90c5-5498fffe8c9e","year":2025},"citing_paper":{"arxiv_id":"2506.22370","last_updated":"2025-09-07T11:26:53Z","snapshot_observed_at":"2026-08-06T22:03:10.107310Z","submitted_at":"2025-06-27T16:34:13Z","title":"Can Large Language Models Help Students Prove Software Correctness? An Experimental Study with Dafny","version":4},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-06T22:10:19.614436Z"},"links":{"citing_paper":"/paper/2506.22370"},"observation_digest":"sha256:e823dcab9394b28ebe881856d5f380141d9df788195470676128ab5d81c38ee2","observation_id":"46ed1b5e-4f0f-4645-aa9d-219e279763e4","resolution":{"observed_at":"2026-08-06T22:10:22.803786Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T22:10:22.551266Z","title":"In: NASA Formal Methods Symposium","venue":null,"work_id":"ebe52a74-3429-42cd-bbc9-dee35cf04fa2","year":2022},"citing_paper":{"arxiv_id":"2506.22370","last_updated":"2025-09-07T11:26:53Z","snapshot_observed_at":"2026-08-06T22:03:10.107310Z","submitted_at":"2025-06-27T16:34:13Z","title":"Can Large Language Models Help Students Prove Software Correctness? An Experimental Study with Dafny","version":4},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-06T22:10:19.712977Z"},"links":{"citing_paper":"/paper/2506.22370"},"observation_digest":"sha256:d11464c826316e7e00bef55175ee338be005c143ed226648304a2b9cbfe87c5b","observation_id":"97463543-45bf-4e0f-bb42-b083cbb0bec9","resolution":{"observed_at":"2026-08-06T22:10:22.635095Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T22:10:22.381774Z","title":"In: International Conference on Fundamentals of Software Engineering","venue":null,"work_id":"ca2382ef-2c66-4b74-b675-cc89d686f763","year":2025},"citing_paper":{"arxiv_id":"2506.22370","last_updated":"2025-09-07T11:26:53Z","snapshot_observed_at":"2026-08-06T22:03:10.107310Z","submitted_at":"2025-06-27T16:34:13Z","title":"Can Large Language Models Help Students Prove Software Correctness? An Experimental Study with Dafny","version":4},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-06T22:10:19.840113Z"},"links":{"citing_paper":"/paper/2506.22370"},"observation_digest":"sha256:8c3bffbcc7897bb8edc1755192e5cb5c343b888b1541c539f393b17f6e5f9543","observation_id":"68cd8d5a-b1e2-4995-9a80-49f1b69cee0c","resolution":{"observed_at":"2026-08-06T22:10:22.475039Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2411.15143","last_updated":"2024-11-05T19:27:56Z","snapshot_observed_at":"2026-08-07T20:32:30.401426Z","submitted_at":"2024-11-05T19:27:56Z","title":"dafny-annotator: AI-Assisted Verification of Dafny Programs","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2411.15143","snapshot_observed_at":"2026-08-06T22:10:19.945974Z","title":"arXiv preprint arXiv:2411.15143 (2024)","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.22370","last_updated":"2025-09-07T11:26:53Z","snapshot_observed_at":"2026-08-06T22:03:10.107310Z","submitted_at":"2025-06-27T16:34:13Z","title":"Can Large Language Models Help Students Prove Software Correctness? An Experimental Study with Dafny","version":4},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-06T22:10:19.945974Z"},"links":{"cited_paper":"/paper/2411.15143","citing_paper":"/paper/2506.22370"},"observation_digest":"sha256:09e8b9cc497173251e946bfec20878e80cfbefa1cfa6fa4c892ef447114a27f0","observation_id":"77098d97-8aa7-4a23-9456-0ae2eec6f4d3","resolution":{"observed_at":"2026-08-06T22:10:19.945974Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T22:10:22.228462Z","title":"In: Proceedings of the ACM Con- ference on Global Computing Education Vol 1","venue":null,"work_id":"c8e57a79-f93a-47db-bc63-677e6a946e66","year":2023},"citing_paper":{"arxiv_id":"2506.22370","last_updated":"2025-09-07T11:26:53Z","snapshot_observed_at":"2026-08-06T22:03:10.107310Z","submitted_at":"2025-06-27T16:34:13Z","title":"Can Large Language Models Help Students Prove Software Correctness? An Experimental Study with Dafny","version":4},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-06T22:10:20.024034Z"},"links":{"citing_paper":"/paper/2506.22370"},"observation_digest":"sha256:2124acd5f5e4c7fb422b426c890e3aacbc8ba65ff593f42c4e97c5ca5d9c79c2","observation_id":"6e4b1156-728a-455f-a3d4-d2cd697e98b6","resolution":{"observed_at":"2026-08-06T22:10:22.307291Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T22:10:22.092296Z","title":"It’s weird that it knows what I want","venue":null,"work_id":"662f1ad9-ec41-4dfc-bec9-2252bd026345","year":2023},"citing_paper":{"arxiv_id":"2506.22370","last_updated":"2025-09-07T11:26:53Z","snapshot_observed_at":"2026-08-06T22:03:10.107310Z","submitted_at":"2025-06-27T16:34:13Z","title":"Can Large Language Models Help Students Prove Software Correctness? An Experimental Study with Dafny","version":4},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-06T22:10:20.103786Z"},"links":{"citing_paper":"/paper/2506.22370"},"observation_digest":"sha256:a846e94164ebaefbee8c4be2ac7b03de9c904c4d65c24b836a5cda56d3331724","observation_id":"810191bd-857c-4856-9448-11c0964c1342","resolution":{"observed_at":"2026-08-06T22:10:22.151705Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T22:10:21.934307Z","title":"In: Proceedings of the 2024 ACM Conference on International Computing Education Research-Volume 1","venue":null,"work_id":"6eceed02-febf-4501-b315-76eaa6a59c78","year":2024},"citing_paper":{"arxiv_id":"2506.22370","last_updated":"2025-09-07T11:26:53Z","snapshot_observed_at":"2026-08-06T22:03:10.107310Z","submitted_at":"2025-06-27T16:34:13Z","title":"Can Large Language Models Help Students Prove Software Correctness? An Experimental Study with Dafny","version":4},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-06T22:10:20.159230Z"},"links":{"citing_paper":"/paper/2506.22370"},"observation_digest":"sha256:f44583319dccf01d1badf34fc4eae8461a3075be205010014924c8882d7bbb90","observation_id":"6ba07279-6064-41a6-9c1e-752edd9fbfa4","resolution":{"observed_at":"2026-08-06T22:10:22.026977Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2304.11214","last_updated":"2023-04-16T21:04:52Z","snapshot_observed_at":"2026-08-05T18:23:16.314109Z","submitted_at":"2023-04-16T21:04:52Z","title":"Exploring the Use of ChatGPT as a Tool for Learning and Assessment in Undergraduate Computer Science Curriculum: Opportunities and Challenges","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2304.11214","snapshot_observed_at":"2026-08-06T22:10:20.208512Z","title":"arXiv preprint arXiv:2304.11214 (2023)","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.22370","last_updated":"2025-09-07T11:26:53Z","snapshot_observed_at":"2026-08-06T22:03:10.107310Z","submitted_at":"2025-06-27T16:34:13Z","title":"Can Large Language Models Help Students Prove Software Correctness? An Experimental Study with Dafny","version":4},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-06T22:10:20.208512Z"},"links":{"cited_paper":"/paper/2304.11214","citing_paper":"/paper/2506.22370"},"observation_digest":"sha256:3adc24753496c281e91f0a671e2198c3744928ea99624c61142a2d5d3d121deb","observation_id":"e952621b-b581-44f9-b11e-b93d7decd9b3","resolution":{"observed_at":"2026-08-06T22:10:20.208512Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T22:10:21.716908Z","title":null,"venue":null,"work_id":"8bd260eb-dc52-42bb-96d8-6839772c2ddf","year":2022},"citing_paper":{"arxiv_id":"2506.22370","last_updated":"2025-09-07T11:26:53Z","snapshot_observed_at":"2026-08-06T22:03:10.107310Z","submitted_at":"2025-06-27T16:34:13Z","title":"Can Large Language Models Help Students Prove Software Correctness? An Experimental Study with Dafny","version":4},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-06T22:10:20.225656Z"},"links":{"citing_paper":"/paper/2506.22370"},"observation_digest":"sha256:9411db1f9334d9bc152e65df80cd555cb0a7818bf5dd819b0bc555cfde8bfbea","observation_id":"1c3fc999-e5e6-4e00-81fe-42c42ecf9885","resolution":{"observed_at":"2026-08-06T22:10:21.812585Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T22:10:21.550543Z","title":"In: Proceedings of the 2024 IEEE/ACM 12th International Conference on Formal Methods in Software Engineering (FormaliSE)","venue":null,"work_id":"79e5b6f9-b87b-40e7-9d3a-9e142b7a0c67","year":2024},"citing_paper":{"arxiv_id":"2506.22370","last_updated":"2025-09-07T11:26:53Z","snapshot_observed_at":"2026-08-06T22:03:10.107310Z","submitted_at":"2025-06-27T16:34:13Z","title":"Can Large Language Models Help Students Prove Software Correctness? An Experimental Study with Dafny","version":4},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-06T22:10:20.255691Z"},"links":{"citing_paper":"/paper/2506.22370"},"observation_digest":"sha256:8776a471244ecbed25deff41587ca63bb8afa022296394afd109661ea96cfffe","observation_id":"ec95da06-b675-4a8f-a4f9-072de1460dcd","resolution":{"observed_at":"2026-08-06T22:10:21.643639Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T22:10:21.305357Z","title":"In: International Symposium on AI Verification","venue":null,"work_id":"e5f4fab7-101c-4014-badc-24b8b8c1d0c0","year":2024},"citing_paper":{"arxiv_id":"2506.22370","last_updated":"2025-09-07T11:26:53Z","snapshot_observed_at":"2026-08-06T22:03:10.107310Z","submitted_at":"2025-06-27T16:34:13Z","title":"Can Large Language Models Help Students Prove Software Correctness? An Experimental Study with Dafny","version":4},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-06T22:10:20.273290Z"},"links":{"citing_paper":"/paper/2506.22370"},"observation_digest":"sha256:3326132fc083fc751d50e63ce14179bac2b3c0b975c34f2d5b8389c38e82ddd2","observation_id":"bf8e2275-0295-4037-b569-d1d6f8c44caa","resolution":{"observed_at":"2026-08-06T22:10:21.420615Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T22:10:21.090457Z","title":"International Journal of Educational Technology in Higher Education21(1), 14 (2024)","venue":null,"work_id":"0973dac4-91e1-4f6a-a239-db7c289eaa84","year":2024},"citing_paper":{"arxiv_id":"2506.22370","last_updated":"2025-09-07T11:26:53Z","snapshot_observed_at":"2026-08-06T22:03:10.107310Z","submitted_at":"2025-06-27T16:34:13Z","title":"Can Large Language Models Help Students Prove Software Correctness? An Experimental Study with Dafny","version":4},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-06T22:10:20.292400Z"},"links":{"citing_paper":"/paper/2506.22370"},"observation_digest":"sha256:a548bce0ce002dca7c52ea24b2137c77b4a607ea0ed4678b2a359328bdf06624","observation_id":"0b6970fd-49df-42d5-98c6-229f88c64605","resolution":{"observed_at":"2026-08-06T22:10:21.182525Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T22:10:20.935392Z","title":"In: 23rd International Conference on Software Engineering and Formal Methods (SEFM) (2025)","venue":null,"work_id":"c13b4349-0e2d-40a3-96d6-179d2e86bda0","year":2025},"citing_paper":{"arxiv_id":"2506.22370","last_updated":"2025-09-07T11:26:53Z","snapshot_observed_at":"2026-08-06T22:03:10.107310Z","submitted_at":"2025-06-27T16:34:13Z","title":"Can Large Language Models Help Students Prove Software Correctness? An Experimental Study with Dafny","version":4},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-06T22:10:20.317739Z"},"links":{"citing_paper":"/paper/2506.22370"},"observation_digest":"sha256:42b3aab83e5b68e4da6d7ac5143093963772238e55fbb18ab098ff7cbc0e50af","observation_id":"e6684581-ff6d-4104-8e4d-7a10a94d377d","resolution":{"observed_at":"2026-08-06T22:10:20.989008Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T22:10:20.698689Z","title":"In: Proceedings of the 46th International Conference on Software Engineering: Software Engineering Education and Training","venue":null,"work_id":"55b99445-1c48-4ae8-bc18-e050701590b0","year":2024},"citing_paper":{"arxiv_id":"2506.22370","last_updated":"2025-09-07T11:26:53Z","snapshot_observed_at":"2026-08-06T22:03:10.107310Z","submitted_at":"2025-06-27T16:34:13Z","title":"Can Large Language Models Help Students Prove Software Correctness? An Experimental Study with Dafny","version":4},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-06T22:10:20.375361Z"},"links":{"citing_paper":"/paper/2506.22370"},"observation_digest":"sha256:5a981699a975f6a237a754a65a49b9251837e07452589349918deeda90606751","observation_id":"f0e85018-93e9-4385-a92c-0cef8b8d3ce6","resolution":{"observed_at":"2026-08-06T22:10:20.839689Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-08T06:32:00.761636+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2506.22370","last_updated":"2025-09-07T11:26:53Z","latest_version":4,"primary_category":"cs.SE","snapshot_observed_at":"2026-08-06T22:03:10.107310Z","submitted_at":"2025-06-27T16:34:13Z","title":"Can Large Language Models Help Students Prove Software Correctness? An Experimental Study with Dafny"},"reference_resolution":{"displayed":25,"state_counts":{"malformed_identifier":0,"metadata_mismatch":1,"parse_uncertain":0,"unresolved":8,"verified_exact":0,"verified_fuzzy":16},"total_outbound_references":25},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-08T06:32:00.761636+00:00","source":"crossref"},{"observed_at":"2026-08-08T06:31:55.24221+00:00","source":"retraction_watch"}],"thesis":"As of 8 August 2026, this Paper Citation Record lists 25 of 25 outbound references and 0 inbound Pith citation observations for arXiv:2506.22370."}