{"as_of":"2026-08-21T06:53:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:a38f4a01361a5db5c350546a19e6836da8b6f097ed6b8688ef5eabff25398feb","coverage":[{"denominator":31,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":31,"source":"paper_references, paper_reference_links","source_observed_at":"2026-07-03T19:23:50.577675Z","state":"measured"},{"denominator":32,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":32,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-21T06:32:19.484+00:00","state":"measured"},{"denominator":1,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":1,"source":"paper_references, paper_reference_links","source_observed_at":"2026-07-31T10:52:17.791456Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"pith","source_observed_at":"2026-07-31T12:16:09.833992Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2607.01360","last_updated":"2026-07-01T18:20:27Z","snapshot_observed_at":"2026-08-15T16:04:59.762247Z","submitted_at":"2026-07-01T18:20:27Z","title":"Benchmarking Code Improvement with Progressive, Adaptive, and Interactive Feedback","version":1},"cited_work":{"arxiv_id":"2607.01360","doi":"10.48550/arxiv.2607.01360","metadata_source":"pith","pith_arxiv_id":"2607.01360","snapshot_observed_at":"2026-07-31T12:16:09.833992Z","title":"Benchmarking Code Improvement with Progressive, Adaptive, and Interactive Feedback","venue":"cs.SE","work_id":"5e37e315-5d56-4a6e-b955-13b9c2541121","year":2026},"citing_paper":{"arxiv_id":"2607.24604","last_updated":"2026-07-27T16:05:23Z","snapshot_observed_at":"2026-08-13T15:37:49.163509Z","submitted_at":"2026-07-27T16:05:23Z","title":"Looping Is Not Reliability: State-Bound Evidence and Typed Revision Contracts for Agentic Code Repair","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-07-31T10:52:17.791456Z"},"links":{"cited_paper":"/paper/2607.01360","citing_paper":"/paper/2607.24604"},"observation_digest":"sha256:49b3f2ae437ebe9295f0c22dd714ef9ee38a84abdc47d61b5e10742e51846699","observation_id":"eac9c070-c0f6-4b28-a262-7a09080005d4","resolution":{"observed_at":"2026-07-31T10:56:25.047648Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2607.01360/citation-record","integrity":"/paper/2607.01360/integrity","json":"/paper/2607.01360/citation-record.json","paper":"/paper/2607.01360"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2107.03374","last_updated":"2021-07-14T17:16:02Z","snapshot_observed_at":"2026-08-20T20:50:23.483838Z","submitted_at":"2021-07-07T17:41:24Z","title":"Evaluating Large Language Models Trained on Code","version":2},"cited_work":{"arxiv_id":"2107.03374","doi":"10.48550/arxiv.2107.03374","metadata_source":"pith","pith_arxiv_id":"2107.03374","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Evaluating Large Language Models Trained on Code","venue":"cs.LG","work_id":"042493e9-b26f-4b4e-bbde-382072ca9b08","year":2021},"citing_paper":{"arxiv_id":"2607.01360","last_updated":"2026-07-01T18:20:27Z","snapshot_observed_at":"2026-08-15T16:04:59.762247Z","submitted_at":"2026-07-01T18:20:27Z","title":"Benchmarking Code Improvement with Progressive, Adaptive, and Interactive Feedback","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-07-03T19:23:50.577675Z"},"links":{"cited_paper":"/paper/2107.03374","citing_paper":"/paper/2607.01360"},"observation_digest":"sha256:4e495abe08bba2c98ce01cc49ea8d7048aacb9eb03249ea80b474255e63cc5fd","observation_id":"0b104bec-47e9-47f2-b16c-fad90c66df98","resolution":{"observed_at":"2026-07-03T19:28:51.853327Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-01T08:08:23.404839+00:00","source":"crossref_status_cache"},{"observed_at":"2026-08-01T08:08:23.404839+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-05T07:50:48.541419Z","title":"Livecodebench: Holistic and contamination free evaluation of large language models for code,","venue":null,"work_id":"22be2a97-8b3d-4df0-b00e-fb7dfedec03a","year":null},"citing_paper":{"arxiv_id":"2607.01360","last_updated":"2026-07-01T18:20:27Z","snapshot_observed_at":"2026-08-15T16:04:59.762247Z","submitted_at":"2026-07-01T18:20:27Z","title":"Benchmarking Code Improvement with Progressive, Adaptive, and Interactive Feedback","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-07-03T19:23:50.577675Z"},"links":{"citing_paper":"/paper/2607.01360"},"observation_digest":"sha256:6b2f752d6ed74cfb5301180ad50a1d997b07b5b92898686cfca459946515ad80","observation_id":"238fd308-0870-4db3-a60e-caeee6dcbfe5","resolution":{"observed_at":"2026-07-05T03:00:39.474479Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-05T07:50:48.539523Z","title":"Available: https://openreview.net/forum?id=chfJJYC3iL","venue":null,"work_id":"d53eec81-a245-41f1-986b-6cefe8baf554","year":null},"citing_paper":{"arxiv_id":"2607.01360","last_updated":"2026-07-01T18:20:27Z","snapshot_observed_at":"2026-08-15T16:04:59.762247Z","submitted_at":"2026-07-01T18:20:27Z","title":"Benchmarking Code Improvement with Progressive, Adaptive, and Interactive Feedback","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-07-03T19:23:50.577675Z"},"links":{"citing_paper":"/paper/2607.01360"},"observation_digest":"sha256:f3059873944e83ed5275ef4d2c366613af9bcabf6b1597adc6b8935e9807af59","observation_id":"7181e20a-21d8-4fbc-bf2e-9af7d2a2b46a","resolution":{"observed_at":"2026-07-05T03:00:39.472657Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-05T03:00:39.475178Z","title":"SWE-bench: Can language models resolve real-world github issues?","venue":null,"work_id":"79d48cc5-2537-4200-80fd-8d70ed3973fb","year":2024},"citing_paper":{"arxiv_id":"2607.01360","last_updated":"2026-07-01T18:20:27Z","snapshot_observed_at":"2026-08-15T16:04:59.762247Z","submitted_at":"2026-07-01T18:20:27Z","title":"Benchmarking Code Improvement with Progressive, Adaptive, and Interactive Feedback","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-07-03T19:23:50.577675Z"},"links":{"citing_paper":"/paper/2607.01360"},"observation_digest":"sha256:c8420ec49a0f0233f21962b58983352ef57056fa1c514cf2f00dfc2df40d3f3f","observation_id":"9a61d9d4-3092-4501-97f7-6005e33af4d2","resolution":{"observed_at":"2026-07-05T03:00:39.476465Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"0212.368038","doi":"10.1145/3650212.3680381","metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"In Proceedings of the 33rd ACM SIGSOFT International Symposium on Software Testing and Analysis (Vienna, Austria) (ISSTA 2024)","venue":null,"work_id":"89646ac3-6811-4a11-9005-7942fb6e88d7","year":2024},"citing_paper":{"arxiv_id":"2607.01360","last_updated":"2026-07-01T18:20:27Z","snapshot_observed_at":"2026-08-15T16:04:59.762247Z","submitted_at":"2026-07-01T18:20:27Z","title":"Benchmarking Code Improvement with Progressive, Adaptive, and Interactive Feedback","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-07-03T19:23:50.577675Z"},"links":{"citing_paper":"/paper/2607.01360"},"observation_digest":"sha256:1710a809e8d431c84ee1d8aeab995cf7583c0532a8c21004c9c2cf718051d4b1","observation_id":"211008f7-c255-4e3f-938b-edde3869310e","resolution":{"observed_at":"2026-07-03T19:28:51.391555Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-05T03:00:39.484176Z","title":"Agentless: Demystifying llm-based software engineering agents,","venue":null,"work_id":"1b57dfc3-5888-4f50-8afa-c90cfaf88722","year":2024},"citing_paper":{"arxiv_id":"2607.01360","last_updated":"2026-07-01T18:20:27Z","snapshot_observed_at":"2026-08-15T16:04:59.762247Z","submitted_at":"2026-07-01T18:20:27Z","title":"Benchmarking Code Improvement with Progressive, Adaptive, and Interactive Feedback","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-07-03T19:23:50.577675Z"},"links":{"citing_paper":"/paper/2607.01360"},"observation_digest":"sha256:5d41ff35943e88901ed59cef77fb956b9a9a1890cd4282b5a8ccb0ca88fe8d5e","observation_id":"63672f72-25d8-4cc5-a4fb-1b04d029f994","resolution":{"observed_at":"2026-07-05T03:00:39.485329Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"6805.278682","doi":"10.1145/2786805.2786825","metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"K., Barr, E","venue":null,"work_id":"a8aa7336-a39d-4d7b-83c3-73d6730eb486","year":2015},"citing_paper":{"arxiv_id":"2607.01360","last_updated":"2026-07-01T18:20:27Z","snapshot_observed_at":"2026-08-15T16:04:59.762247Z","submitted_at":"2026-07-01T18:20:27Z","title":"Benchmarking Code Improvement with Progressive, Adaptive, and Interactive Feedback","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-07-03T19:23:50.577675Z"},"links":{"citing_paper":"/paper/2607.01360"},"observation_digest":"sha256:85dc8d78ea384cc0ca37cc4af0aa7b59ad79c9b122974763d032f46c092a3453","observation_id":"944e2715-b456-45b5-90ab-049c8b511a2c","resolution":{"observed_at":"2026-07-03T19:28:51.396131Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-05-22T14:53:02.966355+00:00","source":"crossref_status_cache"},{"observed_at":"2026-05-22T14:53:02.966355+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2503.15223","last_updated":"2025-09-09T02:59:54Z","snapshot_observed_at":"2026-08-20T23:36:56.136190Z","submitted_at":"2025-03-19T14:02:21Z","title":"Are \"Solved Issues\" in SWE-bench Really Solved Correctly? An Empirical Study","version":2},"cited_work":{"arxiv_id":"2503.15223","doi":"10.48550/arxiv.2503.15223","metadata_source":"pith","pith_arxiv_id":"2503.15223","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Solved Issues","venue":"cs.SE","work_id":"c6c784d6-07f6-492d-9ed7-f5565484586a","year":2025},"citing_paper":{"arxiv_id":"2607.01360","last_updated":"2026-07-01T18:20:27Z","snapshot_observed_at":"2026-08-15T16:04:59.762247Z","submitted_at":"2026-07-01T18:20:27Z","title":"Benchmarking Code Improvement with Progressive, Adaptive, and Interactive Feedback","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-07-03T19:23:50.577675Z"},"links":{"cited_paper":"/paper/2503.15223","citing_paper":"/paper/2607.01360"},"observation_digest":"sha256:b11596c5fef75e8773b4305cf35b988e7fb07e8dc5fd4cac95e4279983ef3646","observation_id":"0525c222-92df-4326-a555-5af6b5b84fee","resolution":{"observed_at":"2026-07-03T19:28:51.859271Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-07-11T01:50:47.219299+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T01:50:47.219299+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2503.15223","last_updated":"2025-09-09T02:59:54Z","snapshot_observed_at":"2026-08-20T23:36:56.136190Z","submitted_at":"2025-03-19T14:02:21Z","title":"Are \"Solved Issues\" in SWE-bench Really Solved Correctly? An Empirical Study","version":2},"cited_work":{"arxiv_id":"2503.15223","doi":"10.48550/arxiv.2503.15223","metadata_source":"pith","pith_arxiv_id":"2503.15223","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Solved Issues","venue":"cs.SE","work_id":"c6c784d6-07f6-492d-9ed7-f5565484586a","year":2025},"citing_paper":{"arxiv_id":"2607.01360","last_updated":"2026-07-01T18:20:27Z","snapshot_observed_at":"2026-08-15T16:04:59.762247Z","submitted_at":"2026-07-01T18:20:27Z","title":"Benchmarking Code Improvement with Progressive, Adaptive, and Interactive Feedback","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-07-03T19:23:50.577675Z"},"links":{"cited_paper":"/paper/2503.15223","citing_paper":"/paper/2607.01360"},"observation_digest":"sha256:2c1d7ba2f4c6ca2bd180147e852a3cb6d09409deabad5ce3b6f7d0af5c2d6de4","observation_id":"f9d4c3f7-8477-4a86-ab49-eb5f2c3cf23c","resolution":{"observed_at":"2026-07-03T19:28:51.387866Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-07-11T01:50:47.219299+00:00","source":"crossref_status_cache"},{"observed_at":"2026-07-11T01:50:47.219299+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-05T03:00:39.459257Z","title":"Introducing SWE-bench verified,","venue":null,"work_id":"6e630aa1-f87c-4a2e-ab6e-a661a0cd41aa","year":2024},"citing_paper":{"arxiv_id":"2607.01360","last_updated":"2026-07-01T18:20:27Z","snapshot_observed_at":"2026-08-15T16:04:59.762247Z","submitted_at":"2026-07-01T18:20:27Z","title":"Benchmarking Code Improvement with Progressive, Adaptive, and Interactive Feedback","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-07-03T19:23:50.577675Z"},"links":{"citing_paper":"/paper/2607.01360"},"observation_digest":"sha256:37557669daea822adf15f2f5c16c73eb29e1fa16691feb5576dbd2ce64960800","observation_id":"91093192-025f-4f59-84e6-92ef1f35b9b0","resolution":{"observed_at":"2026-07-05T03:00:39.460584Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-05T03:00:39.466166Z","title":"Graph-based, self-supervised program repair from diagnostic feedback,","venue":null,"work_id":"c8de91e6-1ec0-4b9e-b7b5-5590afe480f1","year":2020},"citing_paper":{"arxiv_id":"2607.01360","last_updated":"2026-07-01T18:20:27Z","snapshot_observed_at":"2026-08-15T16:04:59.762247Z","submitted_at":"2026-07-01T18:20:27Z","title":"Benchmarking Code Improvement with Progressive, Adaptive, and Interactive Feedback","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-07-03T19:23:50.577675Z"},"links":{"citing_paper":"/paper/2607.01360"},"observation_digest":"sha256:f34bd098af545c1e14e8ae24d5307510a7329366e1f864fa9ecdc71762060a07","observation_id":"673c6e0f-2a8e-4be6-b917-52fb9ed00a85","resolution":{"observed_at":"2026-07-05T03:00:39.467339Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2504.06939","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-03T19:28:51.854444Z","title":"Feedbackeval: A benchmark for evaluating large language models in feedback-driven code repair tasks","venue":null,"work_id":"b2a2c35c-92b4-499c-a008-e65890c2eaee","year":2025},"citing_paper":{"arxiv_id":"2607.01360","last_updated":"2026-07-01T18:20:27Z","snapshot_observed_at":"2026-08-15T16:04:59.762247Z","submitted_at":"2026-07-01T18:20:27Z","title":"Benchmarking Code Improvement with Progressive, Adaptive, and Interactive Feedback","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-07-03T19:23:50.577675Z"},"links":{"citing_paper":"/paper/2607.01360"},"observation_digest":"sha256:c68d9c8ccdfafdf2109118fc9a31cf5732b90b6c10889a56b60fca207627f5f5","observation_id":"5f356c19-650a-4d27-a5bf-362d4f1fb968","resolution":{"observed_at":"2026-07-03T19:28:51.856434Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":"10.18653/v1/2023.acl-long.45","metadata_source":"openalex","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Self-edit: Fault-aware code editor for code generation,","venue":null,"work_id":"47385af0-29eb-42f2-9ef9-72fa62c54ddb","year":2023},"citing_paper":{"arxiv_id":"2607.01360","last_updated":"2026-07-01T18:20:27Z","snapshot_observed_at":"2026-08-15T16:04:59.762247Z","submitted_at":"2026-07-01T18:20:27Z","title":"Benchmarking Code Improvement with Progressive, Adaptive, and Interactive Feedback","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-07-03T19:23:50.577675Z"},"links":{"citing_paper":"/paper/2607.01360"},"observation_digest":"sha256:f1395c9a09800dba8e4d717e0823e841c84fb5192ddff884742c58a89024e459","observation_id":"433e75ac-23f1-4466-bd7c-27d7b5882ffd","resolution":{"observed_at":"2026-07-03T19:28:51.389944Z","resolver_source":"doi","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-17T19:08:30.757148+00:00","source":"crossref_status_cache"},{"observed_at":"2026-08-17T19:08:30.757148+00:00","source":"openalex_status_cache"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-05T03:00:39.467851Z","title":"Teaching large language models to self-debug,","venue":null,"work_id":"3358491a-cbc8-4788-885e-40af075821eb","year":2024},"citing_paper":{"arxiv_id":"2607.01360","last_updated":"2026-07-01T18:20:27Z","snapshot_observed_at":"2026-08-15T16:04:59.762247Z","submitted_at":"2026-07-01T18:20:27Z","title":"Benchmarking Code Improvement with Progressive, Adaptive, and Interactive Feedback","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-07-03T19:23:50.577675Z"},"links":{"citing_paper":"/paper/2607.01360"},"observation_digest":"sha256:8e85e9f25b5bbb14f6ae715f555b335cc7487f47b74830eb05402bfbdbceb164","observation_id":"7f9ce8dc-90c4-49bd-b33e-477909ad5e02","resolution":{"observed_at":"2026-07-05T03:00:39.469176Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2502.02928","last_updated":"2025-06-21T05:27:49Z","snapshot_observed_at":"2026-08-19T16:57:09.258890Z","submitted_at":"2025-02-05T06:43:40Z","title":"Large Language Model Guided Self-Debugging Code Generation","version":2},"cited_work":{"arxiv_id":"2502.02928","doi":"10.48550/arxiv.2502.02928","metadata_source":"arxiv_reference","pith_arxiv_id":"2502.02928","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Large language model guided self-debugging code generation.arXiv preprint arXiv:2502.02928","venue":"ArXiv.org","work_id":"418dca89-6599-471b-85c5-d854ea2413d2","year":2025},"citing_paper":{"arxiv_id":"2607.01360","last_updated":"2026-07-01T18:20:27Z","snapshot_observed_at":"2026-08-15T16:04:59.762247Z","submitted_at":"2026-07-01T18:20:27Z","title":"Benchmarking Code Improvement with Progressive, Adaptive, and Interactive Feedback","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-07-03T19:23:50.577675Z"},"links":{"cited_paper":"/paper/2502.02928","citing_paper":"/paper/2607.01360"},"observation_digest":"sha256:d9bbb573bb25b9a6bd0eff6074a0f5e9c88cc1f19d046e67e492ea392942620c","observation_id":"f80555c1-b40b-4d1a-aad1-05dade6dd969","resolution":{"observed_at":"2026-07-03T19:28:51.382722Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-05T03:00:39.464461Z","title":"ConvCodeWorld: Benchmarking conversational code generation in reproducible feedback environments,","venue":null,"work_id":"e5ee7c1a-1900-4e3b-85ae-e9fd3a9edbba","year":2025},"citing_paper":{"arxiv_id":"2607.01360","last_updated":"2026-07-01T18:20:27Z","snapshot_observed_at":"2026-08-15T16:04:59.762247Z","submitted_at":"2026-07-01T18:20:27Z","title":"Benchmarking Code Improvement with Progressive, Adaptive, and Interactive Feedback","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-07-03T19:23:50.577675Z"},"links":{"citing_paper":"/paper/2607.01360"},"observation_digest":"sha256:42c305f436f7902eb4d3423b6142c0704aff3cd842989af478ddea03c4c7d598","observation_id":"eb871329-fe48-4fe4-b4e5-425382ddab16","resolution":{"observed_at":"2026-07-05T03:00:39.465607Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-05T03:00:39.489555Z","title":"When benchmarks talk: Re-evaluating code LLMs with interactive feedback,","venue":null,"work_id":"a0c1f4b3-d9d2-461c-93ab-adbffda9e63c","year":2025},"citing_paper":{"arxiv_id":"2607.01360","last_updated":"2026-07-01T18:20:27Z","snapshot_observed_at":"2026-08-15T16:04:59.762247Z","submitted_at":"2026-07-01T18:20:27Z","title":"Benchmarking Code Improvement with Progressive, Adaptive, and Interactive Feedback","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-07-03T19:23:50.577675Z"},"links":{"citing_paper":"/paper/2607.01360"},"observation_digest":"sha256:02f67bdfa9b133bfcb8477b20a5dfabb2736c3ecdb455a0fea9f23ef53417192","observation_id":"86e9ecba-ede8-4cbf-b4f4-7a8224d1ddb4","resolution":{"observed_at":"2026-07-05T03:00:39.490698Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-05T03:00:39.477076Z","title":"Available: https://pairbench.site","venue":null,"work_id":"4aa77e50-6350-456d-92b4-994dfbb1d23b","year":null},"citing_paper":{"arxiv_id":"2607.01360","last_updated":"2026-07-01T18:20:27Z","snapshot_observed_at":"2026-08-15T16:04:59.762247Z","submitted_at":"2026-07-01T18:20:27Z","title":"Benchmarking Code Improvement with Progressive, Adaptive, and Interactive Feedback","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-07-03T19:23:50.577675Z"},"links":{"citing_paper":"/paper/2607.01360"},"observation_digest":"sha256:256453f1c870c7ff8edab3cdb8db7b3c15dc76bf99009921672b52927fc1345a","observation_id":"20ec7849-b1de-4a85-98c2-8431755ea605","resolution":{"observed_at":"2026-07-05T03:00:39.478192Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2108.07732","last_updated":"2021-08-16T03:57:30Z","snapshot_observed_at":"2026-08-15T17:40:38.050939Z","submitted_at":"2021-08-16T03:57:30Z","title":"Program Synthesis with Large Language Models","version":1},"cited_work":{"arxiv_id":"2108.07732","doi":"10.1007/s11390-025-5518-5","metadata_source":"pith","pith_arxiv_id":"2108.07732","snapshot_observed_at":"2026-07-11T11:50:26.030339Z","title":"Program Synthesis with Large Language Models","venue":"cs.PL","work_id":"fd241a05-03b9-4de2-9588-9d77ce176125","year":2021},"citing_paper":{"arxiv_id":"2607.01360","last_updated":"2026-07-01T18:20:27Z","snapshot_observed_at":"2026-08-15T16:04:59.762247Z","submitted_at":"2026-07-01T18:20:27Z","title":"Benchmarking Code Improvement with Progressive, Adaptive, and Interactive Feedback","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-07-03T19:23:50.577675Z"},"links":{"cited_paper":"/paper/2108.07732","citing_paper":"/paper/2607.01360"},"observation_digest":"sha256:66251a23f9d77b4fad1f789e6aae59405a1c1e15956b00a4601e9ea897beca0c","observation_id":"77bb7840-4d74-452a-a413-580400cd71c8","resolution":{"observed_at":"2026-07-03T19:28:51.849550Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-05T03:00:39.478776Z","title":"InterCode: Stan- dardizing and benchmarking interactive coding with execution feed- back,","venue":null,"work_id":"b2bed944-9a5e-4b0e-a4ef-66e1e35c8f3d","year":2023},"citing_paper":{"arxiv_id":"2607.01360","last_updated":"2026-07-01T18:20:27Z","snapshot_observed_at":"2026-08-15T16:04:59.762247Z","submitted_at":"2026-07-01T18:20:27Z","title":"Benchmarking Code Improvement with Progressive, Adaptive, and Interactive Feedback","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-07-03T19:23:50.577675Z"},"links":{"citing_paper":"/paper/2607.01360"},"observation_digest":"sha256:8a58773548af88e602e5f3e4be3d04e5965f9394500aa1555bd3d8c735a63bcf","observation_id":"d6f5e396-2f71-41c9-ae06-8e4e6cb56def","resolution":{"observed_at":"2026-07-05T03:00:39.479965Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-05T03:00:39.491307Z","title":"Measuring coding challenge competence with APPS,","venue":null,"work_id":"eed75a04-3d74-464d-a75e-c3ea9d7f03d5","year":2021},"citing_paper":{"arxiv_id":"2607.01360","last_updated":"2026-07-01T18:20:27Z","snapshot_observed_at":"2026-08-15T16:04:59.762247Z","submitted_at":"2026-07-01T18:20:27Z","title":"Benchmarking Code Improvement with Progressive, Adaptive, and Interactive Feedback","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-07-03T19:23:50.577675Z"},"links":{"citing_paper":"/paper/2607.01360"},"observation_digest":"sha256:5591d9e0e6d603913cac303eb1e729542ab522ed701587da59d0615b456007b2","observation_id":"a305db9c-3781-4498-83f0-c12467b17b6c","resolution":{"observed_at":"2026-07-05T03:00:39.492617Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-05T03:00:39.482338Z","title":"Is your code generated by chatGPT really correct? rigorous evaluation of large language models for code generation,","venue":null,"work_id":"689aac1e-fd6f-4d6d-8094-d6f29d48e6b6","year":2023},"citing_paper":{"arxiv_id":"2607.01360","last_updated":"2026-07-01T18:20:27Z","snapshot_observed_at":"2026-08-15T16:04:59.762247Z","submitted_at":"2026-07-01T18:20:27Z","title":"Benchmarking Code Improvement with Progressive, Adaptive, and Interactive Feedback","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-07-03T19:23:50.577675Z"},"links":{"citing_paper":"/paper/2607.01360"},"observation_digest":"sha256:447be234c1473c2512b18086a55646f94c79b9c458d7167ae54b85da52c47210","observation_id":"c5ab4027-6899-41ec-8c59-51ec2b5bdba4","resolution":{"observed_at":"2026-07-05T03:00:39.483590Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-06T05:11:57.159799Z","title":"Evaluating language models for efficient code generation,","venue":null,"work_id":"a97d757b-88c9-4dd0-9b9d-a7aa9bc4bbb6","year":2024},"citing_paper":{"arxiv_id":"2607.01360","last_updated":"2026-07-01T18:20:27Z","snapshot_observed_at":"2026-08-15T16:04:59.762247Z","submitted_at":"2026-07-01T18:20:27Z","title":"Benchmarking Code Improvement with Progressive, Adaptive, and Interactive Feedback","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-07-03T19:23:50.577675Z"},"links":{"citing_paper":"/paper/2607.01360"},"observation_digest":"sha256:bf9c4ee7d54d15327b90c0e4de22c026be46d7be02a9139f9433db2ae74c3253","observation_id":"26ca6827-d294-488c-b529-6aad35550b00","resolution":{"observed_at":"2026-07-05T03:00:39.470869Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"0384.262805","doi":"10.1145/2610384.2628053","metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Ernst, Reid Holmes, and Gordon Fraser","venue":null,"work_id":"f1561e83-91e1-48d8-9809-6fbab03c1cfc","year":2014},"citing_paper":{"arxiv_id":"2607.01360","last_updated":"2026-07-01T18:20:27Z","snapshot_observed_at":"2026-08-15T16:04:59.762247Z","submitted_at":"2026-07-01T18:20:27Z","title":"Benchmarking Code Improvement with Progressive, Adaptive, and Interactive Feedback","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-07-03T19:23:50.577675Z"},"links":{"citing_paper":"/paper/2607.01360"},"observation_digest":"sha256:760e8670a65affd5f93e86dd6f08202ee265d5ee5d0719391c49d7c34f8547ed","observation_id":"1e886f20-d2a3-4955-b8b6-c837745430af","resolution":{"observed_at":"2026-07-03T19:28:51.399385Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"8089.341794","doi":"10.1145/3368089.3417940","metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Yiheng Xiong, Ting Su, Jue Wang, Jingling Sun, Geguang Pu, and Zhendong Su","venue":null,"work_id":"c5dd3bf4-64dc-4365-a833-cedf9668cfc2","year":2020},"citing_paper":{"arxiv_id":"2607.01360","last_updated":"2026-07-01T18:20:27Z","snapshot_observed_at":"2026-08-15T16:04:59.762247Z","submitted_at":"2026-07-01T18:20:27Z","title":"Benchmarking Code Improvement with Progressive, Adaptive, and Interactive Feedback","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-07-03T19:23:50.577675Z"},"links":{"citing_paper":"/paper/2607.01360"},"observation_digest":"sha256:32e03f147af04b1879ba6c590ad435d0b68d0be48e782a3acf07b95fb41f39d3","observation_id":"8e737cee-a16b-4f6d-aba8-7595810b137f","resolution":{"observed_at":"2026-07-03T19:28:51.397826Z","resolver_source":"arxiv_id","status":"metadata_mismatch"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"5932.313594","doi":"10.1145/3135932.3135941","metadata_source":"arxiv_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"Quixbugs: A multi-lingual program repair benchmark set based on the quixey challenge","venue":null,"work_id":"928a3af6-b762-49af-8359-d155c3cd0f4b","year":2017},"citing_paper":{"arxiv_id":"2607.01360","last_updated":"2026-07-01T18:20:27Z","snapshot_observed_at":"2026-08-15T16:04:59.762247Z","submitted_at":"2026-07-01T18:20:27Z","title":"Benchmarking Code Improvement with Progressive, Adaptive, and Interactive Feedback","version":1},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-07-03T19:23:50.577675Z"},"links":{"citing_paper":"/paper/2607.01360"},"observation_digest":"sha256:c6e9d96747b9c720b3c08d23af22f1319bd8060a410d95b69cafd8d0fbea1675","observation_id":"516ffc33-0f2f-4f51-bfdf-fefeb65ed89e","resolution":{"observed_at":"2026-07-03T19:28:51.392860Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-05T03:00:39.461176Z","title":"The power of feedback,","venue":null,"work_id":"10117617-9fde-4b3d-bd84-15af4652d485","year":2007},"citing_paper":{"arxiv_id":"2607.01360","last_updated":"2026-07-01T18:20:27Z","snapshot_observed_at":"2026-08-15T16:04:59.762247Z","submitted_at":"2026-07-01T18:20:27Z","title":"Benchmarking Code Improvement with Progressive, Adaptive, and Interactive Feedback","version":1},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-07-03T19:23:50.577675Z"},"links":{"citing_paper":"/paper/2607.01360"},"observation_digest":"sha256:0f6821dc15f3d94df171484cd65af3ddb8ca6bf001b1ae45533ce4b7fe013955","observation_id":"41db9097-3227-4ddb-8186-1431dc74e1fa","resolution":{"observed_at":"2026-07-05T03:00:39.462274Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-05T03:00:39.487693Z","title":"Focus on formative feedback,","venue":null,"work_id":"ff7a933a-dee1-41cf-bd64-8c8c9e36831a","year":2008},"citing_paper":{"arxiv_id":"2607.01360","last_updated":"2026-07-01T18:20:27Z","snapshot_observed_at":"2026-08-15T16:04:59.762247Z","submitted_at":"2026-07-01T18:20:27Z","title":"Benchmarking Code Improvement with Progressive, Adaptive, and Interactive Feedback","version":1},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-07-03T19:23:50.577675Z"},"links":{"citing_paper":"/paper/2607.01360"},"observation_digest":"sha256:c64abf03a8557b4b9ef3493dd0c2e4f76034efcbd1100a54aad34e7f7616cb4d","observation_id":"db73f802-0a32-4d85-8e48-d02ee76f06d7","resolution":{"observed_at":"2026-07-05T03:00:39.488965Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-08T17:55:21.242584Z","title":"The role of tutoring in problem solving","venue":null,"work_id":"41eb2c35-bec0-400e-a37d-04dfca25516a","year":1976},"citing_paper":{"arxiv_id":"2607.01360","last_updated":"2026-07-01T18:20:27Z","snapshot_observed_at":"2026-08-15T16:04:59.762247Z","submitted_at":"2026-07-01T18:20:27Z","title":"Benchmarking Code Improvement with Progressive, Adaptive, and Interactive Feedback","version":1},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-07-03T19:23:50.577675Z"},"links":{"citing_paper":"/paper/2607.01360"},"observation_digest":"sha256:78d210949db0ae7bbedb0a15b49416d4823c93ee6746dc2c3b58184624015ebe","observation_id":"ce7d2be9-fec2-4b9a-aa1d-144e0292bc4d","resolution":{"observed_at":"2026-07-05T03:00:39.463916Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-08T17:55:21.250376Z","title":null,"venue":null,"work_id":"9c188ba2-3e0b-4c63-beb7-cbd868c2c295","year":1978},"citing_paper":{"arxiv_id":"2607.01360","last_updated":"2026-07-01T18:20:27Z","snapshot_observed_at":"2026-08-15T16:04:59.762247Z","submitted_at":"2026-07-01T18:20:27Z","title":"Benchmarking Code Improvement with Progressive, Adaptive, and Interactive Feedback","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-07-03T19:23:50.577675Z"},"links":{"citing_paper":"/paper/2607.01360"},"observation_digest":"sha256:48e2aad440c9b9bc165247a02ef5b7fa997c91f98f75201a1f977a9df02d3917","observation_id":"4d17eba8-5ab9-41e5-86ab-e89874eab233","resolution":{"observed_at":"2026-07-05T03:00:39.481731Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-07-05T03:00:39.485963Z","title":"Codeforces-python-submissions,","venue":null,"work_id":"aa81a147-265a-44f2-9e7d-b3d4a33b2144","year":2024},"citing_paper":{"arxiv_id":"2607.01360","last_updated":"2026-07-01T18:20:27Z","snapshot_observed_at":"2026-08-15T16:04:59.762247Z","submitted_at":"2026-07-01T18:20:27Z","title":"Benchmarking Code Improvement with Progressive, Adaptive, and Interactive Feedback","version":1},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-07-03T19:23:50.577675Z"},"links":{"citing_paper":"/paper/2607.01360"},"observation_digest":"sha256:ef871b2dcfaaf368cccfb396686556c5ffc178429e338a888a8799df7249857e","observation_id":"2159dfc9-8ec1-479c-add0-a48e31fa5206","resolution":{"observed_at":"2026-07-05T03:00:39.487111Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-21T06:32:19.484+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2607.01360","last_updated":"2026-07-01T18:20:27Z","latest_version":1,"primary_category":"cs.SE","snapshot_observed_at":"2026-08-15T16:04:59.762247Z","submitted_at":"2026-07-01T18:20:27Z","title":"Benchmarking Code Improvement with Progressive, Adaptive, and Interactive Feedback"},"reference_resolution":{"displayed":31,"state_counts":{"malformed_identifier":0,"metadata_mismatch":5,"parse_uncertain":0,"unresolved":1,"verified_exact":7,"verified_fuzzy":18},"total_outbound_references":31},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-21T06:32:19.484+00:00","source":"crossref"},{"observed_at":"2026-08-21T06:32:16.066871+00:00","source":"retraction_watch"}],"thesis":"As of 21 August 2026, this Paper Citation Record lists 31 of 31 outbound references and 1 inbound Pith citation observation for arXiv:2607.01360."}