{"as_of":"2026-08-07T19:47:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:862b1b34c583783814d3476acf2f3cb67ffb4871c3f616c3f20d2391a945d897","coverage":[{"denominator":66,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":66,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-06T19:02:40.696853Z","state":"measured"},{"denominator":66,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":66,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-07T06:34:17.273281+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2608.04682/citation-record","integrity":"/paper/2608.04682/integrity","json":"/paper/2608.04682/citation-record.json","paper":"/paper/2608.04682"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:42.028193Z","title":"ICLR , year=","venue":null,"work_id":"95bc2946-81a7-45db-932e-9eb98778d7cc","year":null},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-07T19:20:24.139330Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":1,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.290931Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:532c033df676d4569e790076db3d91dfff531c31b808236dd0c172835749d402","observation_id":"0565f63a-f9be-4791-8c58-d2f804216905","resolution":{"observed_at":"2026-08-06T19:02:42.033982Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2302.13971","last_updated":"2023-02-27T17:11:15Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-02-27T17:11:15Z","title":"LLaMA: Open and Efficient Foundation Language Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2302.13971","snapshot_observed_at":"2026-08-06T19:02:40.308496Z","title":"arXiv preprint arXiv:2302.13971 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-07T19:20:24.139330Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":3,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.308496Z"},"links":{"cited_paper":"/paper/2302.13971","citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:3fe397307501e67edf59255c06808f950d52140561cb12709df656d5a0bfa552","observation_id":"49516cd8-6ba3-43f7-b0fa-143c03450ce4","resolution":{"observed_at":"2026-08-06T19:02:40.308496Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:42.010349Z","title":"ACM transactions on intelligent systems and technology , year=","venue":null,"work_id":"ed67f5a8-4f7f-41cb-a669-0bd940034f47","year":null},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-07T19:20:24.139330Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":4,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.313909Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:598efab8df85cb4e81fc473fa41beff5802b3f071c52a256af6d3c61c6a31991","observation_id":"b9a61cc0-4534-45ca-80c1-abe1104f57b9","resolution":{"observed_at":"2026-08-06T19:02:42.015769Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.989366Z","title":"Frontiers of Computer Science , year=","venue":null,"work_id":"f3328172-a2c1-406f-8a25-ff859d48a194","year":null},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-07T19:20:24.139330Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":5,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.319687Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:1ae0667e51b8f4011570b958ad7e0d716b2f3767545277fe501d9617302c5efc","observation_id":"dadae06c-eb74-4022-a9c4-f036a2014744","resolution":{"observed_at":"2026-08-06T19:02:41.994927Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:40.325309Z","title":"ICSE-FoSE , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-07T19:20:24.139330Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":6,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.325309Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:f6788f0ca69cee040705a759356bcce75ede56cc645fb53e9d4ef4e00bee3251","observation_id":"e81df24c-c8ba-4769-a303-c616ca9b4ce5","resolution":{"observed_at":"2026-08-06T19:02:40.325309Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.956015Z","title":"2024 , howpublished =","venue":null,"work_id":"14776416-a9fa-46d5-ae15-1f7e39242e07","year":2024},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-07T19:20:24.139330Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.336945Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:c76e27d451ec77bf363d0beb228c061a7937058b9d5b3ddc846ca7f178ed15f1","observation_id":"8d662c10-40b9-4050-970d-2e35e528b05c","resolution":{"observed_at":"2026-08-06T19:02:41.964031Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.938374Z","title":"2024 , howpublished =","venue":null,"work_id":"0ee82197-7525-4e6d-86d8-1e8e74f96509","year":2024},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-07T19:20:24.139330Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":9,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.342006Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:f56a3beff8c6a4706c73039a256f4e3cfc943dd0992a0734127658c185dc20b3","observation_id":"9115f5ca-b6f8-4f03-b6f3-a2b45d3ebe7d","resolution":{"observed_at":"2026-08-06T19:02:41.943464Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.921314Z","title":"2024 , howpublished =","venue":null,"work_id":"39e6b0b1-876c-456f-958b-86fafdefefe4","year":2024},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-07T19:20:24.139330Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":10,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.346986Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:a1f875e385a3309c1faf8066e802aaeab5fb3ffbefdec7fbb03c612cd8abd134","observation_id":"891edd44-be57-4a68-8a60-b294434e3997","resolution":{"observed_at":"2026-08-06T19:02:41.927060Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.903476Z","title":"2026 , howpublished =","venue":null,"work_id":"ca2daf4f-12f4-4e7c-a6ea-90cfbe418ce1","year":2026},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-07T19:20:24.139330Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.351934Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:3fe1aec2ca57a2cf69d0a0cc86b9a237eef0f762db0a359d6c51dc5462a877e1","observation_id":"c9d8cead-7ff6-4820-b947-b63811c2c720","resolution":{"observed_at":"2026-08-06T19:02:41.908802Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.884439Z","title":"2026 , howpublished =","venue":null,"work_id":"34e93196-43d0-4ec1-b6c3-f566fe067950","year":2026},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-07T19:20:24.139330Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":12,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.356669Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:44ac8aa4df5119d57104fd7bad0d361c4502f56ac7e70e2bfb1d9eaefae7b8c9","observation_id":"b30d7ecf-2130-4d80-9f68-d40884c106d2","resolution":{"observed_at":"2026-08-06T19:02:41.890389Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.864849Z","title":"NeurIPS , year=","venue":null,"work_id":"9b8b39d5-ba3b-45d2-b7db-827aeb39a94a","year":null},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-07T19:20:24.139330Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":13,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.361764Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:1ddb2f209b004a61947410b11883bb59ccd03affdf621fdf731bf67bed4dc8ea","observation_id":"85e92ddc-fa13-435f-915e-714e05b68d1c","resolution":{"observed_at":"2026-08-06T19:02:41.870937Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.846962Z","title":"NeurIPS , year=","venue":null,"work_id":"a188d8bf-f235-4ce3-b94a-e10707a713c5","year":null},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-07T19:20:24.139330Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":15,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.372026Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:731a0f2f638ef99e7f2289a875bf94af77bb0bee1974b66f88e6433c32c1e340","observation_id":"0a0e2002-0825-4760-9fee-7dac66b10620","resolution":{"observed_at":"2026-08-06T19:02:41.852951Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.06992","last_updated":"2024-10-10T13:13:09Z","snapshot_observed_at":"2026-08-05T03:15:38.088849Z","submitted_at":"2024-10-09T15:38:53Z","title":"SWE-Bench+: Enhanced Coding Benchmark for LLMs","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.06992","snapshot_observed_at":"2026-08-06T19:02:40.376652Z","title":"arXiv preprint arXiv:2410.06992 , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-07T19:20:24.139330Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":16,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.376652Z"},"links":{"cited_paper":"/paper/2410.06992","citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:d64bbb5dbee84b8d167094bdbb5968926a8e56e7e307c2fb8a6fcb6235477940","observation_id":"64ba4e39-b21a-4b0a-9c41-5cddd61a9a18","resolution":{"observed_at":"2026-08-06T19:02:40.376652Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.830062Z","title":"ACL , year=","venue":null,"work_id":"1a61c50c-a94c-465b-a57b-120d7809ac04","year":null},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-07T19:20:24.139330Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":17,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.382378Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:b5157bf18d2821a9e9b2d24412150dc7cc6fed2735aa988fbf613b3ed8babfbc","observation_id":"6a627b34-1f43-409d-81a5-4827a2338a3c","resolution":{"observed_at":"2026-08-06T19:02:41.835097Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.813383Z","title":"NeurIPS , year=","venue":null,"work_id":"f5b61f48-9a66-4f59-be1c-fd1159cfd6d5","year":null},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-07T19:20:24.139330Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":21,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.403082Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:f50bf3026f73545e829594ef89450d9e08c8b590014ae1b62415a66a7f9e1815","observation_id":"9a0930f9-0958-4d81-a76d-bd09a0406e57","resolution":{"observed_at":"2026-08-06T19:02:41.818320Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:40.413636Z","title":"Advances in Neural Information Processing Systems , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-07T19:20:24.139330Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":23,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.413636Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:04bb249db0adac3f5ace699b0ca2dbbef391ae1295b6149dd7b52a148f8032c9","observation_id":"83bd7b1c-0ffa-42f8-85af-8929782a01e3","resolution":{"observed_at":"2026-08-06T19:02:40.413636Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.784335Z","title":"ACL , year=","venue":null,"work_id":"1acb400d-2b03-4766-868a-b0b7bdf05a2c","year":null},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-07T19:20:24.139330Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":24,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.418361Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:e7a96c17489af512f6dca0830b6fb8c0b69d46a5917af6f98a20ec8e017a2d03","observation_id":"1cfb059a-d1d4-4f63-bb5a-695d361d32a7","resolution":{"observed_at":"2026-08-06T19:02:41.790730Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.766474Z","title":"ICLR , year=","venue":null,"work_id":"eda0e4b2-253d-4e54-9e2a-06bc48726214","year":null},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-07T19:20:24.139330Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":25,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.423733Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:fa1ea38741142e3b105a45a428ea04cbbb793f8999f638fcd9f5f84c3965c0e5","observation_id":"0619187b-1846-4afe-9b3b-eaa0d2de5f5a","resolution":{"observed_at":"2026-08-06T19:02:41.771010Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.749701Z","title":"ACL , year=","venue":null,"work_id":"db474a50-73ac-41a9-8777-85b1f10655fc","year":null},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-07T19:20:24.139330Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":26,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.428427Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:887acc0a3b2de69186a4b5661a34527b27f1991f89ccd0ccd085aca7ca6c7940","observation_id":"9a479890-d1b4-42f0-a840-d4b7fd6359c7","resolution":{"observed_at":"2026-08-06T19:02:41.754414Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.731244Z","title":"NeurIPS , year=","venue":null,"work_id":"b0a1b423-5145-4b6e-973b-c8c038a168bd","year":null},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-07T19:20:24.139330Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":27,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.433649Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:d1545c6fa901cc9ee753acbab5e5d5d7240a8c885fe270ce7a7acf71a2ef925d","observation_id":"d098cacf-1a18-4d77-99eb-6662078e2e95","resolution":{"observed_at":"2026-08-06T19:02:41.737619Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.709353Z","title":"2026 , howpublished =","venue":null,"work_id":"f1e7aa1c-dfa1-4ec8-a99f-0d1fcf700fb9","year":2026},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-07T19:20:24.139330Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":30,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.448812Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:ac462214c72459f92cc7cd394b69ec5e386295141855ab0c0fc14e67be16935d","observation_id":"6742c5c2-6174-4d6b-9714-f98aa2d79b7b","resolution":{"observed_at":"2026-08-06T19:02:41.714824Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.690635Z","title":"2026 , howpublished =","venue":null,"work_id":"bafc57fc-44c5-4d85-a3ac-6324ad51e34c","year":2026},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-07T19:20:24.139330Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":32,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.459486Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:aa85c38caa31278ca509bf0a82d9f4393d135082493e039578e278f256e98a4a","observation_id":"517c0c83-b80f-41fe-aaa4-0e8cc3177cfe","resolution":{"observed_at":"2026-08-06T19:02:41.696949Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.668298Z","title":"Qwen3.5: Accelerating Productivity with Native Multimodal Agents , howpublished =","venue":null,"work_id":"cf732404-8ad0-477e-851d-41bc56191711","year":null},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-07T19:20:24.139330Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":35,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.474687Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:f37b186d07fe7a250527bd3d035f89b7e6522614221be8b38355ab39c95a50a0","observation_id":"7fda5a46-3f01-4293-bf9e-341e083e61f5","resolution":{"observed_at":"2026-08-06T19:02:41.678185Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:40.479235Z","title":"2026 , howpublished =","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-07T19:20:24.139330Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":36,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.479235Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:477e1db4b5a42517a8705754a3391bd95f89f2019844a12a10d3e1cc45852fd3","observation_id":"d437dc11-78b9-4574-80e2-b06b942efee2","resolution":{"observed_at":"2026-08-06T19:02:40.479235Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.640573Z","title":"2026 , howpublished =","venue":null,"work_id":"ddf35601-d618-4fc7-841c-666443317641","year":2026},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-07T19:20:24.139330Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":37,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.484085Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:283f63d7087e621893357366ec0f60e47738cbce633f5b2236a15a0577c35dbf","observation_id":"ba7c3479-1fa3-49a1-a12a-6b79a7f87839","resolution":{"observed_at":"2026-08-06T19:02:41.645910Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.624139Z","title":"2026 , howpublished =","venue":null,"work_id":"55dee3de-6ff6-4d09-9994-1eaadbe4e419","year":2026},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-07T19:20:24.139330Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":38,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.489325Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:2d6dc1090ad720e32238e79c78d9453cc018ad5a6b4fc4ceb216393bac10c05a","observation_id":"dd58faa9-7394-48e2-81ce-f65cf8b8bdf0","resolution":{"observed_at":"2026-08-06T19:02:41.629172Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:40.509167Z","title":null,"venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-07T19:20:24.139330Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":42,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.509167Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:d3755e9bc1a049b99cc76cd97c4e3e3eb0fd8292d437ad0390c920a445e20d6d","observation_id":"b0c83089-c242-4d7b-bfbd-0785c326eb3b","resolution":{"observed_at":"2026-08-06T19:02:40.509167Z","resolver_source":null,"status":"parse_uncertain"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.596249Z","title":"System card: Claude sonnet 4.6","venue":null,"work_id":"665e9a3e-b727-4567-99ce-265d44c4fbd2","year":2026},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-07T19:20:24.139330Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":43,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.513795Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:126803c55722b10cba0c6444eb1eaf4ce8dac7d40074301b95782f910dad15dc","observation_id":"89c6e5dd-035d-4ee2-a1ce-842f3486327d","resolution":{"observed_at":"2026-08-06T19:02:41.601259Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.579587Z","title":"System card: Claude opus 4.8","venue":null,"work_id":"f29257e8-7b29-4df5-9477-a8cbc240a6e0","year":2026},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-07T19:20:24.139330Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":44,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.518702Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:639000975fc86d94a5c45d0c553db2cabb097f2397298d86147874ec44ee8985","observation_id":"06eb38fb-7222-444f-b0d7-7abd67c27903","resolution":{"observed_at":"2026-08-06T19:02:41.585132Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.563749Z","title":"A survey on evaluation of large language models","venue":null,"work_id":"eb9c5f75-b36b-4c11-91ae-3b02d80917e8","year":2024},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-07T19:20:24.139330Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":45,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.523406Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:4773982463f0bf2a12e4c342aa8c095e085d8419e933b6ac3e9566f2dce484d6","observation_id":"d0187244-64c2-4c7c-9476-308866a0ec00","resolution":{"observed_at":"2026-08-06T19:02:41.568599Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2605.26494","last_updated":"2026-07-30T05:04:36Z","snapshot_observed_at":"2026-08-06T03:12:19.331370Z","submitted_at":"2026-05-26T03:16:11Z","title":"The MiniMax-M2 Series: Mini Activations Unleashing Max Real-World Intelligence","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2605.26494","snapshot_observed_at":"2026-08-06T19:02:40.528271Z","title":"The minimax-m2 series: Mini activations unleashing max real-world intelligence","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-07T19:20:24.139330Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":46,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.528271Z"},"links":{"cited_paper":"/paper/2605.26494","citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:61a774f007169071aa1f836134856b5725c0e6f5ea9339fb947ff3ac4e82527f","observation_id":"72be4632-bdf0-49e8-82ed-74b805e6b46b","resolution":{"observed_at":"2026-08-06T19:02:40.528271Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.546902Z","title":"External technical root cause analysis — channel file 291","venue":null,"work_id":"8905016b-c670-4154-877b-1d86541c1a42","year":2024},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-07T19:20:24.139330Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":47,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.532761Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:4e692f50b0c286793d89473c0cdb9e3a0214d4fb1385e8c42ad5fb6a8e2387c3","observation_id":"6e8b870a-04a4-444f-b81c-ebaca12c6134","resolution":{"observed_at":"2026-08-06T19:02:41.552095Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.529610Z","title":"Deepseek-v4: Towards highly efficient million-token context intelligence","venue":null,"work_id":"0c2a254f-6c21-4d16-b098-633525427a40","year":2026},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-07T19:20:24.139330Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":48,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.537453Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:0d3dd61010cca06dca360d0c6fada8bd47209853f69c0d0abd28a09c719b6057","observation_id":"beb6dfa7-3e20-4817-85df-dfa4e8d756d7","resolution":{"observed_at":"2026-08-06T19:02:41.535235Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2509.16941","last_updated":"2025-11-14T22:00:03Z","snapshot_observed_at":"2026-08-06T21:34:46.041153Z","submitted_at":"2025-09-21T06:28:17Z","title":"SWE-Bench Pro: Can AI Agents Solve Long-Horizon Software Engineering Tasks?","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2509.16941","snapshot_observed_at":"2026-08-06T19:02:40.542646Z","title":"Swe-bench pro: Can ai agents solve long-horizon software engineering tasks? arXiv preprint arXiv:2509.16941, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-07T19:20:24.139330Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":49,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.542646Z"},"links":{"cited_paper":"/paper/2509.16941","citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:e0463b10f185dc7f5c13554e3f5c2ca9c42253ff65aeb5ed9aaeb6c0fd849c82","observation_id":"49bd99bc-1278-4d1d-ae47-14d10ac73729","resolution":{"observed_at":"2026-08-06T19:02:40.542646Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.512817Z","title":"Large language models for software engineering: Survey and open problems","venue":null,"work_id":"f40d98be-f376-4eb1-a32a-8436b9f21dbc","year":2023},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-07T19:20:24.139330Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":50,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.547534Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:7d84d938e5c1dec27642b9ed8f58dbcc76a33e12a7b1536b30727a50ab086c78","observation_id":"128b5bd0-da9f-491c-bb5a-2f80a28b5e2e","resolution":{"observed_at":"2026-08-06T19:02:41.518388Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.497271Z","title":"Gemini 3.1 pro model card","venue":null,"work_id":"61a2ee32-58e1-4e64-a9b9-2f70f652d209","year":2026},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-07T19:20:24.139330Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":51,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.552364Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:57deb56f24b22a72b9738d720b2df61f2c49da9bec5d4086101e6e782641d99f","observation_id":"d4703c88-aaa4-4104-9a54-5d37d6225d5c","resolution":{"observed_at":"2026-08-06T19:02:41.502015Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2312.13010","last_updated":"2024-05-24T11:47:24Z","snapshot_observed_at":"2026-07-06T17:05:56.282078Z","submitted_at":"2023-12-20T13:22:41Z","title":"AgentCoder: Multi-Agent-based Code Generation with Iterative Testing and Optimisation","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.13010","snapshot_observed_at":"2026-08-06T19:02:40.557097Z","title":"Agentcoder: Multi-agent-based code generation with iterative testing and optimisation","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-07T19:20:24.139330Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":52,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.557097Z"},"links":{"cited_paper":"/paper/2312.13010","citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:49afa7a205559e067c1fa858202ac17b54f4fb5067ea5b5f036e9f995021bfc8","observation_id":"904c9bb4-aaa4-46e1-9782-cd4833834b9d","resolution":{"observed_at":"2026-08-06T19:02:40.557097Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.481043Z","title":"Mapcoder: Multi-agent code generation for competitive problem solving","venue":null,"work_id":"ea0ec46d-8871-41a2-b1ee-117c9b04ee29","year":2024},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-07T19:20:24.139330Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":53,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.561483Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:a5524e74736395e2794ed7611dd132bd66f628f974b5e825ae324c5b808c31b8","observation_id":"bfed2013-810d-438f-979a-4fccabf85356","resolution":{"observed_at":"2026-08-06T19:02:41.486321Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.462801Z","title":"Swe-bench: Can language models resolve real-world github issues? In ICLR, 2024","venue":null,"work_id":"101e9a78-45b5-4a15-aca7-65515e89fe8c","year":2024},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-07T19:20:24.139330Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":54,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.566348Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:61d2df2d16919fe3f5ddf109474e5ae313265736e3cb6e596d4332f68a728b75","observation_id":"753f1b08-6ee4-468e-8039-99aa741f11c0","resolution":{"observed_at":"2026-08-06T19:02:41.468934Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2408.02479","last_updated":"2025-04-13T09:42:30Z","snapshot_observed_at":"2026-07-06T18:56:57.369673Z","submitted_at":"2024-08-05T14:01:15Z","title":"From LLMs to LLM-based Agents for Software Engineering: A Survey of Current, Challenges and Future","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2408.02479","snapshot_observed_at":"2026-08-06T19:02:40.571286Z","title":"From llms to llm-based agents for software engineering: A survey of current, challenges and future","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-07T19:20:24.139330Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":55,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.571286Z"},"links":{"cited_paper":"/paper/2408.02479","citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:9bc59ab202be04d197b0d55d2e9dbdfa795bd3eb8379bf1dfe7658fe8a834e5d","observation_id":"1b1f2c14-9151-43ab-bbf6-bc47a76ad8bf","resolution":{"observed_at":"2026-08-06T19:02:40.571286Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2606.15079","last_updated":"2026-06-13T03:21:49Z","snapshot_observed_at":"2026-07-06T23:52:28.557859Z","submitted_at":"2026-06-13T03:21:49Z","title":"Ling and Ring 2.6 Technical Report: Efficient and Instant Agentic Intelligence at Trillion-Parameter Scale","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2606.15079","snapshot_observed_at":"2026-08-06T19:02:40.575847Z","title":"Ling and ring 2.6 technical report: Efficient and instant agentic intelligence at trillion-parameter scale","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-07T19:20:24.139330Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":56,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.575847Z"},"links":{"cited_paper":"/paper/2606.15079","citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:c57f95dbb334475b9425f5d1a896a66260deb86ddf8fd4fb1e10991a795518c4","observation_id":"764ae076-5e9f-4c84-a758-364f9bfbb011","resolution":{"observed_at":"2026-08-06T19:02:40.575847Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":"2602.10471","doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:40.998397Z","title":"Testexplora: Benchmarking llms for proactive bug discovery via repository-level test generation","venue":null,"work_id":"92b91b58-d47f-45bd-b17f-1d4912341b58","year":2026},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-07T19:20:24.139330Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":57,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.580855Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:59684bf80145cd6945e4c765e988fed5fc8d9faa9c0dd8162eae48dbbc533f15","observation_id":"f5745e75-a471-43d7-ba00-7f27b1a049dc","resolution":{"observed_at":"2026-08-06T19:02:41.009629Z","resolver_source":"raw_fallback","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.446354Z","title":"Helping our customers through the crowdstrike outage","venue":null,"work_id":"44382069-3647-47db-8052-fb318fac8c5c","year":2024},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-07T19:20:24.139330Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":58,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.585756Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:7deaee68a54e055f908d1775f4d5a92f2805db691f43139ab0edbdb82f5a6188","observation_id":"60060c8c-6308-4b6d-abe2-43168cc9d958","resolution":{"observed_at":"2026-08-06T19:02:41.451448Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.429637Z","title":"u ndler, Mark N M \\","venue":null,"work_id":"1a022e9b-2db3-456d-b5c6-1370b594a626","year":2024},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-07T19:20:24.139330Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":59,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.591162Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:a2de0556de5812ffe4a1a79142dd8b4200dad750b155086f91b8da690ab60717","observation_id":"24b7c483-1b3f-4c62-a17b-e3ea15790a6d","resolution":{"observed_at":"2026-08-06T19:02:41.435745Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.413281Z","title":"Why swe-bench verified no longer measures frontier coding capabilities","venue":null,"work_id":"61a376c1-015b-437f-8797-304a6160b1a7","year":2026},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-07T19:20:24.139330Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":60,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.595627Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:7cf679555b5ba51e62ecc2f995b49b2c49b64088237a838871e3fd919007176e","observation_id":"fac8908f-b69b-4567-8be1-3f35844fff57","resolution":{"observed_at":"2026-08-06T19:02:41.418435Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.396668Z","title":"Separating signal from noise in coding evaluations","venue":null,"work_id":"ba9ab390-b7d1-4b02-b4e7-1e4904eb0560","year":2026},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-07T19:20:24.139330Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":61,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.600347Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:adb792d8531e4f81d8765217c175dd7e6b5ac5ab3650ed3aeeb965cb11874077","observation_id":"06da3624-803e-471c-8eb2-b598b707635a","resolution":{"observed_at":"2026-08-06T19:02:41.402190Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.377276Z","title":"Crowdstrike to cost fortune 500 \\ 5.4b; insured loss range of \\ 0.54b - \\ 1.08b","venue":null,"work_id":"111ae049-b901-4558-90a9-af626ba8eb7a","year":2024},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-07T19:20:24.139330Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":62,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.605278Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:bd53a6948b71bbce4375231c5ea8937344f31fa26f387687859d09399f00cb61","observation_id":"d813e5dd-8a68-49dd-abdd-1a9943392915","resolution":{"observed_at":"2026-08-06T19:02:41.383196Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.360783Z","title":"Qwen3.7 : The agent frontier, May 2026","venue":null,"work_id":"a25b81af-0b60-4c68-b803-3efd631ad40c","year":2026},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-07T19:20:24.139330Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":63,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.609842Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:88cd12a88f60e5ad53db018d5df6660474e5bca1915038217dc15ed19362ac08","observation_id":"25ba78ca-7e44-4a4c-a537-292814a882fc","resolution":{"observed_at":"2026-08-06T19:02:41.365835Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2605.08366","last_updated":"2026-05-08T18:21:44Z","snapshot_observed_at":"2026-07-06T23:20:38.385189Z","submitted_at":"2026-05-08T18:21:44Z","title":"SWE Atlas: Benchmarking Coding Agents Beyond Issue Resolution","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2605.08366","snapshot_observed_at":"2026-08-06T19:02:40.614800Z","title":"Swe atlas: Benchmarking coding agents beyond issue resolution","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-07T19:20:24.139330Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":64,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.614800Z"},"links":{"cited_paper":"/paper/2605.08366","citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:a43fd72ab4f5c284c616a7f70e8e2803176bc67a9dce47dbcfca6295f99bf175","observation_id":"36983989-f6c6-44e7-9295-216d3effa092","resolution":{"observed_at":"2026-08-06T19:02:40.614800Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2607.00248","last_updated":"2026-06-30T22:57:43Z","snapshot_observed_at":"2026-08-03T15:00:12.134373Z","submitted_at":"2026-06-30T22:57:43Z","title":"Seed2.0 Model Card: Towards Intelligence Frontier for Real-World Complexity","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2607.00248","snapshot_observed_at":"2026-08-06T19:02:40.619737Z","title":null,"venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-07T19:20:24.139330Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":65,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.619737Z"},"links":{"cited_paper":"/paper/2607.00248","citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:9b497a2cc37e7930f9571ab1e8464d4bca06e6eb147ac32af5315f2930c6d0b8","observation_id":"7e520ca1-e326-40d5-9a5d-33f4815a7fba","resolution":{"observed_at":"2026-08-06T19:02:40.619737Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2601.03267","last_updated":"2026-05-01T23:55:43Z","snapshot_observed_at":"2026-08-02T10:52:10.211700Z","submitted_at":"2025-12-19T07:05:38Z","title":"OpenAI GPT-5 System Card","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2601.03267","snapshot_observed_at":"2026-08-06T19:02:40.624433Z","title":"Openai gpt-5 system card","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-07T19:20:24.139330Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":66,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.624433Z"},"links":{"cited_paper":"/paper/2601.03267","citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:fa202472424f66bfca5a0a40b873ac7ef197a42fb2b31c7c81f1182cde6573f2","observation_id":"6bff20fc-1e23-499d-aef2-306784211d9f","resolution":{"observed_at":"2026-08-06T19:02:40.624433Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2602.02276","last_updated":"2026-02-02T16:17:38Z","snapshot_observed_at":"2026-07-06T22:44:09.804048Z","submitted_at":"2026-02-02T16:17:38Z","title":"Kimi K2.5: Visual Agentic Intelligence","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2602.02276","snapshot_observed_at":"2026-08-06T19:02:40.629242Z","title":null,"venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-07T19:20:24.139330Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":67,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.629242Z"},"links":{"cited_paper":"/paper/2602.02276","citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:ae7ef71e1fc71bdbaa131909cc1dcd42b55c390bd899ed14439ad502340e63b6","observation_id":"4166d0b0-61c3-460c-9d67-305dae3fa442","resolution":{"observed_at":"2026-08-06T19:02:40.629242Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.344226Z","title":"Qwen3.5: Accelerating productivity with native multimodal agents","venue":null,"work_id":"3aa5193a-d358-4e78-bc85-ea2bac52587d","year":2026},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-07T19:20:24.139330Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":68,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.633864Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:1ff12c957e82fcdd0e231f8376574d0ce71a095d7df5d74b2b9b9cd15f9f9ae0","observation_id":"888e5023-ee2f-4712-9181-fc809203c592","resolution":{"observed_at":"2026-08-06T19:02:41.349448Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.326988Z","title":null,"venue":null,"work_id":"537e14a4-282f-41c1-9b47-0d68a63b0819","year":2026},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-07T19:20:24.139330Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":69,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.638972Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:aad93884583f5092a2340ada29be7dd29c73fec882569d94dc3ff312d15c7a46","observation_id":"8d751679-314d-461e-9cb2-18c912352ff8","resolution":{"observed_at":"2026-08-06T19:02:41.332028Z","resolver_source":"raw_fallback","status":"unresolved"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.310448Z","title":"Openhands: An open platform for ai software developers as generalist agents","venue":null,"work_id":"e82b0409-6ffa-4a44-972a-2d4db4c08fe9","year":2025},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-07T19:20:24.139330Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":70,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.643906Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:e94d4abd2905c8ccd9d26132da6ae5958408d1d871122639902e72e9bd66cd68","observation_id":"4913b48a-561f-4138-9269-3c513df3004f","resolution":{"observed_at":"2026-08-06T19:02:41.315666Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:40.648923Z","title":"Live-swe-agent: Can software engineering agents self-evolve on the fly? arXiv preprint arXiv:2511.13646, 2025","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-07T19:20:24.139330Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":71,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.648923Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:4b01f0fe990b4678904d7caf7396806e1e49cfc79dd1b0355e86e86e9030cce6","observation_id":"c95108be-bc6a-4c3f-9943-0e1ad7feaf8f","resolution":{"observed_at":"2026-08-06T19:02:40.648923Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.293219Z","title":"Swe-agent: Agent-computer interfaces enable automated software engineering","venue":null,"work_id":"0d6af1be-81d9-4e1b-9817-a87ec64abf3e","year":2024},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-07T19:20:24.139330Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":72,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.653538Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:0b125458d03450bcb4826ceab35c230d8ea7a673e6431ab21c31ea1eddb65802","observation_id":"70896e99-4ea8-4a36-aaed-baf635d8ea05","resolution":{"observed_at":"2026-08-06T19:02:41.298970Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.03859","last_updated":"2024-10-04T18:48:58Z","snapshot_observed_at":"2026-07-06T19:28:08.374354Z","submitted_at":"2024-10-04T18:48:58Z","title":"SWE-bench Multimodal: Do AI Systems Generalize to Visual Software Domains?","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2410.03859","snapshot_observed_at":"2026-08-06T19:02:40.658366Z","title":"Swe-bench multimodal: Do ai systems generalize to visual software domains? arXiv preprint arXiv:2410.03859, 2024 b","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-07T19:20:24.139330Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":73,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.658366Z"},"links":{"cited_paper":"/paper/2410.03859","citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:3e53f9cc12ae5aa4fc4ee3a6645c2f0f61383cfadc173c04fbb7c7167aa2f198","observation_id":"474521f1-6871-4316-9674-9fed3c956337","resolution":{"observed_at":"2026-08-06T19:02:40.658366Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.275494Z","title":"Swe-smith: Scaling data for software engineering agents","venue":null,"work_id":"a2bbbae5-0de7-44c2-8cf5-003ff227f24e","year":2026},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-07T19:20:24.139330Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":74,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.662708Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:fd03f7849a3e5de403bec3676be9e11b03dd37e738b740ea45fc560b267e198a","observation_id":"b3543968-4e74-4956-8813-d4946b3df476","resolution":{"observed_at":"2026-08-06T19:02:41.280645Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2210.03629","last_updated":"2023-03-10T01:00:17Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2022-10-06T01:00:32Z","title":"ReAct: Synergizing Reasoning and Acting in Language Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2210.03629","snapshot_observed_at":"2026-08-06T19:02:40.667485Z","title":"React: Synergizing reasoning and acting in language models","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-07T19:20:24.139330Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":75,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.667485Z"},"links":{"cited_paper":"/paper/2210.03629","citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:dee0fceadb642cacd2936992fe71aa4dc94709ebed8eacf412eb03007fcc1ec9","observation_id":"f28180df-8257-44fc-aa16-5f74e3e2d7fe","resolution":{"observed_at":"2026-08-06T19:02:40.667485Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.258952Z","title":"Multi-swe-bench: A multilingual benchmark for issue resolving","venue":null,"work_id":"66172575-dd52-4e0d-aa5f-8af50d5c1163","year":2026},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-07T19:20:24.139330Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":76,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.672186Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:e7bf677fdf69e9e0e81a308422998d14cff1d666c25d04a6609e6f2061876eb5","observation_id":"1d73afb6-5baf-41b4-8e2e-f8662083c1a6","resolution":{"observed_at":"2026-08-06T19:02:41.264160Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2602.15763","last_updated":"2026-02-24T10:44:44Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2026-02-17T17:50:56Z","title":"GLM-5: from Vibe Coding to Agentic Engineering","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2602.15763","snapshot_observed_at":"2026-08-06T19:02:40.677154Z","title":"Glm-5: from vibe coding to agentic engineering","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-07T19:20:24.139330Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":77,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.677154Z"},"links":{"cited_paper":"/paper/2602.15763","citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:46706470a918d0c00f49e21bfcbd4ae3878390ecce31ddc858ce439b91d54473","observation_id":"9b55618c-b92d-4c42-a61b-1751e94b601e","resolution":{"observed_at":"2026-08-06T19:02:40.677154Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.241313Z","title":"Codeagent: Enhancing code generation with tool-integrated agent systems for real-world repo-level coding challenges","venue":null,"work_id":"169da89d-a4af-459d-a902-e2dd8dd405c3","year":2024},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-07T19:20:24.139330Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":78,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.681643Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:51120a74218499e899f7a3ae8473de71d8fcd2e7380553e2586db9ee612e8139","observation_id":"37d59c2f-702c-447e-a703-8d4c3d9a7959","resolution":{"observed_at":"2026-08-06T19:02:41.247502Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.224458Z","title":"Swe-bench goes live! In NeurIPS, 2026","venue":null,"work_id":"830b05dd-1e22-4148-8d5f-eb93ac5163ce","year":2026},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-07T19:20:24.139330Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":79,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.686763Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:e7a80d74d23a4ef96c921c0ce678f65548bd9549203763e0c7edc48b4fc47bf3","observation_id":"772dcec3-c004-4870-91db-1d8927a27f78","resolution":{"observed_at":"2026-08-06T19:02:41.229614Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:41.207255Z","title":"A survey of large language models","venue":null,"work_id":"8b5fb974-b079-4cc0-a371-57e0f08c9c77","year":2026},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-07T19:20:24.139330Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":80,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.691634Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:1287f607c3d36e23204a5464378926b30efd339a1726bae5cbe3ac34f4876d69","observation_id":"8eebe9aa-1454-4e4b-b045-82b5beb7fa9f","resolution":{"observed_at":"2026-08-06T19:02:41.212617Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-07T06:34:17.273281+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T19:02:40.696853Z","title":"Featurebench: Benchmarking agentic coding for complex feature development","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","snapshot_observed_at":"2026-08-07T19:20:24.139330Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports","version":1},"reference_index":81,"source":"arxiv_source","source_observed_at":"2026-08-06T19:02:40.696853Z"},"links":{"citing_paper":"/paper/2608.04682"},"observation_digest":"sha256:af36be82c39c358ffd202416048def7ad1a7bd0b64d650468ae545ce967a0859","observation_id":"8e9f7764-9d6f-4452-9c46-28fe032d255a","resolution":{"observed_at":"2026-08-06T19:02:40.696853Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2608.04682","last_updated":"2026-08-05T10:50:01Z","latest_version":1,"primary_category":"cs.SE","snapshot_observed_at":"2026-08-07T19:20:24.139330Z","submitted_at":"2026-08-05T10:50:01Z","title":"Active-SWE: Benchmarking Coding Agents for Proactive Bug Fixing without Issue Reports"},"reference_resolution":{"displayed":66,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":1,"unresolved":20,"verified_exact":1,"verified_fuzzy":44},"total_outbound_references":66},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"thesis":"As of 7 August 2026, this Paper Citation Record lists 66 of 66 outbound references and 0 inbound Pith citation observations for arXiv:2608.04682."}