{"as_of":"2026-08-07T07:11:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:bf0596ca88de8a16ec1e74bd2d643a6772569c1b4cd9d640787806eeb0e9103a","coverage":[{"denominator":100,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":100,"source":"paper_references, paper_reference_links","source_observed_at":"2026-07-31T04:48:47.456192Z","state":"measured"},{"denominator":100,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":100,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-07T06:34:17.273281+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2607.28520/citation-record","integrity":"/paper/2607.28520/integrity","json":"/paper/2607.28520/citation-record.json","paper":"/paper/2607.28520"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.191281Z","title":"Bowling , title =","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":1,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.191281Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:a50ca2ce31747886f913549bd25263a614ee5f84c13683ced853a877725da66e","observation_id":"c8fbee1f-737b-425a-ad86-6c37619fe01b","resolution":{"observed_at":"2026-07-31T04:48:47.191281Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.194494Z","title":"Artificial Intelligence and Statistics , pages=","venue":null,"work_id":null,"year":2009},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":2,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.194494Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:bda8b4c5ae01ea0d4edef65799a1378cd42bf3d005e250580645604bb6ed4dce","observation_id":"2145736d-1962-4b09-af5a-8b557e5cb668","resolution":{"observed_at":"2026-07-31T04:48:47.194494Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.197341Z","title":"ACM Transactions on Economics and Computation (TEAC) , volume=","venue":null,"work_id":null,"year":2015},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":3,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.197341Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:b0141a74378c1c3dbcbcc2fb5bb97697383fdec30170e1ac5fa1c554e1011ecd","observation_id":"bd731ca9-ab3f-4a14-86fd-86d4850ef2d7","resolution":{"observed_at":"2026-07-31T04:48:47.197341Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.200068Z","title":"Advances in neural information processing systems , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":4,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.200068Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:93799bc313aff2a6a5a474f7ddd64d7b42a02e93465ae2c170b88ff25724ac7a","observation_id":"9b0a1b1f-9fb7-4332-9360-2e0e91d20d94","resolution":{"observed_at":"2026-07-31T04:48:47.200068Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.206062Z","title":"Science , volume=","venue":null,"work_id":null,"year":2015},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":6,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.206062Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:5ccdbf04f46c7158281aac43da05660e72c9d7630b50f9512736cc1cbb488701","observation_id":"62304136-5b85-414d-8045-0003d8ad6de4","resolution":{"observed_at":"2026-07-31T04:48:47.206062Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.208931Z","title":"Advances in Neural Information Processing Systems , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":7,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.208931Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:9b5c8026306c0fc673e68f4f36480edea118441de1209f471ef0d910ba8768fb","observation_id":"02307224-e046-4fe3-908c-13a0167ce6a4","resolution":{"observed_at":"2026-07-31T04:48:47.208931Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.211562Z","title":"International Conference on Machine Learning , pages=","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":8,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.211562Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:c8f188d2b00c458fa150a1bda6684bff02eea788d66972d4c518c43efaa81993","observation_id":"f35b3d1e-57df-4231-bcec-b3bd1debbfca","resolution":{"observed_at":"2026-07-31T04:48:47.211562Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.214270Z","title":"Proceedings of the 24th International Conference on Autonomous Agents and Multiagent Systems , pages=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":9,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.214270Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:97230b83ea1af172c5a9dfe290b1e10bed5ca2ece95fb9e8a9bffe69a23ea3d9","observation_id":"c8013b70-776d-41f9-bf38-4beb10b7945a","resolution":{"observed_at":"2026-07-31T04:48:47.214270Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.217252Z","title":"Forty-first International Conference on Machine Learning , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":10,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.217252Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:c4f497e84dd0c13a2536e1a405eece95450d05e187983fbe33a568806653a792","observation_id":"6785c2f6-11e4-4b7a-875f-36b50233a49b","resolution":{"observed_at":"2026-07-31T04:48:47.217252Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.220009Z","title":"Science , volume=","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":11,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.220009Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:d86efcfe9add79a3f4ac2597a8c7e9005fb2377dce744307902d84d97b96ace8","observation_id":"f57299c5-decc-4f89-9663-84beb93cf073","resolution":{"observed_at":"2026-07-31T04:48:47.220009Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.222502Z","title":"Science , volume=","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":12,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.222502Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:a59e8f21818f82abfeedefbde05b7e3ee9f989a4d537e543602411a5010d9388","observation_id":"362d9309-1580-44e5-b153-b0d2c7301133","resolution":{"observed_at":"2026-07-31T04:48:47.222502Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.225137Z","title":"Science , volume=","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":13,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.225137Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:20da8750b351e69b0fdbfea834416c59501bb1e423c231c543c26a0c7d1d0495","observation_id":"67ae75b7-fbb2-48cc-944c-dd98fda1b9d1","resolution":{"observed_at":"2026-07-31T04:48:47.225137Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.227481Z","title":"Advances in neural information processing systems , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":14,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.227481Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:e12dd89fb954cd653b8974efa0c7366da5ce141824a740468a6bfd223dbc8284","observation_id":"5ccb8e8a-207f-4898-96db-518d6021f182","resolution":{"observed_at":"2026-07-31T04:48:47.227481Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.229923Z","title":"Advances in neural information processing systems , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":15,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.229923Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:d00de9c3033c793744e100e8546ba706a9a988d0a4db31053c4aa5486118630c","observation_id":"a1486383-e520-4bb3-8d31-e0748df522a2","resolution":{"observed_at":"2026-07-31T04:48:47.229923Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.232290Z","title":"Advances in neural information processing systems , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":16,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.232290Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:c6563cf3a6d63c7bcb07f4d49f21bddc0b003d92d59eab9f5027855afada3e21","observation_id":"428bda40-ef63-4edf-8a60-66a02719f73a","resolution":{"observed_at":"2026-07-31T04:48:47.232290Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.234761Z","title":"Proceedings of the AAAI conference on artificial intelligence , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":17,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.234761Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:55a319885d309d52c9bcfef8e08aacc820b34e4199e9333bc69a8a0d9a8ebd7a","observation_id":"37b97009-bb19-48f1-b405-6ab69f6422c7","resolution":{"observed_at":"2026-07-31T04:48:47.234761Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.237440Z","title":"Proceedings of the AAAI Conference on Artificial Intelligence , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":18,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.237440Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:cab46fc4d9db9a6eb17b91460b555401b30ba261106b416228bab52c287a0138","observation_id":"d14a0dbe-c1b4-41d6-9766-f1aba462be1c","resolution":{"observed_at":"2026-07-31T04:48:47.237440Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.239726Z","title":"International conference on machine learning , pages=","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":19,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.239726Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:3c143a4ec61197e3714405eda40959be43f479c3144b405424d4435678aa9c60","observation_id":"e45043cf-f62d-40f8-b8d4-212a15a9dc23","resolution":{"observed_at":"2026-07-31T04:48:47.239726Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.242214Z","title":"Advances in neural information processing systems , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":20,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.242214Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:73330080a351f77b7430b2d3f2b9b7908946b8ff1efb724df710ef5e0dea7b67","observation_id":"f3bb24ad-8208-44f8-a28b-ad3f9d28fac2","resolution":{"observed_at":"2026-07-31T04:48:47.242214Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.244547Z","title":"Proceedings of the Twenty-First Conference on Uncertainty in Artificial Intelligence , pages=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":21,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.244547Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:dc29f43b31367423a8e9b53a5854d5bf72f0b0199c0fdf4297db04fe7919e52a","observation_id":"53e1ba60-9f8e-4c1b-8610-efc8ab4557d8","resolution":{"observed_at":"2026-07-31T04:48:47.244547Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.246808Z","title":"The Annals of Statistics , volume=","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":22,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.246808Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:398ff6294edbdaf6e859178c4389b5a487d5b6fe2cf662f222191ccfdb3172c2","observation_id":"e15db01b-900b-4a20-943b-67523188fba3","resolution":{"observed_at":"2026-07-31T04:48:47.246808Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.249155Z","title":"Journal of the Royal Statistical Society Series B: Statistical Methodology , volume=","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":23,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.249155Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:ac94d0b41ea0a4a8c2bf1d757e402527e98c8cf424a1a73048f983d03ccc635e","observation_id":"e023aea7-1616-40d8-85e7-fada2b8a44e8","resolution":{"observed_at":"2026-07-31T04:48:47.249155Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.256765Z","title":"CoRR , volume =","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":26,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.256765Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:1d1f126fc4173ef01b1f3d75dab1644e5982b69dcd58c3eeebd15bca37a4abdf","observation_id":"20423ccb-0642-4200-bc96-65235570d0c2","resolution":{"observed_at":"2026-07-31T04:48:47.256765Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.259014Z","title":"The International FLAIRS Conference Proceedings , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":27,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.259014Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:c65c1bf4bcda0cc9f33a0c3c5c7ea0650868389df595e01d6afd7ec027a2dd94","observation_id":"51a16d91-a2b0-42b6-9dd8-0403a9e60647","resolution":{"observed_at":"2026-07-31T04:48:47.259014Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.261397Z","title":"AAAI Technical Report (2) , pages=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":28,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.261397Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:75aeff4c2704f77bfd5bbded4d96e718244edb86849ccc2bffc55cf03361f629","observation_id":"0456e291-f538-4392-8ecb-2e3c6dcfeb0a","resolution":{"observed_at":"2026-07-31T04:48:47.261397Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.263661Z","title":"AAAI , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":29,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.263661Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:01c3982022567930dc25a620983acb0ab57c0ccde58c5de4a16d97b15e89c4c1","observation_id":"3c69ef12-e3b2-4fdd-93f8-8cfb998b0232","resolution":{"observed_at":"2026-07-31T04:48:47.263661Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.266175Z","title":"Proceedings of the 3rd AAAI Conference on Interactive Decision Theory and Game Theory , pages=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":30,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.266175Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:ee574206e30096389ea4323d639cb730396b2bba0797c824866a793b0334f72f","observation_id":"1f9d38ea-fbdf-43e2-9cfd-6c347751d246","resolution":{"observed_at":"2026-07-31T04:48:47.266175Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.268527Z","title":"Proceedings of the 41st International Conference on Machine Learning , pages=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":31,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.268527Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:8c75170ebf6fdbba15f291632f2b140cdcb8bcffaceec94d67d8f4ae4c973eb2","observation_id":"8fc4e63e-b87e-42de-a790-eacefd6577a3","resolution":{"observed_at":"2026-07-31T04:48:47.268527Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.270806Z","title":"The Thirteenth International Conference on Learning Representations , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":32,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.270806Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:8a825de1d95b748d59eeaba3f1e561b0b85007a55a0030df0e7c4307e4b4413d","observation_id":"7c4d4408-2f6d-4a37-b544-57ff779a07fa","resolution":{"observed_at":"2026-07-31T04:48:47.270806Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.281175Z","title":"Proceedings of the AAAI Conference on Artificial Intelligence , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":36,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.281175Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:6da5dba61a6653e04aa43c9726100055655b9311d7faba79259d5a421c9cf054","observation_id":"9ba0f7d9-892c-4a51-8578-4d3ff79a4116","resolution":{"observed_at":"2026-07-31T04:48:47.281175Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.283487Z","title":"Advances in Neural Information Processing Systems , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":37,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.283487Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:0b41bd370f8d59018bc799c38603b08ff41485ffa2e647b6fdbdb9e80470e45e","observation_id":"ee395685-1186-400f-b78b-2fa10ac078bf","resolution":{"observed_at":"2026-07-31T04:48:47.283487Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.285735Z","title":"IEEE Transactions on Cybernetics , volume=","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":38,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.285735Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:a836a75fd9e52b7d4cfe02220d97f98d9fc2e4ed08591ef60e1e748e919887f0","observation_id":"4e21a80e-dd8f-44ab-bbc3-aa0a09387974","resolution":{"observed_at":"2026-07-31T04:48:47.285735Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.288232Z","title":"Journal of Artificial Intelligence Research , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":39,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.288232Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:563ce69a7294be4ae31fa645eb977973a57c454bfcb7537a8f6fb18d3f206454","observation_id":"918a4a66-ccb3-49f1-9971-33148f8303a8","resolution":{"observed_at":"2026-07-31T04:48:47.288232Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.290482Z","title":"Advances in Neural Information Processing Systems , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":40,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.290482Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:81c9fd5951b7fe2fd2aa19653bfb641365b1ccbcdd565acb41c0341c73895542","observation_id":"029a4655-f316-4d29-9e85-79fc2a56d6c1","resolution":{"observed_at":"2026-07-31T04:48:47.290482Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.292873Z","title":"Advances in Neural Information Processing Systems , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":41,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.292873Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:3df3c35a0a1700da9deb24b22d1b7812b051611d36da1ed99a8bfd7ce84b02e8","observation_id":"42ac89b8-38f1-4559-80c0-11afd659abc8","resolution":{"observed_at":"2026-07-31T04:48:47.292873Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.295160Z","title":"The Twelfth International Conference on Learning Representations , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":42,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.295160Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:4bdd03664f500cfe5d9b1bbb40943dca312b05958e06750eb68210ff245b8618","observation_id":"be395faf-b533-4731-908f-f64cc2bb686d","resolution":{"observed_at":"2026-07-31T04:48:47.295160Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.297482Z","title":"International Conference on Machine Learning , pages=","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":43,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.297482Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:09658601b290c8aa4c03eb697f13884705743b6233e0a09842aca3efbe64430b","observation_id":"e526642d-2566-4375-8cdf-f94d4a5aa6d5","resolution":{"observed_at":"2026-07-31T04:48:47.297482Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.299690Z","title":"Proceedings of the AAAI Conference on Artificial Intelligence , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":44,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.299690Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:7d74ab692bbd49a54e6e0bb5e63a93a75231166f75b845d79789f5fe7ce7bdc2","observation_id":"a6fc2c57-292a-489f-985c-eac2cb1b41f8","resolution":{"observed_at":"2026-07-31T04:48:47.299690Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.301990Z","title":"The Twelfth International Conference on Learning Representations , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":45,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.301990Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:c57a7edd1f34cb92e6f5b0a6c2cb180e025c35a0d0054daf8059f16c4f90cf1f","observation_id":"4f548956-5396-4d64-b3d5-6e059cccb14e","resolution":{"observed_at":"2026-07-31T04:48:47.301990Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.304302Z","title":"Advances in Neural Information Processing Systems , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":46,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.304302Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:213ec31c53eb84bde9a9d8c21a6e3b57f70b12f8641dae4121af14973edc9bea","observation_id":"b217de98-af45-46eb-8576-3f20d746ac52","resolution":{"observed_at":"2026-07-31T04:48:47.304302Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.306782Z","title":"Expert Systems with Applications , volume=","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":47,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.306782Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:1c49062e63707da8f7fcbca0ef527cd4df49c760734b85c88666dbb3a799cd18","observation_id":"efb569d8-74fc-4385-90fe-17d39d556fe2","resolution":{"observed_at":"2026-07-31T04:48:47.306782Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.309069Z","title":"Proceedings of the AAAI Conference on Artificial Intelligence , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":48,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.309069Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:a1d2665c9587bfb7b7a65d2a6338f21292d768483f8809cc48af09dfe040532f","observation_id":"b3fdd216-3594-424e-a048-bad677e3f6c6","resolution":{"observed_at":"2026-07-31T04:48:47.309069Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.312132Z","title":"The Thirty-ninth Annual Conference on Neural Information Processing Systems , year=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":49,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.312132Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:b4626aae318be14f6a823864c0e006a2016f7d27a247c53594640002246f17dc","observation_id":"51160048-8a13-404e-8bfc-9ddf938d701f","resolution":{"observed_at":"2026-07-31T04:48:47.312132Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.314571Z","title":"Proceedings of the AAAI Conference on Artificial Intelligence , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":50,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.314571Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:321bcb26945584bd85de1b2d31f6bd0dafaba72758ae00f1af8292a66ef6df9d","observation_id":"e576e758-459c-462f-b370-19bb22c93c66","resolution":{"observed_at":"2026-07-31T04:48:47.314571Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.316953Z","title":"Proceedings of the AAAI Conference on Artificial Intelligence , volume=","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":51,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.316953Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:4302a84780aa4dca5e933b617314c73ba90c6a986df042978e61c47a5f686f71","observation_id":"c32e7f24-c8b6-4096-b506-5a3cdf311680","resolution":{"observed_at":"2026-07-31T04:48:47.316953Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.324527Z","title":"2026 , eprint=","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":54,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.324527Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:79a96f77cb32f8e73ceb434940ca9adbb1ef395198ce34d131b2775888659617","observation_id":"07816d28-fb5e-4ff3-b86d-760688954725","resolution":{"observed_at":"2026-07-31T04:48:47.324527Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.327029Z","title":"Exploiting opponents under utility constraints in sequential games","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":55,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.327029Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:365064e8b4ac4517fd3a96516040eb642cdb3c099294789d84b997c33f58b6a0","observation_id":"bf9eca4a-bf34-4fc1-8e1b-802d62661485","resolution":{"observed_at":"2026-07-31T04:48:47.327029Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.329452Z","title":"Heads-up limit hold'em poker is solved","venue":null,"work_id":null,"year":2015},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":56,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.329452Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:86380171654bac0d8d0e52f1676e19bf827bddf42cc4e2caea9e46bed7f5a6a1","observation_id":"6ab6f38b-0fa6-466d-a44f-699d93970a83","resolution":{"observed_at":"2026-07-31T04:48:47.329452Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.331971Z","title":"Regret-based pruning in extensive-form games","venue":null,"work_id":null,"year":2015},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":57,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.331971Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:f222d3d488533d04193d12e4373692398820891588972a0a71bf8db4036c3715","observation_id":"4e7d203f-8794-4b9d-8c25-97291aceaf8a","resolution":{"observed_at":"2026-07-31T04:48:47.331971Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.334352Z","title":"Safe and nested subgame solving for imperfect-information games","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":58,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.334352Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:45ae74b0a0adf8f9868cbf966b4bd67482ec4a2c854abc3de9b9dbd8bf87d2c6","observation_id":"9f13f61f-2533-42ae-a566-0c343c97eb7b","resolution":{"observed_at":"2026-07-31T04:48:47.334352Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.336639Z","title":"Superhuman ai for heads-up no-limit poker: Libratus beats top professionals","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":59,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.336639Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:1ad27559f83f7819e20390322617d0c61319b90e142fe2aa0ce6f32934f99c74","observation_id":"a5cd1d36-fa7d-4dc9-ba8f-14a91cedccbb","resolution":{"observed_at":"2026-07-31T04:48:47.336639Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.338989Z","title":"Solving imperfect-information games via discounted regret minimization","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":60,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.338989Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:b206537907ea2f3a55884cd6250f2f8a509b778cbe5eb61cdecba3d48d884147","observation_id":"d21ca11a-5ea9-443d-b96f-df986a320d57","resolution":{"observed_at":"2026-07-31T04:48:47.338989Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.341388Z","title":"Superhuman ai for multiplayer poker","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":61,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.341388Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:f2b730376e1c7c7e8a1ddaa8cb2e8756b003a12d36ec35a49613d301ec07ca86","observation_id":"1684e01c-9738-49e1-8f5c-a57893999323","resolution":{"observed_at":"2026-07-31T04:48:47.341388Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.343696Z","title":"Dynamic thresholding and pruning for regret minimization","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":62,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.343696Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:974971001e1c87504283e5172aa2dd18c54ebb6a366e95c1cfeb4843a481de76","observation_id":"5fbdcc35-499d-4085-a2be-12501ad5fbd7","resolution":{"observed_at":"2026-07-31T04:48:47.343696Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.345938Z","title":"Depth-limited solving for imperfect-information games","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":63,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.345938Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:96e95230e785e0cdfcb94e932d8df93080f01759f537a970bb7460d0c8ab2092","observation_id":"8d7a1f34-a9ff-4c3f-aa2d-f98b170f88b8","resolution":{"observed_at":"2026-07-31T04:48:47.345938Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.348215Z","title":"Deep counterfactual regret minimization","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":64,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.348215Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:1eb79d17f9754254ecd831cd95e38659ab10049ae556e93123328810538f147d","observation_id":"13f9e895-270f-44cc-a357-78ad9a6dbc50","resolution":{"observed_at":"2026-07-31T04:48:47.348215Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.350681Z","title":"Combining deep reinforcement learning and search for imperfect-information games","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":65,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.350681Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:87565484d551fa46fb716a0637cd86c5b1a262983feb2c4305676746dc95bea8","observation_id":"e25a7da1-7264-4780-a207-564d029ca8ef","resolution":{"observed_at":"2026-07-31T04:48:47.350681Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2604.25796","last_updated":"2026-04-28T16:03:14Z","snapshot_observed_at":"2026-07-06T23:11:34.854305Z","submitted_at":"2026-04-28T16:03:14Z","title":"StratFormer: Adaptive Opponent Modeling and Exploitation in Imperfect-Information Games","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2604.25796","snapshot_observed_at":"2026-07-31T04:48:47.352939Z","title":"Stratformer: Adaptive opponent modeling and exploitation in imperfect-information games","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":66,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.352939Z"},"links":{"cited_paper":"/paper/2604.25796","citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:fde905ab99a1302d0f6dd751c679d178a42ad8e5340964730adda4dc4882fdbf","observation_id":"60860922-b7b2-4aa0-951b-169d26ec1b3e","resolution":{"observed_at":"2026-07-31T04:48:47.352939Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.355340Z","title":"Test-then-punish: A statistical approach to repeated games","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":67,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.355340Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:ce591208c850db9cb150a506a788ef95e061bfc223f77a902d0b2a0bbce547ca","observation_id":"fb627b75-bb06-48b7-b737-14560d1066f6","resolution":{"observed_at":"2026-07-31T04:48:47.355340Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.357580Z","title":"Faster game solving via predictive blackwell approachability: Connecting regret matching and mirror descent","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":68,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.357580Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:d0a88e2a0192b47cc5c03b6cec1dca22b0392983875c5682fca2e836d73b6d2d","observation_id":"9a7f5ecc-47df-4264-8e7e-a84f97a047f6","resolution":{"observed_at":"2026-07-31T04:48:47.357580Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.359894Z","title":"Regret matching+:(in) stability and fast convergence in games","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":69,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.359894Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:b43ad543ce8170d1bf3ecba1153f01fe7204739480bd1b1cbe0d8187f4f7c390","observation_id":"5909783a-a95a-4973-ade1-106ebb0a8bbb","resolution":{"observed_at":"2026-07-31T04:48:47.359894Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.362064Z","title":"Greedy when sure and conservative when uncertain about the opponents","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":70,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.362064Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:dfb962d3e831ab31f1cd43532b4a125abbde8de1b04db9804c09b97aeeae8ac6","observation_id":"44dc54f6-4101-4a0c-a91b-cb1d52494560","resolution":{"observed_at":"2026-07-31T04:48:47.362064Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2508.17671","last_updated":"2026-07-26T21:15:12Z","snapshot_observed_at":"2026-08-07T06:10:50.682271Z","submitted_at":"2025-08-25T05:08:49Z","title":"Consistent Opponent Modeling in Imperfect-Information Games","version":8},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2508.17671","snapshot_observed_at":"2026-07-31T04:48:47.364343Z","title":"Consistent opponent modeling of static opponents in imperfect-information games","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":71,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.364343Z"},"links":{"cited_paper":"/paper/2508.17671","citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:ff9290a3dd88a00a21d9b3044556f750dc4831554d0e908582cf6ce56b1094ea","observation_id":"32afac23-d830-4e68-9d3f-0aa4f6c1692f","resolution":{"observed_at":"2026-07-31T04:48:47.364343Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.366877Z","title":"Nonparametric strategy test","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":72,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.366877Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:a47bb977f576ee05863d67e435b78e928a1e172a0caa0cce955c715433a4d184","observation_id":"3055ca0f-38f2-4789-a549-2d6f37f7d6a6","resolution":{"observed_at":"2026-07-31T04:48:47.366877Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.369505Z","title":"Safe opponent exploitation","venue":null,"work_id":null,"year":2015},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":73,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.369505Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:ca65656dc612762f2f4b8792c21c959799603a29f82d45c52df61d4906674d10","observation_id":"6126bee9-c13b-4d0f-bd5b-26e8540ecbab","resolution":{"observed_at":"2026-07-31T04:48:47.369505Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2601.05427","last_updated":"2026-05-22T16:37:56Z","snapshot_observed_at":"2026-07-31T05:26:20.127534Z","submitted_at":"2026-01-08T23:26:05Z","title":"Anytime Detection of Strategic Deviations in Multi-Agent Systems","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2601.05427","snapshot_observed_at":"2026-07-31T04:48:47.372071Z","title":"Betting on equilibrium: Monitoring strategic behavior in multi-agent systems","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":74,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.372071Z"},"links":{"cited_paper":"/paper/2601.05427","citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:9287e1414860307a928ef2300bd17f3a88a9f576dcaf75aab8d6c635d0107604","observation_id":"509423cb-1d88-4611-9e23-675f6ef702ac","resolution":{"observed_at":"2026-07-31T04:48:47.372071Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.374449Z","title":"Modeling rationality: Toward better performance against unknown agents in sequential games","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":75,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.374449Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:718466f826e63f402f356813c2655ccb50a15cfc2510870cb567da0c5cbe737e","observation_id":"16e03f78-fcd7-4305-9560-2460f4685c92","resolution":{"observed_at":"2026-07-31T04:48:47.374449Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.376811Z","title":"Efficient subgame refinement for extensive-form games","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":76,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.376811Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:61ffa648f37082c654cc9686cc08ea996b9fa79fe33ea60b909a22a03cc16ba9","observation_id":"be6533dc-f740-4fee-a582-d6ed481f9e3d","resolution":{"observed_at":"2026-07-31T04:48:47.376811Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.379022Z","title":"Safe and robust subgame exploitation in imperfect information games","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":77,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.379022Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:c960a5587bc8bfc9b388ec58b6b473cc4ea9f4596ff80ba62a9ad6d7ea83b7e0","observation_id":"c4e12683-7496-4807-8c67-4193e47ef89c","resolution":{"observed_at":"2026-07-31T04:48:47.379022Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.381403Z","title":"Proposer of the vote of thanks to waudy-smith and ramdas and contribution to the discussion of `estimating means of bounded random variables by betting'","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":78,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.381403Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:c8f470e0cdc4d28cc32ba71b786cf0befe4c50c75d32ecb720f0fef596676c38","observation_id":"67d7222f-ba3a-4cea-bb71-d285ad555004","resolution":{"observed_at":"2026-07-31T04:48:47.381403Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.383750Z","title":"Effective short-term opponent exploitation in simplified poker","venue":null,"work_id":null,"year":2005},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":79,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.383750Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:abc709a0885719fec94684bdd1979c2ebf02a4ef3843ede5340b820347c33933","observation_id":"559cedc7-0e83-4b73-a813-ea21725b0868","resolution":{"observed_at":"2026-07-31T04:48:47.383750Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.386236Z","title":"Time-uniform, nonparametric, nonasymptotic confidence sequences","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":80,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.386236Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:8328e541d2a44946b718b3e545ec1d2532d280be1841d36aa3a969d1d64c2070","observation_id":"3a51915e-9bd6-4e5d-8720-da5cfdc9e1a6","resolution":{"observed_at":"2026-07-31T04:48:47.386236Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.388962Z","title":"Towards offline opponent modeling with in-context learning","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":81,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.388962Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:062782da10cbdddd1db58a826f2c792b30997104032c6597b299f654a4935cc3","observation_id":"cd9039aa-61de-4b44-b3dd-b54516cbabc1","resolution":{"observed_at":"2026-07-31T04:48:47.388962Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.391267Z","title":"Opponent modeling with in-context search","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":82,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.391267Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:44e6306e1675db80d54925b6ca998d1fba416bf1050f9c0a76a47e423b1bf98f","observation_id":"8bfc98e1-b0a4-4def-a576-b1d0c2be4cc9","resolution":{"observed_at":"2026-07-31T04:48:47.391267Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.393533Z","title":"An open-ended learning framework for opponent modeling","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":83,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.393533Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:f9e230c66445267bbe8986bf0e0d797028c1eab0ba26a7cec8c7ba92ef4e4534","observation_id":"640c34b6-a159-44b7-a408-fdf24e3bbc46","resolution":{"observed_at":"2026-07-31T04:48:47.393533Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.395918Z","title":"Data biased robust counter strategies","venue":null,"work_id":null,"year":2009},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":84,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.395918Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:0ae3cad78baf0595b27a6ff209468cd90ac3a1967773366e73501e93056678f8","observation_id":"9a35002d-e7f6-4fe9-ab91-65eb4269029d","resolution":{"observed_at":"2026-07-31T04:48:47.395918Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.398332Z","title":null,"venue":null,"work_id":null,"year":2007},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":85,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.398332Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:2bdf5759570e3e71d6afa9bdfaea27e23e63b842ad890056e4ab6edcf6b818aa","observation_id":"a4d838ec-50b9-4980-8756-824582a1de83","resolution":{"observed_at":"2026-07-31T04:48:47.398332Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.400828Z","title":"Efficient online pruning and abstraction for imperfect information extensive-form games","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":86,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.400828Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:ffd3fc3845f0115dda2acb0e7193d9b50b1d018ac31a02252fe3b6f4043e1167","observation_id":"59837463-2886-45da-8cb4-68a99c471fa9","resolution":{"observed_at":"2026-07-31T04:48:47.400828Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2605.10900","last_updated":"2026-05-11T17:41:16Z","snapshot_observed_at":"2026-07-06T23:22:47.781940Z","submitted_at":"2026-05-11T17:41:16Z","title":"Effective, Efficient, and General Information Abstraction for Imperfect-Information Extensive-Form Games","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2605.10900","snapshot_observed_at":"2026-07-31T04:48:47.403159Z","title":"Effective, efficient, and general information abstraction for imperfect-information extensive-form games","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":87,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.403159Z"},"links":{"cited_paper":"/paper/2605.10900","citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:853a9fdc94212c5955742282eef18257bdb929d660b530905722bef621fb008e","observation_id":"6ef40b14-c882-43d0-8137-3e8f8aae0b9a","resolution":{"observed_at":"2026-07-31T04:48:47.403159Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2605.19928","last_updated":"2026-05-19T14:49:30Z","snapshot_observed_at":"2026-08-02T22:50:00.712639Z","submitted_at":"2026-05-19T14:49:30Z","title":"Real-Time Parallel Counterfactual Regret Minimization","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2605.19928","snapshot_observed_at":"2026-07-31T04:48:47.405402Z","title":"Real-time parallel counterfactual regret minimization","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":88,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.405402Z"},"links":{"cited_paper":"/paper/2605.19928","citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:c6d56eb4ce461107b80aed5b1bbe454b764e447e4b1789f2e8c062058ae82fdd","observation_id":"107caa0b-cfe2-40c3-9f78-325970cc94ec","resolution":{"observed_at":"2026-07-31T04:48:47.405402Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.407650Z","title":"Rl-cfr: improving action abstraction for imperfect information extensive-form games with reinforcement learning","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":89,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.407650Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:4b048a36240479ee41581ecae37fd32b662f3b1b141871129934a8225ddd46d8","observation_id":"0c987494-cf00-42b2-a323-9532e723384b","resolution":{"observed_at":"2026-07-31T04:48:47.407650Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2607.27035","last_updated":"2026-07-29T15:28:36Z","snapshot_observed_at":"2026-08-01T23:40:36.850318Z","submitted_at":"2026-07-29T15:28:36Z","title":"Correlated Chance Sampling for Monte Carlo Counterfactual Regret Minimization","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2607.27035","snapshot_observed_at":"2026-07-31T04:48:47.409963Z","title":"Correlated chance sampling for monte carlo counterfactual regret minimization, 2026 a","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":90,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.409963Z"},"links":{"cited_paper":"/paper/2607.27035","citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:859860b0fe1355bfe6259e09e2c2d622f2aae39c932980cee3942891544e08eb","observation_id":"e04489e2-452a-430d-8bd1-fe7c8c8668dd","resolution":{"observed_at":"2026-07-31T04:48:47.409963Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2605.30094","last_updated":"2026-05-28T15:38:33Z","snapshot_observed_at":"2026-07-06T23:39:24.682426Z","submitted_at":"2026-05-28T15:38:33Z","title":"PokerSkill: LLMs Can Play Expert-Level Poker without Training or Solvers","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2605.30094","snapshot_observed_at":"2026-07-31T04:48:47.412479Z","title":"Pokerskill: Llms can play expert-level poker without training or solvers","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":91,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.412479Z"},"links":{"cited_paper":"/paper/2605.30094","citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:c14db3c1aec97909de18cfa35824de1f2098d76c8c101f6a68ccd9163af98800","observation_id":"8d7483f5-77fc-4ff0-92d7-25710fcdf6ea","resolution":{"observed_at":"2026-07-31T04:48:47.412479Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.414777Z","title":"Safe opponent-exploitation subgame refinement","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":92,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.414777Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:a0c374a4b96e80880e875fbd61484da8198c2c0df411d1725bac23a5b5a414f9","observation_id":"a9e2647a-dc66-4c99-b9a6-0af7618f4836","resolution":{"observed_at":"2026-07-31T04:48:47.414777Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.417439Z","title":"Opponent-limited online search for imperfect information games","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":93,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.417439Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:67a55d61af6db234182ecd97030bb43c298ebbcf6fd4fc04c3d9ee28436e55d7","observation_id":"4a5b6f77-6b26-4d9c-aebf-053a2eaf1d0f","resolution":{"observed_at":"2026-07-31T04:48:47.417439Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.419742Z","title":"Safe strategies for agent modelling in games","venue":null,"work_id":null,"year":2004},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":94,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.419742Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:4536f6da2e5f8cbf0a0642453924cffe312b1f3ec4a8c6c7e52a3720d8d7acdc","observation_id":"810e6e45-1e4b-40e7-9c74-0fb40e08d761","resolution":{"observed_at":"2026-07-31T04:48:47.419742Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.422670Z","title":"Efficient last-iterate convergence in solving extensive-form games","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":95,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.422670Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:0d402e22f6f765cc73415997d6e10c343cb9d12e3bc8bac804674d4050e6d35e","observation_id":"d8ac54b5-8b81-4fcb-9c7a-e2e8036ff334","resolution":{"observed_at":"2026-07-31T04:48:47.422670Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.425081Z","title":"Faster game solving via asymmetry of step sizes","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":96,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.425081Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:0f28422f4e113d6884ff3fb949d7a7365b0a11e772f803c54c56e774738b289d","observation_id":"d828de44-0604-4d31-94cc-ed32bc6c1dd5","resolution":{"observed_at":"2026-07-31T04:48:47.425081Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.427532Z","title":"Adapting beyond the depth limit: Counter strategies in large imperfect information games","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":97,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.427532Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:b4b554b1b9396dcc2c90f12520599d4b85a3bc26498333fe9215f93680a2f895","observation_id":"f512d5a2-7c16-43d1-b2f0-dd27ea2f21a1","resolution":{"observed_at":"2026-07-31T04:48:47.427532Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.429806Z","title":"Deepstack: Expert-level artificial intelligence in heads-up no-limit poker","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":98,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.429806Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:b16ff0a8aaa1b8f2a4f591f94c06d44af241a41e0edfb7a4d9c6c166f8a28d70","observation_id":"b7f0a2bb-3fe6-42e5-8735-4d6273787a85","resolution":{"observed_at":"2026-07-31T04:48:47.429806Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2605.09150","last_updated":"2026-05-09T20:26:35Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2026-05-09T20:26:35Z","title":"AlphaExploitem: Going Beyond the Nash Equilibrium in Poker by Learning to Exploit Suboptimal Play","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2605.09150","snapshot_observed_at":"2026-07-31T04:48:47.432065Z","title":"Alphaexploitem: Going beyond the nash equilibrium in poker by learning to exploit suboptimal play","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":99,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.432065Z"},"links":{"cited_paper":"/paper/2605.09150","citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:532ed4445069b1c5a7b06b161220dcdb195f057acc0ce933c3d3be43a885d3a0","observation_id":"edab7af6-664e-45a1-9e30-6ea0d1086392","resolution":{"observed_at":"2026-07-31T04:48:47.432065Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.434218Z","title":"A survey of opponent modeling in adversarial domains","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":100,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.434218Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:adf6fb2397357be45f3cc43fbaeae265b3e8a40ac2f5c0a733ceb66cc0b2e147","observation_id":"09c48c32-efe3-4b33-b1e2-617230f6affc","resolution":{"observed_at":"2026-07-31T04:48:47.434218Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.439864Z","title":"Mcrnr: fast computing of restricted nash responses by means of sampling","venue":null,"work_id":null,"year":2010},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":101,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.439864Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:1f793b333921d29eaadd7d778db0c35b59764a3090c41d8bccdc868cfea076a7","observation_id":"7a046576-bd18-4296-8cf7-25024d858675","resolution":{"observed_at":"2026-07-31T04:48:47.439864Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.442292Z","title":"Bayes' bluff: opponent modelling in poker","venue":null,"work_id":null,"year":2005},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":102,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.442292Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:bf4843aa576f12945b4923757142ec94e22839dd736436ce1b48d12372b60a69","observation_id":"0d18b556-dd6b-4e35-af1d-c69fbfce0f8b","resolution":{"observed_at":"2026-07-31T04:48:47.442292Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.444636Z","title":"Learning not to regret","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":103,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.444636Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:0bff0f134e1d0b96dab9e533da285486b99648dd607e5a807a761e82f4c9f909","observation_id":"94963f48-b4a2-4389-9d1a-8a0bb24c26c4","resolution":{"observed_at":"2026-07-31T04:48:47.444636Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1407.5042","last_updated":"2014-07-18T15:41:28Z","snapshot_observed_at":"2026-07-06T03:49:27.130621Z","submitted_at":"2014-07-18T15:41:28Z","title":"Solving Large Imperfect Information Games Using CFR+","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1407.5042","snapshot_observed_at":"2026-07-31T04:48:47.446905Z","title":"Solving large imperfect information games using cfr+","venue":null,"work_id":null,"year":2014},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":104,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.446905Z"},"links":{"cited_paper":"/paper/1407.5042","citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:bcb4ae32ec756a5c95dcb3a0f0f0c19322e009cbba4da1de67e18158692482e7","observation_id":"9a177a92-e23d-4ab5-8f17-720b4d78d21e","resolution":{"observed_at":"2026-07-31T04:48:47.446905Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.449196Z","title":"Horse-cfr: Hierarchical opponent reasoning for safe exploitation counterfactual regret minimization","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":105,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.449196Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:a12658d21001f2e6590d89aa418253937f4f7e85b87d2895c084a6de9b36ea7c","observation_id":"2aa92267-0347-4619-afe2-1fcd46db5804","resolution":{"observed_at":"2026-07-31T04:48:47.449196Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.451676Z","title":"Dynamic discounted counterfactual regret minimization","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":106,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.451676Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:c575e00d7f634e9654a245c8012e149b40e527e8b85eac4ac60a6110f22d052a","observation_id":"2f52a83a-c96d-400c-9519-9c37b0cf30ac","resolution":{"observed_at":"2026-07-31T04:48:47.451676Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.453889Z","title":"Deep (predictive) discounted counterfactual regret minimization","venue":null,"work_id":null,"year":2026},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":107,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.453889Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:644ffd922e8a25915209e7f4790acac9b6d454f261f36bbbd9a24eb78720a138","observation_id":"13962ccc-08b9-4604-b269-c4817a0a3490","resolution":{"observed_at":"2026-07-31T04:48:47.453889Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-31T04:48:47.456192Z","title":"Regret minimization in games with incomplete information","venue":null,"work_id":null,"year":2007},"citing_paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation","version":1},"reference_index":108,"source":"arxiv_source","source_observed_at":"2026-07-31T04:48:47.456192Z"},"links":{"citing_paper":"/paper/2607.28520"},"observation_digest":"sha256:53282ce6a766c6b541348e7a946c56562f8ae9cfa7d8fd4cb0308006a7380c56","observation_id":"3e8b75cc-0ebe-43d7-912f-2e57e9028432","resolution":{"observed_at":"2026-07-31T04:48:47.456192Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2607.28520","last_updated":"2026-07-30T16:57:57Z","latest_version":1,"primary_category":"cs.GT","snapshot_observed_at":"2026-08-05T08:06:59.993355Z","submitted_at":"2026-07-30T16:57:57Z","title":"Agents That Certify Their Own Exploits: Confidence-Scheduled Restricted Responses for Safe Opponent Exploitation"},"reference_resolution":{"displayed":100,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":100,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":100},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-07T06:34:17.273281+00:00","source":"crossref"},{"observed_at":"2026-08-07T06:34:11.927384+00:00","source":"retraction_watch"}],"thesis":"As of 7 August 2026, this Paper Citation Record lists 100 of 100 outbound references and 0 inbound Pith citation observations for arXiv:2607.28520."}