{"record_type":"pith_number_record","schema_url":"https://pith.science/schemas/pith-number/v1.json","pith_number":"pith:2022:2UZZYFZZ4SZ62QVHAVFIDHUE4I","short_pith_number":"pith:2UZZYFZZ","schema_version":"1.0","canonical_sha256":"d5339c1739e4b3ed42a7054a819e84e2151f427cf04320854cd3b7078586c898","source":{"kind":"arxiv","id":"2202.03609","version":5},"attestation_state":"computed","paper":{"title":"PolicyCleanse: Backdoor Detection and Mitigation in Reinforcement Learning","license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","headline":"","cross_cats":["cs.AI"],"primary_cat":"cs.LG","authors_text":"Ang Li, Cong Liu, Junfeng Guo","submitted_at":"2022-02-08T02:49:09Z","abstract_excerpt":"While real-world applications of reinforcement learning are becoming popular, the security and robustness of RL systems are worthy of more attention and exploration. In particular, recent works have revealed that, in a multi-agent RL environment, backdoor trigger actions can be injected into a victim agent (a.k.a. Trojan agent), which can result in a catastrophic failure as soon as it sees the backdoor trigger action. To ensure the security of RL agents against malicious backdoors, in this work, we propose the problem of Backdoor Detection in a multi-agent competitive reinforcement learning sy"},"verification_status":{"content_addressed":true,"pith_receipt":true,"author_attested":false,"weak_author_claims":0,"strong_author_claims":0,"externally_anchored":false,"storage_verified":false,"citation_signatures":0,"replication_records":0,"graph_snapshot":true,"references_resolved":false,"formal_links_present":false},"canonical_record":{"source":{"id":"2202.03609","kind":"arxiv","version":5},"metadata":{"license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","primary_cat":"cs.LG","submitted_at":"2022-02-08T02:49:09Z","cross_cats_sorted":["cs.AI"],"title_canon_sha256":"65b6b19ae180e05a6aad5bb77cbd387f459c61fdf9eae8d4cd6f7eed0d15e8aa","abstract_canon_sha256":"31e9d4611f9a98894882ec7ec781b7c037fffb7b79c1310f49d0cde4f7e078b7"},"schema_version":"1.0"},"receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-07-05T06:50:36.645895Z","signature_b64":"Y+HxUeKM7lGGabWVrksEJIhodAFdIhMMEgMUtImduClrJC8K2LJglOmX8fOxp0axuaLFF3XEUs7r0AzWnnSABA==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"d5339c1739e4b3ed42a7054a819e84e2151f427cf04320854cd3b7078586c898","last_reissued_at":"2026-07-05T06:50:36.645363Z","signature_status":"signed_v1","first_computed_at":"2026-07-05T06:50:36.645363Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"graph_snapshot":{"paper":{"title":"PolicyCleanse: Backdoor Detection and Mitigation in Reinforcement Learning","license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","headline":"","cross_cats":["cs.AI"],"primary_cat":"cs.LG","authors_text":"Ang Li, Cong Liu, Junfeng Guo","submitted_at":"2022-02-08T02:49:09Z","abstract_excerpt":"While real-world applications of reinforcement learning are becoming popular, the security and robustness of RL systems are worthy of more attention and exploration. In particular, recent works have revealed that, in a multi-agent RL environment, backdoor trigger actions can be injected into a victim agent (a.k.a. Trojan agent), which can result in a catastrophic failure as soon as it sees the backdoor trigger action. To ensure the security of RL agents against malicious backdoors, in this work, we propose the problem of Backdoor Detection in a multi-agent competitive reinforcement learning sy"},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2202.03609","kind":"arxiv","version":5},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2202.03609/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"aliases":[{"alias_kind":"arxiv","alias_value":"2202.03609","created_at":"2026-07-05T06:50:36.645423+00:00"},{"alias_kind":"arxiv_version","alias_value":"2202.03609v5","created_at":"2026-07-05T06:50:36.645423+00:00"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2202.03609","created_at":"2026-07-05T06:50:36.645423+00:00"},{"alias_kind":"pith_short_12","alias_value":"2UZZYFZZ4SZ6","created_at":"2026-07-05T06:50:36.645423+00:00"},{"alias_kind":"pith_short_16","alias_value":"2UZZYFZZ4SZ62QVH","created_at":"2026-07-05T06:50:36.645423+00:00"},{"alias_kind":"pith_short_8","alias_value":"2UZZYFZZ","created_at":"2026-07-05T06:50:36.645423+00:00"}],"events":[],"event_summary":{},"paper_claims":[],"inbound_citations":{"count":1,"internal_anchor_count":1,"sample":[{"citing_arxiv_id":"2505.19532","citing_title":"Fox in the Henhouse: Supply-Chain Backdoor Attacks Against Reinforcement Learning","ref_index":23,"is_internal_anchor":true}]},"formal_canon":{"evidence_count":0,"sample":[],"anchors":[]},"links":{"html":"https://pith.science/pith/2UZZYFZZ4SZ62QVHAVFIDHUE4I","json":"https://pith.science/pith/2UZZYFZZ4SZ62QVHAVFIDHUE4I.json","graph_json":"https://pith.science/api/pith-number/2UZZYFZZ4SZ62QVHAVFIDHUE4I/graph.json","events_json":"https://pith.science/api/pith-number/2UZZYFZZ4SZ62QVHAVFIDHUE4I/events.json","paper":"https://pith.science/paper/2UZZYFZZ"},"agent_actions":{"view_html":"https://pith.science/pith/2UZZYFZZ4SZ62QVHAVFIDHUE4I","download_json":"https://pith.science/pith/2UZZYFZZ4SZ62QVHAVFIDHUE4I.json","view_paper":"https://pith.science/paper/2UZZYFZZ","resolve_alias":"https://pith.science/api/pith-number/resolve?arxiv=2202.03609&json=true","fetch_graph":"https://pith.science/api/pith-number/2UZZYFZZ4SZ62QVHAVFIDHUE4I/graph.json","fetch_events":"https://pith.science/api/pith-number/2UZZYFZZ4SZ62QVHAVFIDHUE4I/events.json","actions":{"anchor_timestamp":"https://pith.science/pith/2UZZYFZZ4SZ62QVHAVFIDHUE4I/action/timestamp_anchor","attest_storage":"https://pith.science/pith/2UZZYFZZ4SZ62QVHAVFIDHUE4I/action/storage_attestation","attest_author":"https://pith.science/pith/2UZZYFZZ4SZ62QVHAVFIDHUE4I/action/author_attestation","sign_citation":"https://pith.science/pith/2UZZYFZZ4SZ62QVHAVFIDHUE4I/action/citation_signature","submit_replication":"https://pith.science/pith/2UZZYFZZ4SZ62QVHAVFIDHUE4I/action/replication_record"}},"created_at":"2026-07-05T06:50:36.645423+00:00","updated_at":"2026-07-05T06:50:36.645423+00:00"}