{"bundle_type":"pith_open_graph_bundle","bundle_version":"1.0","pith_number":"pith:2022:KFGKEPJAT6W2R3GEX4U4RON745","short_pith_number":"pith:KFGKEPJA","canonical_record":{"source":{"id":"2210.04723","kind":"arxiv","version":5},"metadata":{"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.AI","submitted_at":"2022-10-10T14:27:53Z","cross_cats_sorted":["cs.HC"],"title_canon_sha256":"c62c0de35da88261a3c681311eb9d6fb65da4d70ce5fd86efae216bf1581aa72","abstract_canon_sha256":"ce481f17595f7afb5ccebebac09af0fcbcb40258573ccfd6b48c45c02e3126b1"},"schema_version":"1.0"},"canonical_sha256":"514ca23d209fada8ecc4bf29c8b9bfe7578ea9ba4e2ca46df064b55dd1036d69","source":{"kind":"arxiv","id":"2210.04723","version":5},"source_aliases":[{"alias_kind":"arxiv","alias_value":"2210.04723","created_at":"2026-07-05T10:49:02Z"},{"alias_kind":"arxiv_version","alias_value":"2210.04723v5","created_at":"2026-07-05T10:49:02Z"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2210.04723","created_at":"2026-07-05T10:49:02Z"},{"alias_kind":"pith_short_12","alias_value":"KFGKEPJAT6W2","created_at":"2026-07-05T10:49:02Z"},{"alias_kind":"pith_short_16","alias_value":"KFGKEPJAT6W2R3GE","created_at":"2026-07-05T10:49:02Z"},{"alias_kind":"pith_short_8","alias_value":"KFGKEPJA","created_at":"2026-07-05T10:49:02Z"}],"events":[{"event_type":"record_created","subject_pith_number":"pith:2022:KFGKEPJAT6W2R3GEX4U4RON745","target":"record","payload":{"canonical_record":{"source":{"id":"2210.04723","kind":"arxiv","version":5},"metadata":{"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.AI","submitted_at":"2022-10-10T14:27:53Z","cross_cats_sorted":["cs.HC"],"title_canon_sha256":"c62c0de35da88261a3c681311eb9d6fb65da4d70ce5fd86efae216bf1581aa72","abstract_canon_sha256":"ce481f17595f7afb5ccebebac09af0fcbcb40258573ccfd6b48c45c02e3126b1"},"schema_version":"1.0"},"canonical_sha256":"514ca23d209fada8ecc4bf29c8b9bfe7578ea9ba4e2ca46df064b55dd1036d69","receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-07-05T10:49:02.115731Z","signature_b64":"N7nS/K3fUL+2lr8Z8w5CSYZeVtToKZVx6Py3cYmztkCTbBZ2+k9GBvUw/Zhqs2X6Js/TdsRC3ON9ePpVIp75CQ==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"514ca23d209fada8ecc4bf29c8b9bfe7578ea9ba4e2ca46df064b55dd1036d69","last_reissued_at":"2026-07-05T10:49:02.115231Z","signature_status":"signed_v1","first_computed_at":"2026-07-05T10:49:02.115231Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"source_kind":"arxiv","source_id":"2210.04723","source_version":5,"attestation_state":"computed"},"signer":{"signer_id":"pith.science","signer_type":"pith_registry","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"created_at":"2026-07-05T10:49:02Z","supersedes":[],"prev_event":null,"signature":{"signature_status":"signed_v1","algorithm":"ed25519","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signature_b64":"AUACf86aLKIWIcMcYeI48fZ3zd7zzEes1THpdDoLKYMG4bdy+srpZSL2lCk+5SNwVZ3dWERYRwgx1E4Ax4mRDQ==","signed_message":"open_graph_event_sha256_bytes","signed_at":"2026-08-23T05:53:48.340798Z"},"content_sha256":"acf2be9c62992ab76faa960fed5a48998ed089b768d167573467c84f5d131a53","schema_version":"1.0","event_id":"sha256:acf2be9c62992ab76faa960fed5a48998ed089b768d167573467c84f5d131a53"},{"event_type":"graph_snapshot","subject_pith_number":"pith:2022:KFGKEPJAT6W2R3GEX4U4RON745","target":"graph","payload":{"graph_snapshot":{"paper":{"title":"Experiential Explanations for Reinforcement Learning","license":"http://creativecommons.org/licenses/by/4.0/","headline":"","cross_cats":["cs.HC"],"primary_cat":"cs.AI","authors_text":"Amal Alabdulkarim, Gennie Mansi, Kaely Hall, Madhuri Singh, Mark O. Riedl, Upol Ehsan","submitted_at":"2022-10-10T14:27:53Z","abstract_excerpt":"Reinforcement learning (RL) systems can be complex and non-interpretable, making it challenging for non-AI experts to understand or intervene in their decisions. This is due in part to the sequential nature of RL in which actions are chosen because of their likelihood of obtaining future rewards. However, RL agents discard the qualitative features of their training, making it difficult to recover user-understandable information for \"why\" an action is chosen. We propose a technique Experiential Explanations to generate counterfactual explanations by training influence predictors along with the "},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2210.04723","kind":"arxiv","version":5},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2210.04723/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"verdict_id":null},"signer":{"signer_id":"pith.science","signer_type":"pith_registry","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"created_at":"2026-07-05T10:49:02Z","supersedes":[],"prev_event":null,"signature":{"signature_status":"signed_v1","algorithm":"ed25519","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signature_b64":"Yu4v6M2R0S0FFXKoJQ2i93llA6l8nZLZjDYckoqf7NOwbiZhGQxZGXyGPsw98raMYF59xMx8/+XsCG/jEEIoBA==","signed_message":"open_graph_event_sha256_bytes","signed_at":"2026-08-23T05:53:48.341574Z"},"content_sha256":"eb98644f6ada9b0ef633de5d5a47a8381a7781fe388dd09ca74baf318a397a4f","schema_version":"1.0","event_id":"sha256:eb98644f6ada9b0ef633de5d5a47a8381a7781fe388dd09ca74baf318a397a4f"}],"timestamp_proofs":[],"mirror_hints":[{"mirror_type":"https","name":"Pith Resolver","base_url":"https://pith.science","bundle_url":"https://pith.science/pith/KFGKEPJAT6W2R3GEX4U4RON745/bundle.json","state_url":"https://pith.science/pith/KFGKEPJAT6W2R3GEX4U4RON745/state.json","well_known_bundle_url":"https://pith.science/.well-known/pith/KFGKEPJAT6W2R3GEX4U4RON745/bundle.json","status":"primary"}],"public_keys":[{"key_id":"pith-v1-2026-05","algorithm":"ed25519","format":"raw","public_key_b64":"stVStoiQhXFxp4s2pdzPNoqVNBMojDU/fJ2db5S3CbM=","public_key_hex":"b2d552b68890857171a78b36a5dccf368a953413288c353f7c9d9d6f94b709b3","fingerprint_sha256_b32_first128bits":"RVFV5Z2OI2J3ZUO7ERDEBCYNKS","fingerprint_sha256_hex":"8d4b5ee74e4693bcd1df2446408b0d54","rotates_at":null,"url":"https://pith.science/pith-signing-key.json","notes":"Pith uses this Ed25519 key to sign canonical record SHA-256 digests. Verify with: ed25519_verify(public_key, message=canonical_sha256_bytes, signature=base64decode(signature_b64))."}],"merge_version":"pith-open-graph-merge-v1","built_at":"2026-08-23T05:53:48Z","links":{"resolver":"https://pith.science/pith/KFGKEPJAT6W2R3GEX4U4RON745","bundle":"https://pith.science/pith/KFGKEPJAT6W2R3GEX4U4RON745/bundle.json","state":"https://pith.science/pith/KFGKEPJAT6W2R3GEX4U4RON745/state.json","well_known_bundle":"https://pith.science/.well-known/pith/KFGKEPJAT6W2R3GEX4U4RON745/bundle.json"},"state":{"state_type":"pith_open_graph_state","state_version":"1.0","pith_number":"pith:2022:KFGKEPJAT6W2R3GEX4U4RON745","merge_version":"pith-open-graph-merge-v1","event_count":2,"valid_event_count":2,"invalid_event_count":0,"equivocation_count":0,"current":{"canonical_record":{"metadata":{"abstract_canon_sha256":"ce481f17595f7afb5ccebebac09af0fcbcb40258573ccfd6b48c45c02e3126b1","cross_cats_sorted":["cs.HC"],"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.AI","submitted_at":"2022-10-10T14:27:53Z","title_canon_sha256":"c62c0de35da88261a3c681311eb9d6fb65da4d70ce5fd86efae216bf1581aa72"},"schema_version":"1.0","source":{"id":"2210.04723","kind":"arxiv","version":5}},"source_aliases":[{"alias_kind":"arxiv","alias_value":"2210.04723","created_at":"2026-07-05T10:49:02Z"},{"alias_kind":"arxiv_version","alias_value":"2210.04723v5","created_at":"2026-07-05T10:49:02Z"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2210.04723","created_at":"2026-07-05T10:49:02Z"},{"alias_kind":"pith_short_12","alias_value":"KFGKEPJAT6W2","created_at":"2026-07-05T10:49:02Z"},{"alias_kind":"pith_short_16","alias_value":"KFGKEPJAT6W2R3GE","created_at":"2026-07-05T10:49:02Z"},{"alias_kind":"pith_short_8","alias_value":"KFGKEPJA","created_at":"2026-07-05T10:49:02Z"}],"graph_snapshots":[{"event_id":"sha256:eb98644f6ada9b0ef633de5d5a47a8381a7781fe388dd09ca74baf318a397a4f","target":"graph","created_at":"2026-07-05T10:49:02Z","signer":{"key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signer_id":"pith.science","signer_type":"pith_registry"},"payload":{"graph_snapshot":{"author_claims":{"count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","strong_count":0},"builder_version":"pith-number-builder-2026-05-17-v1","claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"integrity":{"available":true,"clean":true,"detectors_run":[],"endpoint":"/pith/2210.04723/integrity.json","findings":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938","summary":{"advisory":0,"by_detector":{},"critical":0,"informational":0}},"paper":{"abstract_excerpt":"Reinforcement learning (RL) systems can be complex and non-interpretable, making it challenging for non-AI experts to understand or intervene in their decisions. This is due in part to the sequential nature of RL in which actions are chosen because of their likelihood of obtaining future rewards. However, RL agents discard the qualitative features of their training, making it difficult to recover user-understandable information for \"why\" an action is chosen. We propose a technique Experiential Explanations to generate counterfactual explanations by training influence predictors along with the ","authors_text":"Amal Alabdulkarim, Gennie Mansi, Kaely Hall, Madhuri Singh, Mark O. Riedl, Upol Ehsan","cross_cats":["cs.HC"],"headline":"","license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.AI","submitted_at":"2022-10-10T14:27:53Z","title":"Experiential Explanations for Reinforcement Learning"},"references":{"count":0,"internal_anchors":0,"resolved_work":0,"sample":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2210.04723","kind":"arxiv","version":5},"verdict":{"created_at":null,"id":null,"model_set":{},"one_line_summary":"","pipeline_version":null,"pith_extraction_headline":"","strongest_claim":"","weakest_assumption":""}},"verdict_id":null}}],"author_attestations":[],"timestamp_anchors":[],"storage_attestations":[],"citation_signatures":[],"replication_records":[],"corrections":[],"mirror_hints":[],"record_created":{"event_id":"sha256:acf2be9c62992ab76faa960fed5a48998ed089b768d167573467c84f5d131a53","target":"record","created_at":"2026-07-05T10:49:02Z","signer":{"key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signer_id":"pith.science","signer_type":"pith_registry"},"payload":{"attestation_state":"computed","canonical_record":{"metadata":{"abstract_canon_sha256":"ce481f17595f7afb5ccebebac09af0fcbcb40258573ccfd6b48c45c02e3126b1","cross_cats_sorted":["cs.HC"],"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.AI","submitted_at":"2022-10-10T14:27:53Z","title_canon_sha256":"c62c0de35da88261a3c681311eb9d6fb65da4d70ce5fd86efae216bf1581aa72"},"schema_version":"1.0","source":{"id":"2210.04723","kind":"arxiv","version":5}},"canonical_sha256":"514ca23d209fada8ecc4bf29c8b9bfe7578ea9ba4e2ca46df064b55dd1036d69","receipt":{"algorithm":"ed25519","builder_version":"pith-number-builder-2026-05-17-v1","canonical_sha256":"514ca23d209fada8ecc4bf29c8b9bfe7578ea9ba4e2ca46df064b55dd1036d69","first_computed_at":"2026-07-05T10:49:02.115231Z","key_id":"pith-v1-2026-05","kind":"pith_receipt","last_reissued_at":"2026-07-05T10:49:02.115231Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","receipt_version":"0.3","signature_b64":"N7nS/K3fUL+2lr8Z8w5CSYZeVtToKZVx6Py3cYmztkCTbBZ2+k9GBvUw/Zhqs2X6Js/TdsRC3ON9ePpVIp75CQ==","signature_status":"signed_v1","signed_at":"2026-07-05T10:49:02.115731Z","signed_message":"canonical_sha256_bytes"},"source_id":"2210.04723","source_kind":"arxiv","source_version":5}}},"equivocations":[],"invalid_events":[],"applied_event_ids":["sha256:acf2be9c62992ab76faa960fed5a48998ed089b768d167573467c84f5d131a53","sha256:eb98644f6ada9b0ef633de5d5a47a8381a7781fe388dd09ca74baf318a397a4f"],"state_sha256":"6e81abf42a6e9157e5a8880bdff786a07b90a735aee96f98469150b0300bbb8a"},"bundle_signature":{"signature_status":"signed_v1","algorithm":"ed25519","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signature_b64":"AKkYWs23K4vI0GV/iYv14nMN66qThEUlB1eC22rsyLt4rIDHP/gfvGiI55LxZ8ha/GWwGfDT82tlavDo+4KEAQ==","signed_message":"bundle_sha256_bytes","signed_at":"2026-08-23T05:53:48.346466Z","bundle_sha256":"fa1999d442023ff7c5fa0fd18116bbd78c1185fba838aed824026399c0de5843"}}