{"bundle_type":"pith_open_graph_bundle","bundle_version":"1.0","pith_number":"pith:2025:I4YU5LSXFIBCZCP4BKITAVSHSY","short_pith_number":"pith:I4YU5LSX","canonical_record":{"source":{"id":"2501.03902","kind":"arxiv","version":1},"metadata":{"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.LG","submitted_at":"2025-01-07T16:10:09Z","cross_cats_sorted":["cs.AI"],"title_canon_sha256":"e277036bf882cde12e2961880954c05b9b3adb3d8660d91675bdc85f8b9e92d4","abstract_canon_sha256":"064b08e776667f1bcf7ea6de2f934d24b6f153973d6b768f9997da7ee3cd455b"},"schema_version":"1.0"},"canonical_sha256":"47314eae572a022c89fc0a913056479613fda881c3ff5d46ec9cba13303be5f9","source":{"kind":"arxiv","id":"2501.03902","version":1},"source_aliases":[{"alias_kind":"arxiv","alias_value":"2501.03902","created_at":"2026-07-05T09:58:03Z"},{"alias_kind":"arxiv_version","alias_value":"2501.03902v1","created_at":"2026-07-05T09:58:03Z"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2501.03902","created_at":"2026-07-05T09:58:03Z"},{"alias_kind":"pith_short_12","alias_value":"I4YU5LSXFIBC","created_at":"2026-07-05T09:58:03Z"},{"alias_kind":"pith_short_16","alias_value":"I4YU5LSXFIBCZCP4","created_at":"2026-07-05T09:58:03Z"},{"alias_kind":"pith_short_8","alias_value":"I4YU5LSX","created_at":"2026-07-05T09:58:03Z"}],"events":[{"event_type":"record_created","subject_pith_number":"pith:2025:I4YU5LSXFIBCZCP4BKITAVSHSY","target":"record","payload":{"canonical_record":{"source":{"id":"2501.03902","kind":"arxiv","version":1},"metadata":{"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.LG","submitted_at":"2025-01-07T16:10:09Z","cross_cats_sorted":["cs.AI"],"title_canon_sha256":"e277036bf882cde12e2961880954c05b9b3adb3d8660d91675bdc85f8b9e92d4","abstract_canon_sha256":"064b08e776667f1bcf7ea6de2f934d24b6f153973d6b768f9997da7ee3cd455b"},"schema_version":"1.0"},"canonical_sha256":"47314eae572a022c89fc0a913056479613fda881c3ff5d46ec9cba13303be5f9","receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-07-05T09:58:03.190429Z","signature_b64":"OlP4HrQxCl9L7scSyKNp7OPK8w9fDeNzQRxjvAKiWlVn4UMuO6FZRdWgp0baQssiGdtstNK6Ije/ofHjNk9RDg==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"47314eae572a022c89fc0a913056479613fda881c3ff5d46ec9cba13303be5f9","last_reissued_at":"2026-07-05T09:58:03.189943Z","signature_status":"signed_v1","first_computed_at":"2026-07-05T09:58:03.189943Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"source_kind":"arxiv","source_id":"2501.03902","source_version":1,"attestation_state":"computed"},"signer":{"signer_id":"pith.science","signer_type":"pith_registry","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"created_at":"2026-07-05T09:58:03Z","supersedes":[],"prev_event":null,"signature":{"signature_status":"signed_v1","algorithm":"ed25519","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signature_b64":"uql1tVk/68NKm3Hrz3zB+IoJvjGKz2xrcFNh8lEg/oEuEKq65a5U3Hid334VN4A4twTIlS+M6Z3iN6uZQEa1Dg==","signed_message":"open_graph_event_sha256_bytes","signed_at":"2026-08-08T08:46:59.593556Z"},"content_sha256":"679931caf67b8024322f46299cf6d14d3448c58d85e41ee45a63cb1d150ad51c","schema_version":"1.0","event_id":"sha256:679931caf67b8024322f46299cf6d14d3448c58d85e41ee45a63cb1d150ad51c"},{"event_type":"graph_snapshot","subject_pith_number":"pith:2025:I4YU5LSXFIBCZCP4BKITAVSHSY","target":"graph","payload":{"graph_snapshot":{"paper":{"title":"Explainable Reinforcement Learning via Temporal Policy Decomposition","license":"http://creativecommons.org/licenses/by/4.0/","headline":"","cross_cats":["cs.AI"],"primary_cat":"cs.LG","authors_text":"Alessio Russo, Franco Ruggeri, Karl Henrik Johansson, Rafia Inam","submitted_at":"2025-01-07T16:10:09Z","abstract_excerpt":"We investigate the explainability of Reinforcement Learning (RL) policies from a temporal perspective, focusing on the sequence of future outcomes associated with individual actions. In RL, value functions compress information about rewards collected across multiple trajectories and over an infinite horizon, allowing a compact form of knowledge representation. However, this compression obscures the temporal details inherent in sequential decision-making, presenting a key challenge for interpretability. We present Temporal Policy Decomposition (TPD), a novel explainability approach that explain"},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2501.03902","kind":"arxiv","version":1},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2501.03902/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"verdict_id":null},"signer":{"signer_id":"pith.science","signer_type":"pith_registry","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"created_at":"2026-07-05T09:58:03Z","supersedes":[],"prev_event":null,"signature":{"signature_status":"signed_v1","algorithm":"ed25519","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signature_b64":"Mt9hfgNb5BEh9sbYoT7Y9rkS6K0pq7pqf25YvHZ1ShmSMyMhIMvFgiJlVQOwLYCB+fu4NHmkeZLjqPoVRWs/BA==","signed_message":"open_graph_event_sha256_bytes","signed_at":"2026-08-08T08:46:59.594050Z"},"content_sha256":"41ed8d78e2d9a92c2a0022a137ada7c1b2db1b93a7b2158327e64d1970bd830b","schema_version":"1.0","event_id":"sha256:41ed8d78e2d9a92c2a0022a137ada7c1b2db1b93a7b2158327e64d1970bd830b"}],"timestamp_proofs":[],"mirror_hints":[{"mirror_type":"https","name":"Pith Resolver","base_url":"https://pith.science","bundle_url":"https://pith.science/pith/I4YU5LSXFIBCZCP4BKITAVSHSY/bundle.json","state_url":"https://pith.science/pith/I4YU5LSXFIBCZCP4BKITAVSHSY/state.json","well_known_bundle_url":"https://pith.science/.well-known/pith/I4YU5LSXFIBCZCP4BKITAVSHSY/bundle.json","status":"primary"}],"public_keys":[{"key_id":"pith-v1-2026-05","algorithm":"ed25519","format":"raw","public_key_b64":"stVStoiQhXFxp4s2pdzPNoqVNBMojDU/fJ2db5S3CbM=","public_key_hex":"b2d552b68890857171a78b36a5dccf368a953413288c353f7c9d9d6f94b709b3","fingerprint_sha256_b32_first128bits":"RVFV5Z2OI2J3ZUO7ERDEBCYNKS","fingerprint_sha256_hex":"8d4b5ee74e4693bcd1df2446408b0d54","rotates_at":null,"url":"https://pith.science/pith-signing-key.json","notes":"Pith uses this Ed25519 key to sign canonical record SHA-256 digests. Verify with: ed25519_verify(public_key, message=canonical_sha256_bytes, signature=base64decode(signature_b64))."}],"merge_version":"pith-open-graph-merge-v1","built_at":"2026-08-08T08:46:59Z","links":{"resolver":"https://pith.science/pith/I4YU5LSXFIBCZCP4BKITAVSHSY","bundle":"https://pith.science/pith/I4YU5LSXFIBCZCP4BKITAVSHSY/bundle.json","state":"https://pith.science/pith/I4YU5LSXFIBCZCP4BKITAVSHSY/state.json","well_known_bundle":"https://pith.science/.well-known/pith/I4YU5LSXFIBCZCP4BKITAVSHSY/bundle.json"},"state":{"state_type":"pith_open_graph_state","state_version":"1.0","pith_number":"pith:2025:I4YU5LSXFIBCZCP4BKITAVSHSY","merge_version":"pith-open-graph-merge-v1","event_count":2,"valid_event_count":2,"invalid_event_count":0,"equivocation_count":0,"current":{"canonical_record":{"metadata":{"abstract_canon_sha256":"064b08e776667f1bcf7ea6de2f934d24b6f153973d6b768f9997da7ee3cd455b","cross_cats_sorted":["cs.AI"],"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.LG","submitted_at":"2025-01-07T16:10:09Z","title_canon_sha256":"e277036bf882cde12e2961880954c05b9b3adb3d8660d91675bdc85f8b9e92d4"},"schema_version":"1.0","source":{"id":"2501.03902","kind":"arxiv","version":1}},"source_aliases":[{"alias_kind":"arxiv","alias_value":"2501.03902","created_at":"2026-07-05T09:58:03Z"},{"alias_kind":"arxiv_version","alias_value":"2501.03902v1","created_at":"2026-07-05T09:58:03Z"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2501.03902","created_at":"2026-07-05T09:58:03Z"},{"alias_kind":"pith_short_12","alias_value":"I4YU5LSXFIBC","created_at":"2026-07-05T09:58:03Z"},{"alias_kind":"pith_short_16","alias_value":"I4YU5LSXFIBCZCP4","created_at":"2026-07-05T09:58:03Z"},{"alias_kind":"pith_short_8","alias_value":"I4YU5LSX","created_at":"2026-07-05T09:58:03Z"}],"graph_snapshots":[{"event_id":"sha256:41ed8d78e2d9a92c2a0022a137ada7c1b2db1b93a7b2158327e64d1970bd830b","target":"graph","created_at":"2026-07-05T09:58:03Z","signer":{"key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signer_id":"pith.science","signer_type":"pith_registry"},"payload":{"graph_snapshot":{"author_claims":{"count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","strong_count":0},"builder_version":"pith-number-builder-2026-05-17-v1","claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"integrity":{"available":true,"clean":true,"detectors_run":[],"endpoint":"/pith/2501.03902/integrity.json","findings":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938","summary":{"advisory":0,"by_detector":{},"critical":0,"informational":0}},"paper":{"abstract_excerpt":"We investigate the explainability of Reinforcement Learning (RL) policies from a temporal perspective, focusing on the sequence of future outcomes associated with individual actions. In RL, value functions compress information about rewards collected across multiple trajectories and over an infinite horizon, allowing a compact form of knowledge representation. However, this compression obscures the temporal details inherent in sequential decision-making, presenting a key challenge for interpretability. We present Temporal Policy Decomposition (TPD), a novel explainability approach that explain","authors_text":"Alessio Russo, Franco Ruggeri, Karl Henrik Johansson, Rafia Inam","cross_cats":["cs.AI"],"headline":"","license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.LG","submitted_at":"2025-01-07T16:10:09Z","title":"Explainable Reinforcement Learning via Temporal Policy Decomposition"},"references":{"count":0,"internal_anchors":0,"resolved_work":0,"sample":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2501.03902","kind":"arxiv","version":1},"verdict":{"created_at":null,"id":null,"model_set":{},"one_line_summary":"","pipeline_version":null,"pith_extraction_headline":"","strongest_claim":"","weakest_assumption":""}},"verdict_id":null}}],"author_attestations":[],"timestamp_anchors":[],"storage_attestations":[],"citation_signatures":[],"replication_records":[],"corrections":[],"mirror_hints":[],"record_created":{"event_id":"sha256:679931caf67b8024322f46299cf6d14d3448c58d85e41ee45a63cb1d150ad51c","target":"record","created_at":"2026-07-05T09:58:03Z","signer":{"key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signer_id":"pith.science","signer_type":"pith_registry"},"payload":{"attestation_state":"computed","canonical_record":{"metadata":{"abstract_canon_sha256":"064b08e776667f1bcf7ea6de2f934d24b6f153973d6b768f9997da7ee3cd455b","cross_cats_sorted":["cs.AI"],"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.LG","submitted_at":"2025-01-07T16:10:09Z","title_canon_sha256":"e277036bf882cde12e2961880954c05b9b3adb3d8660d91675bdc85f8b9e92d4"},"schema_version":"1.0","source":{"id":"2501.03902","kind":"arxiv","version":1}},"canonical_sha256":"47314eae572a022c89fc0a913056479613fda881c3ff5d46ec9cba13303be5f9","receipt":{"algorithm":"ed25519","builder_version":"pith-number-builder-2026-05-17-v1","canonical_sha256":"47314eae572a022c89fc0a913056479613fda881c3ff5d46ec9cba13303be5f9","first_computed_at":"2026-07-05T09:58:03.189943Z","key_id":"pith-v1-2026-05","kind":"pith_receipt","last_reissued_at":"2026-07-05T09:58:03.189943Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","receipt_version":"0.3","signature_b64":"OlP4HrQxCl9L7scSyKNp7OPK8w9fDeNzQRxjvAKiWlVn4UMuO6FZRdWgp0baQssiGdtstNK6Ije/ofHjNk9RDg==","signature_status":"signed_v1","signed_at":"2026-07-05T09:58:03.190429Z","signed_message":"canonical_sha256_bytes"},"source_id":"2501.03902","source_kind":"arxiv","source_version":1}}},"equivocations":[],"invalid_events":[],"applied_event_ids":["sha256:679931caf67b8024322f46299cf6d14d3448c58d85e41ee45a63cb1d150ad51c","sha256:41ed8d78e2d9a92c2a0022a137ada7c1b2db1b93a7b2158327e64d1970bd830b"],"state_sha256":"133184c4ef7901eea72ee6a7360348ab614ef817f8fe5b83b00f6c8a44beed3a"},"bundle_signature":{"signature_status":"signed_v1","algorithm":"ed25519","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signature_b64":"IOcEBIT4mGVIOHlxNMHDAvxcIuiFR2JSRswF3Ij5Y7zXuqfhbmwGbJWkFLVNpTi/eLYXdJ4tEEqcUwiQdjhPAA==","signed_message":"bundle_sha256_bytes","signed_at":"2026-08-08T08:46:59.599054Z","bundle_sha256":"4de9c58abaaeb5b50e8c7ddc54c51dc25c61154702e351a473438127efc787eb"}}