{"record_type":"pith_number_record","schema_url":"https://pith.science/schemas/pith-number/v1.json","pith_number":"pith:2025:XJGZMNSHNVUSG35PPUAO2JUYIU","short_pith_number":"pith:XJGZMNSH","schema_version":"1.0","canonical_sha256":"ba4d9636476d69236faf7d00ed2698453c75d34db23bcb60260bf292e6799e12","source":{"kind":"arxiv","id":"2505.10330","version":1},"attestation_state":"computed","paper":{"title":"Efficient Adaptation of Reinforcement Learning Agents to Sudden Environmental Change","license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","headline":"","cross_cats":["cs.AI"],"primary_cat":"cs.LG","authors_text":"Jonathan Clifford Balloch","submitted_at":"2025-05-15T14:19:01Z","abstract_excerpt":"Real-world autonomous decision-making systems, from robots to recommendation engines, must operate in environments that change over time. While deep reinforcement learning (RL) has shown an impressive ability to learn optimal policies in stationary environments, most methods are data intensive and assume a world that does not change between training and test time. As a result, conventional RL methods struggle to adapt when conditions change. This poses a fundamental challenge: how can RL agents efficiently adapt their behavior when encountering novel environmental changes during deployment wit"},"verification_status":{"content_addressed":true,"pith_receipt":true,"author_attested":false,"weak_author_claims":0,"strong_author_claims":0,"externally_anchored":false,"storage_verified":false,"citation_signatures":0,"replication_records":0,"graph_snapshot":true,"references_resolved":false,"formal_links_present":false},"canonical_record":{"source":{"id":"2505.10330","kind":"arxiv","version":1},"metadata":{"license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","primary_cat":"cs.LG","submitted_at":"2025-05-15T14:19:01Z","cross_cats_sorted":["cs.AI"],"title_canon_sha256":"428df0ff13960eac5b207a8a9962100427878aa585ef9b9686d1aa7797a4a2d0","abstract_canon_sha256":"2240c063a21887686c9d6e8306cb2d66a6bff833f8a601cc0da1bd8a0bfcd417"},"schema_version":"1.0"},"receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-07-05T11:03:37.022999Z","signature_b64":"BbV6Uo+MwOrY/VeE1fu3elr7SFmjy01DwR74+iFtHhgtHrHCGIqp/21BntoQnnsij5Yt9qUkyVj5mWYBWuwMDg==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"ba4d9636476d69236faf7d00ed2698453c75d34db23bcb60260bf292e6799e12","last_reissued_at":"2026-07-05T11:03:37.022537Z","signature_status":"signed_v1","first_computed_at":"2026-07-05T11:03:37.022537Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"graph_snapshot":{"paper":{"title":"Efficient Adaptation of Reinforcement Learning Agents to Sudden Environmental Change","license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","headline":"","cross_cats":["cs.AI"],"primary_cat":"cs.LG","authors_text":"Jonathan Clifford Balloch","submitted_at":"2025-05-15T14:19:01Z","abstract_excerpt":"Real-world autonomous decision-making systems, from robots to recommendation engines, must operate in environments that change over time. While deep reinforcement learning (RL) has shown an impressive ability to learn optimal policies in stationary environments, most methods are data intensive and assume a world that does not change between training and test time. As a result, conventional RL methods struggle to adapt when conditions change. This poses a fundamental challenge: how can RL agents efficiently adapt their behavior when encountering novel environmental changes during deployment wit"},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2505.10330","kind":"arxiv","version":1},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2505.10330/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"aliases":[{"alias_kind":"arxiv","alias_value":"2505.10330","created_at":"2026-07-05T11:03:37.022595+00:00"},{"alias_kind":"arxiv_version","alias_value":"2505.10330v1","created_at":"2026-07-05T11:03:37.022595+00:00"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2505.10330","created_at":"2026-07-05T11:03:37.022595+00:00"},{"alias_kind":"pith_short_12","alias_value":"XJGZMNSHNVUS","created_at":"2026-07-05T11:03:37.022595+00:00"},{"alias_kind":"pith_short_16","alias_value":"XJGZMNSHNVUSG35P","created_at":"2026-07-05T11:03:37.022595+00:00"},{"alias_kind":"pith_short_8","alias_value":"XJGZMNSH","created_at":"2026-07-05T11:03:37.022595+00:00"}],"events":[],"event_summary":{},"paper_claims":[],"inbound_citations":{"count":0,"internal_anchor_count":0,"sample":[]},"formal_canon":{"evidence_count":0,"sample":[],"anchors":[]},"links":{"html":"https://pith.science/pith/XJGZMNSHNVUSG35PPUAO2JUYIU","json":"https://pith.science/pith/XJGZMNSHNVUSG35PPUAO2JUYIU.json","graph_json":"https://pith.science/api/pith-number/XJGZMNSHNVUSG35PPUAO2JUYIU/graph.json","events_json":"https://pith.science/api/pith-number/XJGZMNSHNVUSG35PPUAO2JUYIU/events.json","paper":"https://pith.science/paper/XJGZMNSH"},"agent_actions":{"view_html":"https://pith.science/pith/XJGZMNSHNVUSG35PPUAO2JUYIU","download_json":"https://pith.science/pith/XJGZMNSHNVUSG35PPUAO2JUYIU.json","view_paper":"https://pith.science/paper/XJGZMNSH","resolve_alias":"https://pith.science/api/pith-number/resolve?arxiv=2505.10330&json=true","fetch_graph":"https://pith.science/api/pith-number/XJGZMNSHNVUSG35PPUAO2JUYIU/graph.json","fetch_events":"https://pith.science/api/pith-number/XJGZMNSHNVUSG35PPUAO2JUYIU/events.json","actions":{"anchor_timestamp":"https://pith.science/pith/XJGZMNSHNVUSG35PPUAO2JUYIU/action/timestamp_anchor","attest_storage":"https://pith.science/pith/XJGZMNSHNVUSG35PPUAO2JUYIU/action/storage_attestation","attest_author":"https://pith.science/pith/XJGZMNSHNVUSG35PPUAO2JUYIU/action/author_attestation","sign_citation":"https://pith.science/pith/XJGZMNSHNVUSG35PPUAO2JUYIU/action/citation_signature","submit_replication":"https://pith.science/pith/XJGZMNSHNVUSG35PPUAO2JUYIU/action/replication_record"}},"created_at":"2026-07-05T11:03:37.022595+00:00","updated_at":"2026-07-05T11:03:37.022595+00:00"}