{"bundle_type":"pith_open_graph_bundle","bundle_version":"1.0","pith_number":"pith:2022:J6HRIENBQYYSE3D26N5TUJPBLT","short_pith_number":"pith:J6HRIENB","canonical_record":{"source":{"id":"2203.12759","kind":"arxiv","version":3},"metadata":{"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.RO","submitted_at":"2022-03-23T23:05:28Z","cross_cats_sorted":["cs.AI"],"title_canon_sha256":"04947d4da71e352b7bb2ec550f74e9d03d30ad2cb27e703f2347242006a28d6f","abstract_canon_sha256":"8ef94a01ba1b08b554e4ff3583bcb394fde1a1ebd60e15f8f38f3bca55dd972f"},"schema_version":"1.0"},"canonical_sha256":"4f8f1411a18631226c7af37b3a25e15ce99b17f1bdb9efb6badbae427066242b","source":{"kind":"arxiv","id":"2203.12759","version":3},"source_aliases":[{"alias_kind":"arxiv","alias_value":"2203.12759","created_at":"2026-07-05T04:10:26Z"},{"alias_kind":"arxiv_version","alias_value":"2203.12759v3","created_at":"2026-07-05T04:10:26Z"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2203.12759","created_at":"2026-07-05T04:10:26Z"},{"alias_kind":"pith_short_12","alias_value":"J6HRIENBQYYS","created_at":"2026-07-05T04:10:26Z"},{"alias_kind":"pith_short_16","alias_value":"J6HRIENBQYYSE3D2","created_at":"2026-07-05T04:10:26Z"},{"alias_kind":"pith_short_8","alias_value":"J6HRIENB","created_at":"2026-07-05T04:10:26Z"}],"events":[{"event_type":"record_created","subject_pith_number":"pith:2022:J6HRIENBQYYSE3D26N5TUJPBLT","target":"record","payload":{"canonical_record":{"source":{"id":"2203.12759","kind":"arxiv","version":3},"metadata":{"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.RO","submitted_at":"2022-03-23T23:05:28Z","cross_cats_sorted":["cs.AI"],"title_canon_sha256":"04947d4da71e352b7bb2ec550f74e9d03d30ad2cb27e703f2347242006a28d6f","abstract_canon_sha256":"8ef94a01ba1b08b554e4ff3583bcb394fde1a1ebd60e15f8f38f3bca55dd972f"},"schema_version":"1.0"},"canonical_sha256":"4f8f1411a18631226c7af37b3a25e15ce99b17f1bdb9efb6badbae427066242b","receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-07-05T04:10:26.781640Z","signature_b64":"kbMqRQBc6J8FlaGENsMW3EJAoWGzpTdSAry+bggTGUZhP1hTiIrbpgMRFpJc5KBQj97kIOeRpZNphJkY5VHCCQ==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"4f8f1411a18631226c7af37b3a25e15ce99b17f1bdb9efb6badbae427066242b","last_reissued_at":"2026-07-05T04:10:26.781077Z","signature_status":"signed_v1","first_computed_at":"2026-07-05T04:10:26.781077Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"source_kind":"arxiv","source_id":"2203.12759","source_version":3,"attestation_state":"computed"},"signer":{"signer_id":"pith.science","signer_type":"pith_registry","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"created_at":"2026-07-05T04:10:26Z","supersedes":[],"prev_event":null,"signature":{"signature_status":"signed_v1","algorithm":"ed25519","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signature_b64":"vZaRDyQ8wspBrWPO2Rl0Z7VaMqBG6mzDyRIQHC9qghzf/rgTnV+XhB7qdS88iRhdCmpqFrCqHrXhmHs5jfekDg==","signed_message":"open_graph_event_sha256_bytes","signed_at":"2026-08-09T18:53:20.910827Z"},"content_sha256":"2493de4694c4771b604e0ee7bc901f038e68db2bdfe23b9969cbe77db6c18686","schema_version":"1.0","event_id":"sha256:2493de4694c4771b604e0ee7bc901f038e68db2bdfe23b9969cbe77db6c18686"},{"event_type":"graph_snapshot","subject_pith_number":"pith:2022:J6HRIENBQYYSE3D26N5TUJPBLT","target":"graph","payload":{"graph_snapshot":{"paper":{"title":"Asynchronous Reinforcement Learning for Real-Time Control of Physical Robots","license":"http://creativecommons.org/licenses/by/4.0/","headline":"","cross_cats":["cs.AI"],"primary_cat":"cs.RO","authors_text":"A. Rupam Mahmood, Yufeng Yuan","submitted_at":"2022-03-23T23:05:28Z","abstract_excerpt":"An oft-ignored challenge of real-world reinforcement learning is that the real world does not pause when agents make learning updates. As standard simulated environments do not address this real-time aspect of learning, most available implementations of RL algorithms process environment interactions and learning updates sequentially. As a consequence, when such implementations are deployed in the real world, they may make decisions based on significantly delayed observations and not act responsively. Asynchronous learning has been proposed to solve this issue, but no systematic comparison betw"},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2203.12759","kind":"arxiv","version":3},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2203.12759/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"verdict_id":null},"signer":{"signer_id":"pith.science","signer_type":"pith_registry","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"created_at":"2026-07-05T04:10:26Z","supersedes":[],"prev_event":null,"signature":{"signature_status":"signed_v1","algorithm":"ed25519","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signature_b64":"cjYHZqLBttCkAjjOgzBoE+4cHePQCLpjjUnbyeinh3NU8aGV22/JN3XN11jNZh+xB1JBoEsu/m4sTRd8jRcIDA==","signed_message":"open_graph_event_sha256_bytes","signed_at":"2026-08-09T18:53:20.911367Z"},"content_sha256":"823b44598d47622e2aad7029c156d4527e242cf9d23d600ad2eaf28f99b7fd6c","schema_version":"1.0","event_id":"sha256:823b44598d47622e2aad7029c156d4527e242cf9d23d600ad2eaf28f99b7fd6c"}],"timestamp_proofs":[],"mirror_hints":[{"mirror_type":"https","name":"Pith Resolver","base_url":"https://pith.science","bundle_url":"https://pith.science/pith/J6HRIENBQYYSE3D26N5TUJPBLT/bundle.json","state_url":"https://pith.science/pith/J6HRIENBQYYSE3D26N5TUJPBLT/state.json","well_known_bundle_url":"https://pith.science/.well-known/pith/J6HRIENBQYYSE3D26N5TUJPBLT/bundle.json","status":"primary"}],"public_keys":[{"key_id":"pith-v1-2026-05","algorithm":"ed25519","format":"raw","public_key_b64":"stVStoiQhXFxp4s2pdzPNoqVNBMojDU/fJ2db5S3CbM=","public_key_hex":"b2d552b68890857171a78b36a5dccf368a953413288c353f7c9d9d6f94b709b3","fingerprint_sha256_b32_first128bits":"RVFV5Z2OI2J3ZUO7ERDEBCYNKS","fingerprint_sha256_hex":"8d4b5ee74e4693bcd1df2446408b0d54","rotates_at":null,"url":"https://pith.science/pith-signing-key.json","notes":"Pith uses this Ed25519 key to sign canonical record SHA-256 digests. Verify with: ed25519_verify(public_key, message=canonical_sha256_bytes, signature=base64decode(signature_b64))."}],"merge_version":"pith-open-graph-merge-v1","built_at":"2026-08-09T18:53:20Z","links":{"resolver":"https://pith.science/pith/J6HRIENBQYYSE3D26N5TUJPBLT","bundle":"https://pith.science/pith/J6HRIENBQYYSE3D26N5TUJPBLT/bundle.json","state":"https://pith.science/pith/J6HRIENBQYYSE3D26N5TUJPBLT/state.json","well_known_bundle":"https://pith.science/.well-known/pith/J6HRIENBQYYSE3D26N5TUJPBLT/bundle.json"},"state":{"state_type":"pith_open_graph_state","state_version":"1.0","pith_number":"pith:2022:J6HRIENBQYYSE3D26N5TUJPBLT","merge_version":"pith-open-graph-merge-v1","event_count":2,"valid_event_count":2,"invalid_event_count":0,"equivocation_count":0,"current":{"canonical_record":{"metadata":{"abstract_canon_sha256":"8ef94a01ba1b08b554e4ff3583bcb394fde1a1ebd60e15f8f38f3bca55dd972f","cross_cats_sorted":["cs.AI"],"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.RO","submitted_at":"2022-03-23T23:05:28Z","title_canon_sha256":"04947d4da71e352b7bb2ec550f74e9d03d30ad2cb27e703f2347242006a28d6f"},"schema_version":"1.0","source":{"id":"2203.12759","kind":"arxiv","version":3}},"source_aliases":[{"alias_kind":"arxiv","alias_value":"2203.12759","created_at":"2026-07-05T04:10:26Z"},{"alias_kind":"arxiv_version","alias_value":"2203.12759v3","created_at":"2026-07-05T04:10:26Z"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2203.12759","created_at":"2026-07-05T04:10:26Z"},{"alias_kind":"pith_short_12","alias_value":"J6HRIENBQYYS","created_at":"2026-07-05T04:10:26Z"},{"alias_kind":"pith_short_16","alias_value":"J6HRIENBQYYSE3D2","created_at":"2026-07-05T04:10:26Z"},{"alias_kind":"pith_short_8","alias_value":"J6HRIENB","created_at":"2026-07-05T04:10:26Z"}],"graph_snapshots":[{"event_id":"sha256:823b44598d47622e2aad7029c156d4527e242cf9d23d600ad2eaf28f99b7fd6c","target":"graph","created_at":"2026-07-05T04:10:26Z","signer":{"key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signer_id":"pith.science","signer_type":"pith_registry"},"payload":{"graph_snapshot":{"author_claims":{"count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","strong_count":0},"builder_version":"pith-number-builder-2026-05-17-v1","claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"integrity":{"available":true,"clean":true,"detectors_run":[],"endpoint":"/pith/2203.12759/integrity.json","findings":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938","summary":{"advisory":0,"by_detector":{},"critical":0,"informational":0}},"paper":{"abstract_excerpt":"An oft-ignored challenge of real-world reinforcement learning is that the real world does not pause when agents make learning updates. As standard simulated environments do not address this real-time aspect of learning, most available implementations of RL algorithms process environment interactions and learning updates sequentially. As a consequence, when such implementations are deployed in the real world, they may make decisions based on significantly delayed observations and not act responsively. Asynchronous learning has been proposed to solve this issue, but no systematic comparison betw","authors_text":"A. Rupam Mahmood, Yufeng Yuan","cross_cats":["cs.AI"],"headline":"","license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.RO","submitted_at":"2022-03-23T23:05:28Z","title":"Asynchronous Reinforcement Learning for Real-Time Control of Physical Robots"},"references":{"count":0,"internal_anchors":0,"resolved_work":0,"sample":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2203.12759","kind":"arxiv","version":3},"verdict":{"created_at":null,"id":null,"model_set":{},"one_line_summary":"","pipeline_version":null,"pith_extraction_headline":"","strongest_claim":"","weakest_assumption":""}},"verdict_id":null}}],"author_attestations":[],"timestamp_anchors":[],"storage_attestations":[],"citation_signatures":[],"replication_records":[],"corrections":[],"mirror_hints":[],"record_created":{"event_id":"sha256:2493de4694c4771b604e0ee7bc901f038e68db2bdfe23b9969cbe77db6c18686","target":"record","created_at":"2026-07-05T04:10:26Z","signer":{"key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signer_id":"pith.science","signer_type":"pith_registry"},"payload":{"attestation_state":"computed","canonical_record":{"metadata":{"abstract_canon_sha256":"8ef94a01ba1b08b554e4ff3583bcb394fde1a1ebd60e15f8f38f3bca55dd972f","cross_cats_sorted":["cs.AI"],"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.RO","submitted_at":"2022-03-23T23:05:28Z","title_canon_sha256":"04947d4da71e352b7bb2ec550f74e9d03d30ad2cb27e703f2347242006a28d6f"},"schema_version":"1.0","source":{"id":"2203.12759","kind":"arxiv","version":3}},"canonical_sha256":"4f8f1411a18631226c7af37b3a25e15ce99b17f1bdb9efb6badbae427066242b","receipt":{"algorithm":"ed25519","builder_version":"pith-number-builder-2026-05-17-v1","canonical_sha256":"4f8f1411a18631226c7af37b3a25e15ce99b17f1bdb9efb6badbae427066242b","first_computed_at":"2026-07-05T04:10:26.781077Z","key_id":"pith-v1-2026-05","kind":"pith_receipt","last_reissued_at":"2026-07-05T04:10:26.781077Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","receipt_version":"0.3","signature_b64":"kbMqRQBc6J8FlaGENsMW3EJAoWGzpTdSAry+bggTGUZhP1hTiIrbpgMRFpJc5KBQj97kIOeRpZNphJkY5VHCCQ==","signature_status":"signed_v1","signed_at":"2026-07-05T04:10:26.781640Z","signed_message":"canonical_sha256_bytes"},"source_id":"2203.12759","source_kind":"arxiv","source_version":3}}},"equivocations":[],"invalid_events":[],"applied_event_ids":["sha256:2493de4694c4771b604e0ee7bc901f038e68db2bdfe23b9969cbe77db6c18686","sha256:823b44598d47622e2aad7029c156d4527e242cf9d23d600ad2eaf28f99b7fd6c"],"state_sha256":"83183a7c197324c77c272b6586c4517ccc1a37092b83d733fc57efe083f2126c"},"bundle_signature":{"signature_status":"signed_v1","algorithm":"ed25519","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signature_b64":"67wkFmUynJoDkbM5/wlxkeRsOpOnODRjdmbWLLPM8tCzkOmtnmqowy6j90Lp8nmha3ROwRInL8I65ckP0Rk4Bg==","signed_message":"bundle_sha256_bytes","signed_at":"2026-08-09T18:53:20.917176Z","bundle_sha256":"98986c855dcf4a3102219494d4aa4c4b5f4e20c53e4aeb38418c77b8640faff0"}}