{"state_type":"pith_open_graph_state","state_version":"1.0","pith_number":"pith:2021:Y5POH3PWYHZDIQUBU322ZF56RX","merge_version":"pith-open-graph-merge-v1","event_count":2,"valid_event_count":2,"invalid_event_count":0,"equivocation_count":0,"current":{"canonical_record":{"metadata":{"abstract_canon_sha256":"4782fa0b4ff7d664b122209fee880c4fbf362ff4512eaa11b296853d77177c49","cross_cats_sorted":["cs.AI","stat.ME","stat.ML"],"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.LG","submitted_at":"2021-08-30T14:23:48Z","title_canon_sha256":"6c829853fdb66010aa5af3f0979e859c178ef5ed6a9f2d7ad54348463f3e830e"},"schema_version":"1.0","source":{"id":"2108.13264","kind":"arxiv","version":4}},"source_aliases":[{"alias_kind":"arxiv","alias_value":"2108.13264","created_at":"2026-07-05T03:46:13Z"},{"alias_kind":"arxiv_version","alias_value":"2108.13264v4","created_at":"2026-07-05T03:46:13Z"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2108.13264","created_at":"2026-07-05T03:46:13Z"},{"alias_kind":"pith_short_12","alias_value":"Y5POH3PWYHZD","created_at":"2026-07-05T03:46:13Z"},{"alias_kind":"pith_short_16","alias_value":"Y5POH3PWYHZDIQUB","created_at":"2026-07-05T03:46:13Z"},{"alias_kind":"pith_short_8","alias_value":"Y5POH3PW","created_at":"2026-07-05T03:46:13Z"}],"graph_snapshots":[{"event_id":"sha256:31a41995e97e2bb522f7fd60b7fb9b30046dc9e86437312d39077259f716633f","target":"graph","created_at":"2026-07-05T03:46:13Z","signer":{"key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signer_id":"pith.science","signer_type":"pith_registry"},"payload":{"graph_snapshot":{"author_claims":{"count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","strong_count":0},"builder_version":"pith-number-builder-2026-05-17-v1","claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"integrity":{"available":true,"clean":true,"detectors_run":[],"endpoint":"/pith/2108.13264/integrity.json","findings":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938","summary":{"advisory":0,"by_detector":{},"critical":0,"informational":0}},"paper":{"abstract_excerpt":"Deep reinforcement learning (RL) algorithms are predominantly evaluated by comparing their relative performance on a large suite of tasks. Most published results on deep RL benchmarks compare point estimates of aggregate performance such as mean and median scores across tasks, ignoring the statistical uncertainty implied by the use of a finite number of training runs. Beginning with the Arcade Learning Environment (ALE), the shift towards computationally-demanding benchmarks has led to the practice of evaluating only a small number of runs per task, exacerbating the statistical uncertainty in ","authors_text":"Aaron Courville, Marc G. Bellemare, Max Schwarzer, Pablo Samuel Castro, Rishabh Agarwal","cross_cats":["cs.AI","stat.ME","stat.ML"],"headline":"","license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.LG","submitted_at":"2021-08-30T14:23:48Z","title":"Deep Reinforcement Learning at the Edge of the Statistical Precipice"},"references":{"count":0,"internal_anchors":0,"resolved_work":0,"sample":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2108.13264","kind":"arxiv","version":4},"verdict":{"created_at":null,"id":null,"model_set":{},"one_line_summary":"","pipeline_version":null,"pith_extraction_headline":"","strongest_claim":"","weakest_assumption":""}},"verdict_id":null}}],"author_attestations":[],"timestamp_anchors":[],"storage_attestations":[],"citation_signatures":[],"replication_records":[],"corrections":[],"mirror_hints":[],"record_created":{"event_id":"sha256:7457b263afa7d81e5956f86003f562caffc15a8528bfaf34c613c9bd2f02097a","target":"record","created_at":"2026-07-05T03:46:13Z","signer":{"key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signer_id":"pith.science","signer_type":"pith_registry"},"payload":{"attestation_state":"computed","canonical_record":{"metadata":{"abstract_canon_sha256":"4782fa0b4ff7d664b122209fee880c4fbf362ff4512eaa11b296853d77177c49","cross_cats_sorted":["cs.AI","stat.ME","stat.ML"],"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.LG","submitted_at":"2021-08-30T14:23:48Z","title_canon_sha256":"6c829853fdb66010aa5af3f0979e859c178ef5ed6a9f2d7ad54348463f3e830e"},"schema_version":"1.0","source":{"id":"2108.13264","kind":"arxiv","version":4}},"canonical_sha256":"c75ee3edf6c1f2344281a6f5ac97be8dfef244f3beae6961b037c2f487b496db","receipt":{"algorithm":"ed25519","builder_version":"pith-number-builder-2026-05-17-v1","canonical_sha256":"c75ee3edf6c1f2344281a6f5ac97be8dfef244f3beae6961b037c2f487b496db","first_computed_at":"2026-07-05T03:46:13.826451Z","key_id":"pith-v1-2026-05","kind":"pith_receipt","last_reissued_at":"2026-07-05T03:46:13.826451Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","receipt_version":"0.3","signature_b64":"UDtK+duKTvJk746gq/EfTOAn0al7wVzArmRkLbmoJ07pL/9vrvQAPBRZrTj/C4AbAPvFNgtcpVoCmITzDl68Dw==","signature_status":"signed_v1","signed_at":"2026-07-05T03:46:13.827012Z","signed_message":"canonical_sha256_bytes"},"source_id":"2108.13264","source_kind":"arxiv","source_version":4}}},"equivocations":[],"invalid_events":[],"applied_event_ids":["sha256:7457b263afa7d81e5956f86003f562caffc15a8528bfaf34c613c9bd2f02097a","sha256:31a41995e97e2bb522f7fd60b7fb9b30046dc9e86437312d39077259f716633f"],"state_sha256":"1fc7edc5bdc4e41fdd7d14383c17715e8818f14758550cb45b8134569c16aa47"}