{"state_type":"pith_open_graph_state","state_version":"1.0","pith_number":"pith:2019:O2L23Z2FNSFPFPM6Y57Z3N4H6C","merge_version":"pith-open-graph-merge-v1","event_count":2,"valid_event_count":2,"invalid_event_count":0,"equivocation_count":0,"current":{"canonical_record":{"metadata":{"abstract_canon_sha256":"01a9df855e25fba004e54cbf2b14a4c1f23decc67b4d82e7a5083f4166459413","cross_cats_sorted":["stat.ML"],"license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","primary_cat":"cs.LG","submitted_at":"2019-01-28T15:25:10Z","title_canon_sha256":"d5bbb2744a40fee10d36e917bf5f6991ef5edc1ea881c393c6d1283f3dd8cd3b"},"schema_version":"1.0","source":{"id":"1901.09732","kind":"arxiv","version":2}},"source_aliases":[{"alias_kind":"arxiv","alias_value":"1901.09732","created_at":"2026-05-17T23:55:16Z"},{"alias_kind":"arxiv_version","alias_value":"1901.09732v2","created_at":"2026-05-17T23:55:16Z"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.1901.09732","created_at":"2026-05-17T23:55:16Z"},{"alias_kind":"pith_short_12","alias_value":"O2L23Z2FNSFP","created_at":"2026-05-18T12:33:24Z"},{"alias_kind":"pith_short_16","alias_value":"O2L23Z2FNSFPFPM6","created_at":"2026-05-18T12:33:24Z"},{"alias_kind":"pith_short_8","alias_value":"O2L23Z2F","created_at":"2026-05-18T12:33:24Z"}],"graph_snapshots":[{"event_id":"sha256:9f9d5a89c03f29e8ecbe18c50f325cd767a8e701e428317e233cf01264f05f88","target":"graph","created_at":"2026-05-17T23:55:16Z","signer":{"key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signer_id":"pith.science","signer_type":"pith_registry"},"payload":{"graph_snapshot":{"author_claims":{"count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","strong_count":0},"builder_version":"pith-number-builder-2026-05-17-v1","claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"paper":{"abstract_excerpt":"Despite remarkable successes, Deep Reinforcement Learning (DRL) is not robust to hyperparameterization, implementation details, or small environment changes (Henderson et al. 2017, Zhang et al. 2018). Overcoming such sensitivity is key to making DRL applicable to real world problems. In this paper, we identify sensitivity to time discretization in near continuous-time environments as a critical factor; this covers, e.g., changing the number of frames per second, or the action frequency of the controller. Empirically, we find that Q-learning-based approaches such as Deep Q- learning (Mnih et al","authors_text":"Corentin Tallec, L\\'eonard Blier, Yann Ollivier","cross_cats":["stat.ML"],"headline":"","license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","primary_cat":"cs.LG","submitted_at":"2019-01-28T15:25:10Z","title":"Making Deep Q-learning methods robust to time discretization"},"references":{"count":0,"internal_anchors":0,"resolved_work":0,"sample":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"1901.09732","kind":"arxiv","version":2},"verdict":{"created_at":null,"id":null,"model_set":{},"one_line_summary":"","pipeline_version":null,"pith_extraction_headline":"","strongest_claim":"","weakest_assumption":""}},"verdict_id":null}}],"author_attestations":[],"timestamp_anchors":[],"storage_attestations":[],"citation_signatures":[],"replication_records":[],"corrections":[],"mirror_hints":[],"record_created":{"event_id":"sha256:bb8f0a9a84c48cb93d1a7f8a03bac8540a1afb4ebb508ecdaab1d09265d4a677","target":"record","created_at":"2026-05-17T23:55:16Z","signer":{"key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signer_id":"pith.science","signer_type":"pith_registry"},"payload":{"attestation_state":"computed","canonical_record":{"metadata":{"abstract_canon_sha256":"01a9df855e25fba004e54cbf2b14a4c1f23decc67b4d82e7a5083f4166459413","cross_cats_sorted":["stat.ML"],"license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","primary_cat":"cs.LG","submitted_at":"2019-01-28T15:25:10Z","title_canon_sha256":"d5bbb2744a40fee10d36e917bf5f6991ef5edc1ea881c393c6d1283f3dd8cd3b"},"schema_version":"1.0","source":{"id":"1901.09732","kind":"arxiv","version":2}},"canonical_sha256":"7697ade7456c8af2bd9ec77f9db787f0a0a55f61f8f07390cd9e4bc9b6ddd1f0","receipt":{"algorithm":"ed25519","builder_version":"pith-number-builder-2026-05-17-v1","canonical_sha256":"7697ade7456c8af2bd9ec77f9db787f0a0a55f61f8f07390cd9e4bc9b6ddd1f0","first_computed_at":"2026-05-17T23:55:16.635915Z","key_id":"pith-v1-2026-05","kind":"pith_receipt","last_reissued_at":"2026-05-17T23:55:16.635915Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","receipt_version":"0.3","signature_b64":"ZiJs0iIU7P2112x5WtcMjEF1UNxpb5WeCApbCOvYERklBnBuZmAxILlh4JA33qX+Pjyny2hYJ44QkHFpX1Z/Ag==","signature_status":"signed_v1","signed_at":"2026-05-17T23:55:16.636381Z","signed_message":"canonical_sha256_bytes"},"source_id":"1901.09732","source_kind":"arxiv","source_version":2}}},"equivocations":[],"invalid_events":[],"applied_event_ids":["sha256:bb8f0a9a84c48cb93d1a7f8a03bac8540a1afb4ebb508ecdaab1d09265d4a677","sha256:9f9d5a89c03f29e8ecbe18c50f325cd767a8e701e428317e233cf01264f05f88"],"state_sha256":"5bfa5b581bc2879de7dc62e800b63d199e72b2b0ab2b97e9bfc4d8d5a5f42ef0"}