{"record_type":"pith_number_record","schema_url":"https://pith.science/schemas/pith-number/v1.json","pith_number":"pith:2022:M7NUMIZNWXHOLHZRCM5WF6WTLB","short_pith_number":"pith:M7NUMIZN","schema_version":"1.0","canonical_sha256":"67db46232db5cee59f31133b62fad3586455ebbca8ab4e3dddcf13bb74d7ba0c","source":{"kind":"arxiv","id":"2202.11296","version":2},"attestation_state":"computed","paper":{"title":"Reinforcement Learning in Practice: Opportunities and Challenges","license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","headline":"","cross_cats":["cs.AI"],"primary_cat":"cs.LG","authors_text":"Yuxi Li","submitted_at":"2022-02-23T03:58:46Z","abstract_excerpt":"This article is a gentle discussion about the field of reinforcement learning in practice, about opportunities and challenges, touching a broad range of topics, with perspectives and without technical details. The article is based on both historical and recent research papers, surveys, tutorials, talks, blogs, books, (panel) discussions, and workshops/conferences. Various groups of readers, like researchers, engineers, students, managers, investors, officers, and people wanting to know more about the field, may find the article interesting.\n  In this article, we first give a brief introduction"},"verification_status":{"content_addressed":true,"pith_receipt":true,"author_attested":false,"weak_author_claims":0,"strong_author_claims":0,"externally_anchored":false,"storage_verified":false,"citation_signatures":0,"replication_records":0,"graph_snapshot":true,"references_resolved":false,"formal_links_present":false},"canonical_record":{"source":{"id":"2202.11296","kind":"arxiv","version":2},"metadata":{"license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","primary_cat":"cs.LG","submitted_at":"2022-02-23T03:58:46Z","cross_cats_sorted":["cs.AI"],"title_canon_sha256":"803319bfb9dde4c872547a7988975a7a255e88f3126d4a4b6d186ec58179145b","abstract_canon_sha256":"352ec4d3b8ce541c510173b27180cf1b7c0a6843917c043e956d3b9165c79537"},"schema_version":"1.0"},"receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-07-05T04:16:54.578899Z","signature_b64":"K6LoNvR1PmwVEs+ZukKADgy8rJL6rMLj373PUiWe4N0SMmPbIYGqCMyHwQctYWCczRyT6g1tiUFHcf2x30jRAw==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"67db46232db5cee59f31133b62fad3586455ebbca8ab4e3dddcf13bb74d7ba0c","last_reissued_at":"2026-07-05T04:16:54.578407Z","signature_status":"signed_v1","first_computed_at":"2026-07-05T04:16:54.578407Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"graph_snapshot":{"paper":{"title":"Reinforcement Learning in Practice: Opportunities and Challenges","license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","headline":"","cross_cats":["cs.AI"],"primary_cat":"cs.LG","authors_text":"Yuxi Li","submitted_at":"2022-02-23T03:58:46Z","abstract_excerpt":"This article is a gentle discussion about the field of reinforcement learning in practice, about opportunities and challenges, touching a broad range of topics, with perspectives and without technical details. The article is based on both historical and recent research papers, surveys, tutorials, talks, blogs, books, (panel) discussions, and workshops/conferences. Various groups of readers, like researchers, engineers, students, managers, investors, officers, and people wanting to know more about the field, may find the article interesting.\n  In this article, we first give a brief introduction"},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2202.11296","kind":"arxiv","version":2},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2202.11296/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"aliases":[{"alias_kind":"arxiv","alias_value":"2202.11296","created_at":"2026-07-05T04:16:54.578467+00:00"},{"alias_kind":"arxiv_version","alias_value":"2202.11296v2","created_at":"2026-07-05T04:16:54.578467+00:00"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2202.11296","created_at":"2026-07-05T04:16:54.578467+00:00"},{"alias_kind":"pith_short_12","alias_value":"M7NUMIZNWXHO","created_at":"2026-07-05T04:16:54.578467+00:00"},{"alias_kind":"pith_short_16","alias_value":"M7NUMIZNWXHOLHZR","created_at":"2026-07-05T04:16:54.578467+00:00"},{"alias_kind":"pith_short_8","alias_value":"M7NUMIZN","created_at":"2026-07-05T04:16:54.578467+00:00"}],"events":[],"event_summary":{},"paper_claims":[],"inbound_citations":{"count":2,"internal_anchor_count":0,"sample":[{"citing_arxiv_id":"2605.30719","citing_title":"When are LLMs Sufficient Policy Optimizers for Sequential RL Tasks?","ref_index":14,"is_internal_anchor":false},{"citing_arxiv_id":"2605.30719","citing_title":"When are LLMs Sufficient Policy Optimizers for Sequential RL Tasks?","ref_index":14,"is_internal_anchor":false}]},"formal_canon":{"evidence_count":0,"sample":[],"anchors":[]},"links":{"html":"https://pith.science/pith/M7NUMIZNWXHOLHZRCM5WF6WTLB","json":"https://pith.science/pith/M7NUMIZNWXHOLHZRCM5WF6WTLB.json","graph_json":"https://pith.science/api/pith-number/M7NUMIZNWXHOLHZRCM5WF6WTLB/graph.json","events_json":"https://pith.science/api/pith-number/M7NUMIZNWXHOLHZRCM5WF6WTLB/events.json","paper":"https://pith.science/paper/M7NUMIZN"},"agent_actions":{"view_html":"https://pith.science/pith/M7NUMIZNWXHOLHZRCM5WF6WTLB","download_json":"https://pith.science/pith/M7NUMIZNWXHOLHZRCM5WF6WTLB.json","view_paper":"https://pith.science/paper/M7NUMIZN","resolve_alias":"https://pith.science/api/pith-number/resolve?arxiv=2202.11296&json=true","fetch_graph":"https://pith.science/api/pith-number/M7NUMIZNWXHOLHZRCM5WF6WTLB/graph.json","fetch_events":"https://pith.science/api/pith-number/M7NUMIZNWXHOLHZRCM5WF6WTLB/events.json","actions":{"anchor_timestamp":"https://pith.science/pith/M7NUMIZNWXHOLHZRCM5WF6WTLB/action/timestamp_anchor","attest_storage":"https://pith.science/pith/M7NUMIZNWXHOLHZRCM5WF6WTLB/action/storage_attestation","attest_author":"https://pith.science/pith/M7NUMIZNWXHOLHZRCM5WF6WTLB/action/author_attestation","sign_citation":"https://pith.science/pith/M7NUMIZNWXHOLHZRCM5WF6WTLB/action/citation_signature","submit_replication":"https://pith.science/pith/M7NUMIZNWXHOLHZRCM5WF6WTLB/action/replication_record"}},"created_at":"2026-07-05T04:16:54.578467+00:00","updated_at":"2026-07-05T04:16:54.578467+00:00"}