{"record_type":"pith_number_record","schema_url":"https://pith.science/schemas/pith-number/v1.json","pith_number":"pith:2021:FPQAFWEN7L5IXRQ7U3BX4IGCUF","short_pith_number":"pith:FPQAFWEN","schema_version":"1.0","canonical_sha256":"2be002d88dfafa8bc61fa6c37e20c2a17cbf63d38b8d40b7c0148b011b5d6596","source":{"kind":"arxiv","id":"2102.11271","version":2},"attestation_state":"computed","paper":{"title":"Reinforcement Learning with Prototypical Representations","license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","headline":"","cross_cats":["cs.AI"],"primary_cat":"cs.LG","authors_text":"Alessandro Lazaric, Denis Yarats, Lerrel Pinto, Rob Fergus","submitted_at":"2021-02-22T18:56:34Z","abstract_excerpt":"Learning effective representations in image-based environments is crucial for sample efficient Reinforcement Learning (RL). Unfortunately, in RL, representation learning is confounded with the exploratory experience of the agent -- learning a useful representation requires diverse data, while effective exploration is only possible with coherent representations. Furthermore, we would like to learn representations that not only generalize across tasks but also accelerate downstream exploration for efficient task-specific training. To address these challenges we propose Proto-RL, a self-supervise"},"verification_status":{"content_addressed":true,"pith_receipt":true,"author_attested":false,"weak_author_claims":0,"strong_author_claims":0,"externally_anchored":false,"storage_verified":false,"citation_signatures":0,"replication_records":0,"graph_snapshot":true,"references_resolved":false,"formal_links_present":false},"canonical_record":{"source":{"id":"2102.11271","kind":"arxiv","version":2},"metadata":{"license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","primary_cat":"cs.LG","submitted_at":"2021-02-22T18:56:34Z","cross_cats_sorted":["cs.AI"],"title_canon_sha256":"a8e4c56c375d7a6542aaf8906655ba274ff2ceab3cafc12a45c645d3f731ba6e","abstract_canon_sha256":"b5c605241d918ab25416a5a8d2c557d20020e8cc3b2867b385525ddc63500b80"},"schema_version":"1.0"},"receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-07-05T02:59:15.547676Z","signature_b64":"VeRDuuC9KRY4iwhMnBnMWLtsd8E38vXksyIVAIbl/smgnLEzXdLoD3m3tLFbVriuyg5Von1G9R0BmcupJIkoAw==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"2be002d88dfafa8bc61fa6c37e20c2a17cbf63d38b8d40b7c0148b011b5d6596","last_reissued_at":"2026-07-05T02:59:15.547243Z","signature_status":"signed_v1","first_computed_at":"2026-07-05T02:59:15.547243Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"graph_snapshot":{"paper":{"title":"Reinforcement Learning with Prototypical Representations","license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","headline":"","cross_cats":["cs.AI"],"primary_cat":"cs.LG","authors_text":"Alessandro Lazaric, Denis Yarats, Lerrel Pinto, Rob Fergus","submitted_at":"2021-02-22T18:56:34Z","abstract_excerpt":"Learning effective representations in image-based environments is crucial for sample efficient Reinforcement Learning (RL). Unfortunately, in RL, representation learning is confounded with the exploratory experience of the agent -- learning a useful representation requires diverse data, while effective exploration is only possible with coherent representations. Furthermore, we would like to learn representations that not only generalize across tasks but also accelerate downstream exploration for efficient task-specific training. To address these challenges we propose Proto-RL, a self-supervise"},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2102.11271","kind":"arxiv","version":2},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2102.11271/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"aliases":[{"alias_kind":"arxiv","alias_value":"2102.11271","created_at":"2026-07-05T02:59:15.547310+00:00"},{"alias_kind":"arxiv_version","alias_value":"2102.11271v2","created_at":"2026-07-05T02:59:15.547310+00:00"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2102.11271","created_at":"2026-07-05T02:59:15.547310+00:00"},{"alias_kind":"pith_short_12","alias_value":"FPQAFWEN7L5I","created_at":"2026-07-05T02:59:15.547310+00:00"},{"alias_kind":"pith_short_16","alias_value":"FPQAFWEN7L5IXRQ7","created_at":"2026-07-05T02:59:15.547310+00:00"},{"alias_kind":"pith_short_8","alias_value":"FPQAFWEN","created_at":"2026-07-05T02:59:15.547310+00:00"}],"events":[],"event_summary":{},"paper_claims":[],"inbound_citations":{"count":0,"internal_anchor_count":0,"sample":[]},"formal_canon":{"evidence_count":0,"sample":[],"anchors":[]},"links":{"html":"https://pith.science/pith/FPQAFWEN7L5IXRQ7U3BX4IGCUF","json":"https://pith.science/pith/FPQAFWEN7L5IXRQ7U3BX4IGCUF.json","graph_json":"https://pith.science/api/pith-number/FPQAFWEN7L5IXRQ7U3BX4IGCUF/graph.json","events_json":"https://pith.science/api/pith-number/FPQAFWEN7L5IXRQ7U3BX4IGCUF/events.json","paper":"https://pith.science/paper/FPQAFWEN"},"agent_actions":{"view_html":"https://pith.science/pith/FPQAFWEN7L5IXRQ7U3BX4IGCUF","download_json":"https://pith.science/pith/FPQAFWEN7L5IXRQ7U3BX4IGCUF.json","view_paper":"https://pith.science/paper/FPQAFWEN","resolve_alias":"https://pith.science/api/pith-number/resolve?arxiv=2102.11271&json=true","fetch_graph":"https://pith.science/api/pith-number/FPQAFWEN7L5IXRQ7U3BX4IGCUF/graph.json","fetch_events":"https://pith.science/api/pith-number/FPQAFWEN7L5IXRQ7U3BX4IGCUF/events.json","actions":{"anchor_timestamp":"https://pith.science/pith/FPQAFWEN7L5IXRQ7U3BX4IGCUF/action/timestamp_anchor","attest_storage":"https://pith.science/pith/FPQAFWEN7L5IXRQ7U3BX4IGCUF/action/storage_attestation","attest_author":"https://pith.science/pith/FPQAFWEN7L5IXRQ7U3BX4IGCUF/action/author_attestation","sign_citation":"https://pith.science/pith/FPQAFWEN7L5IXRQ7U3BX4IGCUF/action/citation_signature","submit_replication":"https://pith.science/pith/FPQAFWEN7L5IXRQ7U3BX4IGCUF/action/replication_record"}},"created_at":"2026-07-05T02:59:15.547310+00:00","updated_at":"2026-07-05T02:59:15.547310+00:00"}