{"record_type":"pith_number_record","schema_url":"https://pith.science/schemas/pith-number/v1.json","pith_number":"pith:2021:QXBAC7PD72RWGNRFRDR6ERONEP","short_pith_number":"pith:QXBAC7PD","schema_version":"1.0","canonical_sha256":"85c2017de3fea363362588e3e245cd23e4f40ba6f157ff92bf5995c3e1615ea3","source":{"kind":"arxiv","id":"2110.08470","version":3},"attestation_state":"computed","paper":{"title":"Case-based Reasoning for Better Generalization in Textual Reinforcement Learning","license":"http://creativecommons.org/licenses/by/4.0/","headline":"","cross_cats":["cs.LG"],"primary_cat":"cs.CL","authors_text":"Keerthiram Murugesan, Mattia Atzeni, Mrinmaya Sachan, Shehzaad Dhuliawala","submitted_at":"2021-10-16T04:51:34Z","abstract_excerpt":"Text-based games (TBG) have emerged as promising environments for driving research in grounded language understanding and studying problems like generalization and sample efficiency. Several deep reinforcement learning (RL) methods with varying architectures and learning schemes have been proposed for TBGs. However, these methods fail to generalize efficiently, especially under distributional shifts. In a departure from deep RL approaches, in this paper, we propose a general method inspired by case-based reasoning to train agents and generalize out of the training distribution. The case-based "},"verification_status":{"content_addressed":true,"pith_receipt":true,"author_attested":false,"weak_author_claims":0,"strong_author_claims":0,"externally_anchored":false,"storage_verified":false,"citation_signatures":0,"replication_records":0,"graph_snapshot":true,"references_resolved":false,"formal_links_present":false},"canonical_record":{"source":{"id":"2110.08470","kind":"arxiv","version":3},"metadata":{"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.CL","submitted_at":"2021-10-16T04:51:34Z","cross_cats_sorted":["cs.LG"],"title_canon_sha256":"f67e96600fbd29f51e01c980207630a801f39695c42194033aecb86ec479bce2","abstract_canon_sha256":"18b673791007e12cd6c7f8f796da6c886d38dc3937e9630a2b5e41be19918e17"},"schema_version":"1.0"},"receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-07-05T04:09:15.897760Z","signature_b64":"8ijS966rFJcRlExskC7JLqHvK+GHnXX53cwdNsjIcQdxCqOmHIir+GyTOI+hGfQ67lov1Bv3/eVLbvi9M6DzCA==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"85c2017de3fea363362588e3e245cd23e4f40ba6f157ff92bf5995c3e1615ea3","last_reissued_at":"2026-07-05T04:09:15.897267Z","signature_status":"signed_v1","first_computed_at":"2026-07-05T04:09:15.897267Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"graph_snapshot":{"paper":{"title":"Case-based Reasoning for Better Generalization in Textual Reinforcement Learning","license":"http://creativecommons.org/licenses/by/4.0/","headline":"","cross_cats":["cs.LG"],"primary_cat":"cs.CL","authors_text":"Keerthiram Murugesan, Mattia Atzeni, Mrinmaya Sachan, Shehzaad Dhuliawala","submitted_at":"2021-10-16T04:51:34Z","abstract_excerpt":"Text-based games (TBG) have emerged as promising environments for driving research in grounded language understanding and studying problems like generalization and sample efficiency. Several deep reinforcement learning (RL) methods with varying architectures and learning schemes have been proposed for TBGs. However, these methods fail to generalize efficiently, especially under distributional shifts. In a departure from deep RL approaches, in this paper, we propose a general method inspired by case-based reasoning to train agents and generalize out of the training distribution. The case-based "},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2110.08470","kind":"arxiv","version":3},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2110.08470/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"aliases":[{"alias_kind":"arxiv","alias_value":"2110.08470","created_at":"2026-07-05T04:09:15.897323+00:00"},{"alias_kind":"arxiv_version","alias_value":"2110.08470v3","created_at":"2026-07-05T04:09:15.897323+00:00"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2110.08470","created_at":"2026-07-05T04:09:15.897323+00:00"},{"alias_kind":"pith_short_12","alias_value":"QXBAC7PD72RW","created_at":"2026-07-05T04:09:15.897323+00:00"},{"alias_kind":"pith_short_16","alias_value":"QXBAC7PD72RWGNRF","created_at":"2026-07-05T04:09:15.897323+00:00"},{"alias_kind":"pith_short_8","alias_value":"QXBAC7PD","created_at":"2026-07-05T04:09:15.897323+00:00"}],"events":[],"event_summary":{},"paper_claims":[],"inbound_citations":{"count":1,"internal_anchor_count":1,"sample":[{"citing_arxiv_id":"2506.19527","citing_title":"KnowMap: Efficient Knowledge-Driven Task Adaptation for LLMs","ref_index":4,"is_internal_anchor":true}]},"formal_canon":{"evidence_count":0,"sample":[],"anchors":[]},"links":{"html":"https://pith.science/pith/QXBAC7PD72RWGNRFRDR6ERONEP","json":"https://pith.science/pith/QXBAC7PD72RWGNRFRDR6ERONEP.json","graph_json":"https://pith.science/api/pith-number/QXBAC7PD72RWGNRFRDR6ERONEP/graph.json","events_json":"https://pith.science/api/pith-number/QXBAC7PD72RWGNRFRDR6ERONEP/events.json","paper":"https://pith.science/paper/QXBAC7PD"},"agent_actions":{"view_html":"https://pith.science/pith/QXBAC7PD72RWGNRFRDR6ERONEP","download_json":"https://pith.science/pith/QXBAC7PD72RWGNRFRDR6ERONEP.json","view_paper":"https://pith.science/paper/QXBAC7PD","resolve_alias":"https://pith.science/api/pith-number/resolve?arxiv=2110.08470&json=true","fetch_graph":"https://pith.science/api/pith-number/QXBAC7PD72RWGNRFRDR6ERONEP/graph.json","fetch_events":"https://pith.science/api/pith-number/QXBAC7PD72RWGNRFRDR6ERONEP/events.json","actions":{"anchor_timestamp":"https://pith.science/pith/QXBAC7PD72RWGNRFRDR6ERONEP/action/timestamp_anchor","attest_storage":"https://pith.science/pith/QXBAC7PD72RWGNRFRDR6ERONEP/action/storage_attestation","attest_author":"https://pith.science/pith/QXBAC7PD72RWGNRFRDR6ERONEP/action/author_attestation","sign_citation":"https://pith.science/pith/QXBAC7PD72RWGNRFRDR6ERONEP/action/citation_signature","submit_replication":"https://pith.science/pith/QXBAC7PD72RWGNRFRDR6ERONEP/action/replication_record"}},"created_at":"2026-07-05T04:09:15.897323+00:00","updated_at":"2026-07-05T04:09:15.897323+00:00"}