{"record_type":"pith_number_record","schema_url":"https://pith.science/schemas/pith-number/v1.json","pith_number":"pith:2026:24FB2JDI7UWGTMOLI5KQRO4VQM","short_pith_number":"pith:24FB2JDI","schema_version":"1.0","canonical_sha256":"d70a1d2468fd2c69b1cb475508bb95832952161ac4dc8e6f290ff81ca218bb71","source":{"kind":"arxiv","id":"2607.29419","version":1},"attestation_state":"computed","paper":{"title":"Explore Beyond the Boundary Using Entropic Information","license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","headline":"","cross_cats":["cs.AI"],"primary_cat":"cs.LG","authors_text":"Bumgeun Park, Donghwan Lee","submitted_at":"2026-07-31T13:41:51Z","abstract_excerpt":"In reinforcement learning, exploration with sparse and delayed rewards presents a significant challenge due to the limited feedback available for guiding the learning process. Addressing this issue requires extensive exploration in the state space to discover valuable reward signals. In this paper, we propose Entropic Information for Exploration (ENTINEX), a novel method that enhances exploration by incentivizing agents to explore beyond the boundaries of the state distribution. ENTINEX achieves this by assigning intrinsic rewards to these boundaries, leveraging entropic information to identif"},"verification_status":{"content_addressed":true,"pith_receipt":true,"author_attested":false,"weak_author_claims":0,"strong_author_claims":0,"externally_anchored":false,"storage_verified":false,"citation_signatures":0,"replication_records":0,"graph_snapshot":true,"references_resolved":false,"formal_links_present":false},"canonical_record":{"source":{"id":"2607.29419","kind":"arxiv","version":1},"metadata":{"license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","primary_cat":"cs.LG","submitted_at":"2026-07-31T13:41:51Z","cross_cats_sorted":["cs.AI"],"title_canon_sha256":"d754c31c19e38db182427f0ba726a8700b60ca139d30264d0d22084159e3e3e2","abstract_canon_sha256":"1fc543cf01f6d24ae49dc92dc4ad0b5a4d4cf6fc18c21311bc4a2697f12f029e"},"schema_version":"1.0"},"receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-08-03T01:34:13.596480Z","signature_b64":"ixyKzmWRU04HUzkZmH0B5ydHQ17lmmoln1eI8WyX+hDroIKiAuEU2LsHFyYRTdQ+9DnSXcRJSx/8hRH1iiVaAA==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"d70a1d2468fd2c69b1cb475508bb95832952161ac4dc8e6f290ff81ca218bb71","last_reissued_at":"2026-08-03T01:34:13.594964Z","signature_status":"signed_v1","first_computed_at":"2026-08-03T01:34:13.594964Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"graph_snapshot":{"paper":{"title":"Explore Beyond the Boundary Using Entropic Information","license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","headline":"","cross_cats":["cs.AI"],"primary_cat":"cs.LG","authors_text":"Bumgeun Park, Donghwan Lee","submitted_at":"2026-07-31T13:41:51Z","abstract_excerpt":"In reinforcement learning, exploration with sparse and delayed rewards presents a significant challenge due to the limited feedback available for guiding the learning process. Addressing this issue requires extensive exploration in the state space to discover valuable reward signals. In this paper, we propose Entropic Information for Exploration (ENTINEX), a novel method that enhances exploration by incentivizing agents to explore beyond the boundaries of the state distribution. ENTINEX achieves this by assigning intrinsic rewards to these boundaries, leveraging entropic information to identif"},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2607.29419","kind":"arxiv","version":1},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2607.29419/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"aliases":[{"alias_kind":"arxiv","alias_value":"2607.29419","created_at":"2026-08-03T01:34:13.595860+00:00"},{"alias_kind":"arxiv_version","alias_value":"2607.29419v1","created_at":"2026-08-03T01:34:13.595860+00:00"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2607.29419","created_at":"2026-08-03T01:34:13.595860+00:00"},{"alias_kind":"pith_short_12","alias_value":"24FB2JDI7UWG","created_at":"2026-08-03T01:34:13.595860+00:00"},{"alias_kind":"pith_short_16","alias_value":"24FB2JDI7UWGTMOL","created_at":"2026-08-03T01:34:13.595860+00:00"},{"alias_kind":"pith_short_8","alias_value":"24FB2JDI","created_at":"2026-08-03T01:34:13.595860+00:00"}],"events":[],"event_summary":{},"paper_claims":[],"inbound_citations":{"count":0,"internal_anchor_count":0,"sample":[]},"formal_canon":{"evidence_count":0,"sample":[],"anchors":[]},"links":{"html":"https://pith.science/pith/24FB2JDI7UWGTMOLI5KQRO4VQM","json":"https://pith.science/pith/24FB2JDI7UWGTMOLI5KQRO4VQM.json","graph_json":"https://pith.science/api/pith-number/24FB2JDI7UWGTMOLI5KQRO4VQM/graph.json","events_json":"https://pith.science/api/pith-number/24FB2JDI7UWGTMOLI5KQRO4VQM/events.json","paper":"https://pith.science/paper/24FB2JDI"},"agent_actions":{"view_html":"https://pith.science/pith/24FB2JDI7UWGTMOLI5KQRO4VQM","download_json":"https://pith.science/pith/24FB2JDI7UWGTMOLI5KQRO4VQM.json","view_paper":"https://pith.science/paper/24FB2JDI","resolve_alias":"https://pith.science/api/pith-number/resolve?arxiv=2607.29419&json=true","fetch_graph":"https://pith.science/api/pith-number/24FB2JDI7UWGTMOLI5KQRO4VQM/graph.json","fetch_events":"https://pith.science/api/pith-number/24FB2JDI7UWGTMOLI5KQRO4VQM/events.json","actions":{"anchor_timestamp":"https://pith.science/pith/24FB2JDI7UWGTMOLI5KQRO4VQM/action/timestamp_anchor","attest_storage":"https://pith.science/pith/24FB2JDI7UWGTMOLI5KQRO4VQM/action/storage_attestation","attest_author":"https://pith.science/pith/24FB2JDI7UWGTMOLI5KQRO4VQM/action/author_attestation","sign_citation":"https://pith.science/pith/24FB2JDI7UWGTMOLI5KQRO4VQM/action/citation_signature","submit_replication":"https://pith.science/pith/24FB2JDI7UWGTMOLI5KQRO4VQM/action/replication_record"}},"created_at":"2026-08-03T01:34:13.595860+00:00","updated_at":"2026-08-03T01:34:13.595860+00:00"}