{"record_type":"pith_number_record","schema_url":"https://pith.science/schemas/pith-number/v1.json","pith_number":"pith:2025:OF4HVW2CGHDKSY33OQ6NVUDCRF","short_pith_number":"pith:OF4HVW2C","schema_version":"1.0","canonical_sha256":"71787adb4231c6a9637b743cdad062895c61a69aa9550c6c9aaa0e8768cb52d2","source":{"kind":"arxiv","id":"2505.19637","version":1},"attestation_state":"computed","paper":{"title":"Adaptive Episode Length Adjustment for Multi-agent Reinforcement Learning","license":"http://creativecommons.org/licenses/by/4.0/","headline":"","cross_cats":[],"primary_cat":"cs.MA","authors_text":"Byunghyun Yoo, Euisok Chung, Hyunwoo Kim, Jeongmin Yang, Younghwan Shin","submitted_at":"2025-05-26T07:54:58Z","abstract_excerpt":"In standard reinforcement learning, an episode is defined as a sequence of interactions between agents and the environment, which terminates upon reaching a terminal state or a pre-defined episode length. Setting a shorter episode length enables the generation of multiple episodes with the same number of data samples, thereby facilitating an exploration of diverse states. While shorter episodes may limit the collection of long-term interactions, they may offer significant advantages when properly managed. For example, trajectory truncation in single-agent reinforcement learning has shown how t"},"verification_status":{"content_addressed":true,"pith_receipt":true,"author_attested":false,"weak_author_claims":0,"strong_author_claims":0,"externally_anchored":false,"storage_verified":false,"citation_signatures":0,"replication_records":0,"graph_snapshot":true,"references_resolved":false,"formal_links_present":false},"canonical_record":{"source":{"id":"2505.19637","kind":"arxiv","version":1},"metadata":{"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.MA","submitted_at":"2025-05-26T07:54:58Z","cross_cats_sorted":[],"title_canon_sha256":"571d5b1b3e1ce660cc75661fa57d5d4a559d44bd22f27f88bc3af7303689794f","abstract_canon_sha256":"1e7444cac7fb8b9e15d8e4b4ab97fb23b67d0ae91be1d0bcabd029093cb949a8"},"schema_version":"1.0"},"receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-07-05T11:09:33.850638Z","signature_b64":"TivGNBYGnsu011vfw40XVj/0mKj//zh6S+yH56pEIz0IeWSQ7lEVRMj2+PG2h2VYo7LuXHvLCupeXzg0sYgUAw==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"71787adb4231c6a9637b743cdad062895c61a69aa9550c6c9aaa0e8768cb52d2","last_reissued_at":"2026-07-05T11:09:33.850139Z","signature_status":"signed_v1","first_computed_at":"2026-07-05T11:09:33.850139Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"graph_snapshot":{"paper":{"title":"Adaptive Episode Length Adjustment for Multi-agent Reinforcement Learning","license":"http://creativecommons.org/licenses/by/4.0/","headline":"","cross_cats":[],"primary_cat":"cs.MA","authors_text":"Byunghyun Yoo, Euisok Chung, Hyunwoo Kim, Jeongmin Yang, Younghwan Shin","submitted_at":"2025-05-26T07:54:58Z","abstract_excerpt":"In standard reinforcement learning, an episode is defined as a sequence of interactions between agents and the environment, which terminates upon reaching a terminal state or a pre-defined episode length. Setting a shorter episode length enables the generation of multiple episodes with the same number of data samples, thereby facilitating an exploration of diverse states. While shorter episodes may limit the collection of long-term interactions, they may offer significant advantages when properly managed. For example, trajectory truncation in single-agent reinforcement learning has shown how t"},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2505.19637","kind":"arxiv","version":1},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2505.19637/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"aliases":[{"alias_kind":"arxiv","alias_value":"2505.19637","created_at":"2026-07-05T11:09:33.850202+00:00"},{"alias_kind":"arxiv_version","alias_value":"2505.19637v1","created_at":"2026-07-05T11:09:33.850202+00:00"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2505.19637","created_at":"2026-07-05T11:09:33.850202+00:00"},{"alias_kind":"pith_short_12","alias_value":"OF4HVW2CGHDK","created_at":"2026-07-05T11:09:33.850202+00:00"},{"alias_kind":"pith_short_16","alias_value":"OF4HVW2CGHDKSY33","created_at":"2026-07-05T11:09:33.850202+00:00"},{"alias_kind":"pith_short_8","alias_value":"OF4HVW2C","created_at":"2026-07-05T11:09:33.850202+00:00"}],"events":[],"event_summary":{},"paper_claims":[],"inbound_citations":{"count":0,"internal_anchor_count":0,"sample":[]},"formal_canon":{"evidence_count":0,"sample":[],"anchors":[]},"links":{"html":"https://pith.science/pith/OF4HVW2CGHDKSY33OQ6NVUDCRF","json":"https://pith.science/pith/OF4HVW2CGHDKSY33OQ6NVUDCRF.json","graph_json":"https://pith.science/api/pith-number/OF4HVW2CGHDKSY33OQ6NVUDCRF/graph.json","events_json":"https://pith.science/api/pith-number/OF4HVW2CGHDKSY33OQ6NVUDCRF/events.json","paper":"https://pith.science/paper/OF4HVW2C"},"agent_actions":{"view_html":"https://pith.science/pith/OF4HVW2CGHDKSY33OQ6NVUDCRF","download_json":"https://pith.science/pith/OF4HVW2CGHDKSY33OQ6NVUDCRF.json","view_paper":"https://pith.science/paper/OF4HVW2C","resolve_alias":"https://pith.science/api/pith-number/resolve?arxiv=2505.19637&json=true","fetch_graph":"https://pith.science/api/pith-number/OF4HVW2CGHDKSY33OQ6NVUDCRF/graph.json","fetch_events":"https://pith.science/api/pith-number/OF4HVW2CGHDKSY33OQ6NVUDCRF/events.json","actions":{"anchor_timestamp":"https://pith.science/pith/OF4HVW2CGHDKSY33OQ6NVUDCRF/action/timestamp_anchor","attest_storage":"https://pith.science/pith/OF4HVW2CGHDKSY33OQ6NVUDCRF/action/storage_attestation","attest_author":"https://pith.science/pith/OF4HVW2CGHDKSY33OQ6NVUDCRF/action/author_attestation","sign_citation":"https://pith.science/pith/OF4HVW2CGHDKSY33OQ6NVUDCRF/action/citation_signature","submit_replication":"https://pith.science/pith/OF4HVW2CGHDKSY33OQ6NVUDCRF/action/replication_record"}},"created_at":"2026-07-05T11:09:33.850202+00:00","updated_at":"2026-07-05T11:09:33.850202+00:00"}