{"record_type":"pith_number_record","schema_url":"https://pith.science/schemas/pith-number/v1.json","pith_number":"pith:2024:7PO4OVECGY3WNX7LPRQYSX7SNA","short_pith_number":"pith:7PO4OVEC","schema_version":"1.0","canonical_sha256":"fbddc75482363766dfeb7c61895ff2682a5d984b7f5b0b340458ab388ddcfc4b","source":{"kind":"arxiv","id":"2402.08078","version":1},"attestation_state":"computed","paper":{"title":"Large Language Models as Agents in Two-Player Games","license":"http://creativecommons.org/licenses/by/4.0/","headline":"","cross_cats":["cs.LG"],"primary_cat":"cs.CL","authors_text":"Hang Li, Peng Sun, Yang Liu","submitted_at":"2024-02-12T21:44:32Z","abstract_excerpt":"By formally defining the training processes of large language models (LLMs), which usually encompasses pre-training, supervised fine-tuning, and reinforcement learning with human feedback, within a single and unified machine learning paradigm, we can glean pivotal insights for advancing LLM technologies. This position paper delineates the parallels between the training methods of LLMs and the strategies employed for the development of agents in two-player games, as studied in game theory, reinforcement learning, and multi-agent systems. We propose a re-conceptualization of LLM learning process"},"verification_status":{"content_addressed":true,"pith_receipt":true,"author_attested":false,"weak_author_claims":0,"strong_author_claims":0,"externally_anchored":false,"storage_verified":false,"citation_signatures":0,"replication_records":0,"graph_snapshot":true,"references_resolved":false,"formal_links_present":false},"canonical_record":{"source":{"id":"2402.08078","kind":"arxiv","version":1},"metadata":{"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.CL","submitted_at":"2024-02-12T21:44:32Z","cross_cats_sorted":["cs.LG"],"title_canon_sha256":"94333505da28d44dcbd781ff5d778826efe423c7e6d459cd74422365398586aa","abstract_canon_sha256":"3f8337a962b8bc3879fb509094acc886b4eb7dce3089c1eb720f37d995b0ac0c"},"schema_version":"1.0"},"receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-07-05T07:44:26.915790Z","signature_b64":"fvTcb+AgCHbLx9lBuGT0XFlchEhnf+GeTdPfMqs6AIbKk3dLu+1S0hVwIMQvO5Cvqwpb11syTCGvch8h3YHPAw==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"fbddc75482363766dfeb7c61895ff2682a5d984b7f5b0b340458ab388ddcfc4b","last_reissued_at":"2026-07-05T07:44:26.915383Z","signature_status":"signed_v1","first_computed_at":"2026-07-05T07:44:26.915383Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"graph_snapshot":{"paper":{"title":"Large Language Models as Agents in Two-Player Games","license":"http://creativecommons.org/licenses/by/4.0/","headline":"","cross_cats":["cs.LG"],"primary_cat":"cs.CL","authors_text":"Hang Li, Peng Sun, Yang Liu","submitted_at":"2024-02-12T21:44:32Z","abstract_excerpt":"By formally defining the training processes of large language models (LLMs), which usually encompasses pre-training, supervised fine-tuning, and reinforcement learning with human feedback, within a single and unified machine learning paradigm, we can glean pivotal insights for advancing LLM technologies. This position paper delineates the parallels between the training methods of LLMs and the strategies employed for the development of agents in two-player games, as studied in game theory, reinforcement learning, and multi-agent systems. We propose a re-conceptualization of LLM learning process"},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2402.08078","kind":"arxiv","version":1},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2402.08078/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"aliases":[{"alias_kind":"arxiv","alias_value":"2402.08078","created_at":"2026-07-05T07:44:26.915436+00:00"},{"alias_kind":"arxiv_version","alias_value":"2402.08078v1","created_at":"2026-07-05T07:44:26.915436+00:00"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2402.08078","created_at":"2026-07-05T07:44:26.915436+00:00"},{"alias_kind":"pith_short_12","alias_value":"7PO4OVECGY3W","created_at":"2026-07-05T07:44:26.915436+00:00"},{"alias_kind":"pith_short_16","alias_value":"7PO4OVECGY3WNX7L","created_at":"2026-07-05T07:44:26.915436+00:00"},{"alias_kind":"pith_short_8","alias_value":"7PO4OVEC","created_at":"2026-07-05T07:44:26.915436+00:00"}],"events":[],"event_summary":{},"paper_claims":[],"inbound_citations":{"count":2,"internal_anchor_count":0,"sample":[{"citing_arxiv_id":"2503.21460","citing_title":"Large Language Model Agent: A Survey on Methodology, Applications and Challenges","ref_index":111,"is_internal_anchor":false},{"citing_arxiv_id":"2605.08327","citing_title":"Interactive Critique-Revision Training for Reliable Structured LLM Generation","ref_index":24,"is_internal_anchor":false}]},"formal_canon":{"evidence_count":0,"sample":[],"anchors":[]},"links":{"html":"https://pith.science/pith/7PO4OVECGY3WNX7LPRQYSX7SNA","json":"https://pith.science/pith/7PO4OVECGY3WNX7LPRQYSX7SNA.json","graph_json":"https://pith.science/api/pith-number/7PO4OVECGY3WNX7LPRQYSX7SNA/graph.json","events_json":"https://pith.science/api/pith-number/7PO4OVECGY3WNX7LPRQYSX7SNA/events.json","paper":"https://pith.science/paper/7PO4OVEC"},"agent_actions":{"view_html":"https://pith.science/pith/7PO4OVECGY3WNX7LPRQYSX7SNA","download_json":"https://pith.science/pith/7PO4OVECGY3WNX7LPRQYSX7SNA.json","view_paper":"https://pith.science/paper/7PO4OVEC","resolve_alias":"https://pith.science/api/pith-number/resolve?arxiv=2402.08078&json=true","fetch_graph":"https://pith.science/api/pith-number/7PO4OVECGY3WNX7LPRQYSX7SNA/graph.json","fetch_events":"https://pith.science/api/pith-number/7PO4OVECGY3WNX7LPRQYSX7SNA/events.json","actions":{"anchor_timestamp":"https://pith.science/pith/7PO4OVECGY3WNX7LPRQYSX7SNA/action/timestamp_anchor","attest_storage":"https://pith.science/pith/7PO4OVECGY3WNX7LPRQYSX7SNA/action/storage_attestation","attest_author":"https://pith.science/pith/7PO4OVECGY3WNX7LPRQYSX7SNA/action/author_attestation","sign_citation":"https://pith.science/pith/7PO4OVECGY3WNX7LPRQYSX7SNA/action/citation_signature","submit_replication":"https://pith.science/pith/7PO4OVECGY3WNX7LPRQYSX7SNA/action/replication_record"}},"created_at":"2026-07-05T07:44:26.915436+00:00","updated_at":"2026-07-05T07:44:26.915436+00:00"}