{"record_type":"pith_number_record","schema_url":"https://pith.science/schemas/pith-number/v1.json","pith_number":"pith:2025:C7YBFUM7ZXTHOUH53V3CEA5WCP","short_pith_number":"pith:C7YBFUM7","schema_version":"1.0","canonical_sha256":"17f012d19fcde67750fddd762203b613fd6cb00b625df91260d0d9efa8ea8d44","source":{"kind":"arxiv","id":"2502.09589","version":2},"attestation_state":"computed","paper":{"title":"Logical forms complement probability in understanding language model (and human) performance","license":"http://creativecommons.org/licenses/by/4.0/","headline":"","cross_cats":["cs.LO"],"primary_cat":"cs.CL","authors_text":"Freda Shi, Yixuan Wang","submitted_at":"2025-02-13T18:46:44Z","abstract_excerpt":"With the increasing interest in using large language models (LLMs) for planning in natural language, understanding their behaviors becomes an important research question. This work conducts a systematic investigation of LLMs' ability to perform logical reasoning in natural language. We introduce a controlled dataset of hypothetical and disjunctive syllogisms in propositional and modal logic and use it as the testbed for understanding LLM performance. Our results lead to novel insights in predicting LLM behaviors: in addition to the probability of input (Gonen et al., 2023; McCoy et al., 2024),"},"verification_status":{"content_addressed":true,"pith_receipt":true,"author_attested":false,"weak_author_claims":0,"strong_author_claims":0,"externally_anchored":false,"storage_verified":false,"citation_signatures":0,"replication_records":0,"graph_snapshot":true,"references_resolved":false,"formal_links_present":false},"canonical_record":{"source":{"id":"2502.09589","kind":"arxiv","version":2},"metadata":{"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.CL","submitted_at":"2025-02-13T18:46:44Z","cross_cats_sorted":["cs.LO"],"title_canon_sha256":"85a47e3196dcef4766f0f7b7448d556a9740b18351de160bad3a8b8cc004b320","abstract_canon_sha256":"963b10636be4cc785b484eb6017046803c6dc277fc458e9b146a9a5f4ef227ae"},"schema_version":"1.0"},"receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-07-05T10:15:26.676182Z","signature_b64":"wC6VUsx4bUrM++q30PNr+OKH3mVFZdRsnrRfOcX9sQX3Bwys2qpebBE+Up4mWipUiHF9qGHB1FZ+EJTbYXOaAw==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"17f012d19fcde67750fddd762203b613fd6cb00b625df91260d0d9efa8ea8d44","last_reissued_at":"2026-07-05T10:15:26.675690Z","signature_status":"signed_v1","first_computed_at":"2026-07-05T10:15:26.675690Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"graph_snapshot":{"paper":{"title":"Logical forms complement probability in understanding language model (and human) performance","license":"http://creativecommons.org/licenses/by/4.0/","headline":"","cross_cats":["cs.LO"],"primary_cat":"cs.CL","authors_text":"Freda Shi, Yixuan Wang","submitted_at":"2025-02-13T18:46:44Z","abstract_excerpt":"With the increasing interest in using large language models (LLMs) for planning in natural language, understanding their behaviors becomes an important research question. This work conducts a systematic investigation of LLMs' ability to perform logical reasoning in natural language. We introduce a controlled dataset of hypothetical and disjunctive syllogisms in propositional and modal logic and use it as the testbed for understanding LLM performance. Our results lead to novel insights in predicting LLM behaviors: in addition to the probability of input (Gonen et al., 2023; McCoy et al., 2024),"},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2502.09589","kind":"arxiv","version":2},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2502.09589/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"aliases":[{"alias_kind":"arxiv","alias_value":"2502.09589","created_at":"2026-07-05T10:15:26.675743+00:00"},{"alias_kind":"arxiv_version","alias_value":"2502.09589v2","created_at":"2026-07-05T10:15:26.675743+00:00"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2502.09589","created_at":"2026-07-05T10:15:26.675743+00:00"},{"alias_kind":"pith_short_12","alias_value":"C7YBFUM7ZXTH","created_at":"2026-07-05T10:15:26.675743+00:00"},{"alias_kind":"pith_short_16","alias_value":"C7YBFUM7ZXTHOUH5","created_at":"2026-07-05T10:15:26.675743+00:00"},{"alias_kind":"pith_short_8","alias_value":"C7YBFUM7","created_at":"2026-07-05T10:15:26.675743+00:00"}],"events":[],"event_summary":{},"paper_claims":[],"inbound_citations":{"count":1,"internal_anchor_count":0,"sample":[{"citing_arxiv_id":"2605.09271","citing_title":"Shaping Schema via Language Representation as the Next Frontier for LLM Intelligence Expanding","ref_index":131,"is_internal_anchor":false}]},"formal_canon":{"evidence_count":0,"sample":[],"anchors":[]},"links":{"html":"https://pith.science/pith/C7YBFUM7ZXTHOUH53V3CEA5WCP","json":"https://pith.science/pith/C7YBFUM7ZXTHOUH53V3CEA5WCP.json","graph_json":"https://pith.science/api/pith-number/C7YBFUM7ZXTHOUH53V3CEA5WCP/graph.json","events_json":"https://pith.science/api/pith-number/C7YBFUM7ZXTHOUH53V3CEA5WCP/events.json","paper":"https://pith.science/paper/C7YBFUM7"},"agent_actions":{"view_html":"https://pith.science/pith/C7YBFUM7ZXTHOUH53V3CEA5WCP","download_json":"https://pith.science/pith/C7YBFUM7ZXTHOUH53V3CEA5WCP.json","view_paper":"https://pith.science/paper/C7YBFUM7","resolve_alias":"https://pith.science/api/pith-number/resolve?arxiv=2502.09589&json=true","fetch_graph":"https://pith.science/api/pith-number/C7YBFUM7ZXTHOUH53V3CEA5WCP/graph.json","fetch_events":"https://pith.science/api/pith-number/C7YBFUM7ZXTHOUH53V3CEA5WCP/events.json","actions":{"anchor_timestamp":"https://pith.science/pith/C7YBFUM7ZXTHOUH53V3CEA5WCP/action/timestamp_anchor","attest_storage":"https://pith.science/pith/C7YBFUM7ZXTHOUH53V3CEA5WCP/action/storage_attestation","attest_author":"https://pith.science/pith/C7YBFUM7ZXTHOUH53V3CEA5WCP/action/author_attestation","sign_citation":"https://pith.science/pith/C7YBFUM7ZXTHOUH53V3CEA5WCP/action/citation_signature","submit_replication":"https://pith.science/pith/C7YBFUM7ZXTHOUH53V3CEA5WCP/action/replication_record"}},"created_at":"2026-07-05T10:15:26.675743+00:00","updated_at":"2026-07-05T10:15:26.675743+00:00"}