{"record_type":"pith_number_record","schema_url":"https://pith.science/schemas/pith-number/v1.json","pith_number":"pith:2025:4LSU27AZJZDXDLRPRCW4566FNO","short_pith_number":"pith:4LSU27AZ","schema_version":"1.0","canonical_sha256":"e2e54d7c194e4771ae2f88adcefbc56b986fb3773dcb8d5339bffd8d01656f21","source":{"kind":"arxiv","id":"2502.09858","version":1},"attestation_state":"computed","paper":{"title":"Automated Hypothesis Validation with Agentic Sequential Falsifications","license":"http://creativecommons.org/licenses/by/4.0/","headline":"","cross_cats":["cs.AI","cs.CL","q-bio.QM"],"primary_cat":"cs.LG","authors_text":"Emmanuel Cand\\`es, Jure Leskovec, Kexin Huang, Michael Y. Li, Ryan Li, Ying Jin","submitted_at":"2025-02-14T01:46:00Z","abstract_excerpt":"Hypotheses are central to information acquisition, decision-making, and discovery. However, many real-world hypotheses are abstract, high-level statements that are difficult to validate directly. This challenge is further intensified by the rise of hypothesis generation from Large Language Models (LLMs), which are prone to hallucination and produce hypotheses in volumes that make manual validation impractical. Here we propose Popper, an agentic framework for rigorous automated validation of free-form hypotheses. Guided by Karl Popper's principle of falsification, Popper validates a hypothesis "},"verification_status":{"content_addressed":true,"pith_receipt":true,"author_attested":false,"weak_author_claims":0,"strong_author_claims":0,"externally_anchored":false,"storage_verified":false,"citation_signatures":0,"replication_records":0,"graph_snapshot":true,"references_resolved":false,"formal_links_present":false},"canonical_record":{"source":{"id":"2502.09858","kind":"arxiv","version":1},"metadata":{"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.LG","submitted_at":"2025-02-14T01:46:00Z","cross_cats_sorted":["cs.AI","cs.CL","q-bio.QM"],"title_canon_sha256":"9773c8ec5cb876629da699b00724ac92ab06733a5c7eed8c96c1613b8411105c","abstract_canon_sha256":"b84d4ddcbf6eff781bba385e692f13785197ab2a9470368f0faa0184250fbdd5"},"schema_version":"1.0"},"receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-07-05T10:14:13.428911Z","signature_b64":"uLGcqoDO8Nep9oZ2dedcLB/SVpqDTmXjqdoG9LXnEtH/KLjh1Y+MoapZSlk0opLMz85p3uUdtaCbC0t302U8CQ==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"e2e54d7c194e4771ae2f88adcefbc56b986fb3773dcb8d5339bffd8d01656f21","last_reissued_at":"2026-07-05T10:14:13.428399Z","signature_status":"signed_v1","first_computed_at":"2026-07-05T10:14:13.428399Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"graph_snapshot":{"paper":{"title":"Automated Hypothesis Validation with Agentic Sequential Falsifications","license":"http://creativecommons.org/licenses/by/4.0/","headline":"","cross_cats":["cs.AI","cs.CL","q-bio.QM"],"primary_cat":"cs.LG","authors_text":"Emmanuel Cand\\`es, Jure Leskovec, Kexin Huang, Michael Y. Li, Ryan Li, Ying Jin","submitted_at":"2025-02-14T01:46:00Z","abstract_excerpt":"Hypotheses are central to information acquisition, decision-making, and discovery. However, many real-world hypotheses are abstract, high-level statements that are difficult to validate directly. This challenge is further intensified by the rise of hypothesis generation from Large Language Models (LLMs), which are prone to hallucination and produce hypotheses in volumes that make manual validation impractical. Here we propose Popper, an agentic framework for rigorous automated validation of free-form hypotheses. Guided by Karl Popper's principle of falsification, Popper validates a hypothesis "},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2502.09858","kind":"arxiv","version":1},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2502.09858/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"aliases":[{"alias_kind":"arxiv","alias_value":"2502.09858","created_at":"2026-07-05T10:14:13.428467+00:00"},{"alias_kind":"arxiv_version","alias_value":"2502.09858v1","created_at":"2026-07-05T10:14:13.428467+00:00"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2502.09858","created_at":"2026-07-05T10:14:13.428467+00:00"},{"alias_kind":"pith_short_12","alias_value":"4LSU27AZJZDX","created_at":"2026-07-05T10:14:13.428467+00:00"},{"alias_kind":"pith_short_16","alias_value":"4LSU27AZJZDXDLRP","created_at":"2026-07-05T10:14:13.428467+00:00"},{"alias_kind":"pith_short_8","alias_value":"4LSU27AZ","created_at":"2026-07-05T10:14:13.428467+00:00"}],"events":[],"event_summary":{},"paper_claims":[],"inbound_citations":{"count":4,"internal_anchor_count":0,"sample":[{"citing_arxiv_id":"2605.10224","citing_title":"Hypothesis-Driven Deep Research with Large Language Models: A Structured Methodology for Automated Knowledge Discovery","ref_index":3,"is_internal_anchor":false},{"citing_arxiv_id":"2604.12144","citing_title":"VERITAS: A Multi-Agent Co-Scientist for Verifiable Image-Derived Hypothesis Testing","ref_index":25,"is_internal_anchor":false},{"citing_arxiv_id":"2604.08501","citing_title":"sciwrite-lint: Verification Infrastructure for the Age of Science Vibe-Writing","ref_index":17,"is_internal_anchor":false},{"citing_arxiv_id":"2604.19049","citing_title":"Refute-or-Promote: An Adversarial Stage-Gated Multi-Agent Review Methodology for High-Precision LLM-Assisted Defect Discovery","ref_index":13,"is_internal_anchor":false}]},"formal_canon":{"evidence_count":0,"sample":[],"anchors":[]},"links":{"html":"https://pith.science/pith/4LSU27AZJZDXDLRPRCW4566FNO","json":"https://pith.science/pith/4LSU27AZJZDXDLRPRCW4566FNO.json","graph_json":"https://pith.science/api/pith-number/4LSU27AZJZDXDLRPRCW4566FNO/graph.json","events_json":"https://pith.science/api/pith-number/4LSU27AZJZDXDLRPRCW4566FNO/events.json","paper":"https://pith.science/paper/4LSU27AZ"},"agent_actions":{"view_html":"https://pith.science/pith/4LSU27AZJZDXDLRPRCW4566FNO","download_json":"https://pith.science/pith/4LSU27AZJZDXDLRPRCW4566FNO.json","view_paper":"https://pith.science/paper/4LSU27AZ","resolve_alias":"https://pith.science/api/pith-number/resolve?arxiv=2502.09858&json=true","fetch_graph":"https://pith.science/api/pith-number/4LSU27AZJZDXDLRPRCW4566FNO/graph.json","fetch_events":"https://pith.science/api/pith-number/4LSU27AZJZDXDLRPRCW4566FNO/events.json","actions":{"anchor_timestamp":"https://pith.science/pith/4LSU27AZJZDXDLRPRCW4566FNO/action/timestamp_anchor","attest_storage":"https://pith.science/pith/4LSU27AZJZDXDLRPRCW4566FNO/action/storage_attestation","attest_author":"https://pith.science/pith/4LSU27AZJZDXDLRPRCW4566FNO/action/author_attestation","sign_citation":"https://pith.science/pith/4LSU27AZJZDXDLRPRCW4566FNO/action/citation_signature","submit_replication":"https://pith.science/pith/4LSU27AZJZDXDLRPRCW4566FNO/action/replication_record"}},"created_at":"2026-07-05T10:14:13.428467+00:00","updated_at":"2026-07-05T10:14:13.428467+00:00"}