{"record_type":"pith_number_record","schema_url":"https://pith.science/schemas/pith-number/v1.json","pith_number":"pith:2020:UVCGA2JRDG6GXTH6EIEGYJJKMT","short_pith_number":"pith:UVCGA2JR","schema_version":"1.0","canonical_sha256":"a54460693119bc6bccfe22086c252a64e77c6d08d9e2c2ab87eb2010b68a2a2a","source":{"kind":"arxiv","id":"2004.04228","version":1},"attestation_state":"computed","paper":{"title":"Asking and Answering Questions to Evaluate the Factual Consistency of Summaries","license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","headline":"","cross_cats":[],"primary_cat":"cs.CL","authors_text":"Alex Wang, Kyunghyun Cho, Mike Lewis","submitted_at":"2020-04-08T20:01:09Z","abstract_excerpt":"Practical applications of abstractive summarization models are limited by frequent factual inconsistencies with respect to their input. Existing automatic evaluation metrics for summarization are largely insensitive to such errors. We propose an automatic evaluation protocol called QAGS (pronounced \"kags\") that is designed to identify factual inconsistencies in a generated summary. QAGS is based on the intuition that if we ask questions about a summary and its source, we will receive similar answers if the summary is factually consistent with the source. To evaluate QAGS, we collect human judg"},"verification_status":{"content_addressed":true,"pith_receipt":true,"author_attested":false,"weak_author_claims":0,"strong_author_claims":0,"externally_anchored":false,"storage_verified":false,"citation_signatures":0,"replication_records":0,"graph_snapshot":true,"references_resolved":false,"formal_links_present":false},"canonical_record":{"source":{"id":"2004.04228","kind":"arxiv","version":1},"metadata":{"license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","primary_cat":"cs.CL","submitted_at":"2020-04-08T20:01:09Z","cross_cats_sorted":[],"title_canon_sha256":"bfb3c21ac93a2fc4897a071e3c05ba85c4e5cf3ae71865ee11b8ed22952c244a","abstract_canon_sha256":"8420722fb590e069f49cef7cf87becf37ad5fc138b0355b4bb46898b72ea5988"},"schema_version":"1.0"},"receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-07-05T00:54:05.917357Z","signature_b64":"s2qWTneqdkM/hq+aXWZ3YHOqtDPGfTjGQ7wlaTELD9pHW6ZG9AKoS0wEsQNIPPkUgTFVwpu/C0PcRVuRwNQuCQ==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"a54460693119bc6bccfe22086c252a64e77c6d08d9e2c2ab87eb2010b68a2a2a","last_reissued_at":"2026-07-05T00:54:05.916922Z","signature_status":"signed_v1","first_computed_at":"2026-07-05T00:54:05.916922Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"graph_snapshot":{"paper":{"title":"Asking and Answering Questions to Evaluate the Factual Consistency of Summaries","license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","headline":"","cross_cats":[],"primary_cat":"cs.CL","authors_text":"Alex Wang, Kyunghyun Cho, Mike Lewis","submitted_at":"2020-04-08T20:01:09Z","abstract_excerpt":"Practical applications of abstractive summarization models are limited by frequent factual inconsistencies with respect to their input. Existing automatic evaluation metrics for summarization are largely insensitive to such errors. We propose an automatic evaluation protocol called QAGS (pronounced \"kags\") that is designed to identify factual inconsistencies in a generated summary. QAGS is based on the intuition that if we ask questions about a summary and its source, we will receive similar answers if the summary is factually consistent with the source. To evaluate QAGS, we collect human judg"},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2004.04228","kind":"arxiv","version":1},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2004.04228/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"aliases":[{"alias_kind":"arxiv","alias_value":"2004.04228","created_at":"2026-07-05T00:54:05.916984+00:00"},{"alias_kind":"arxiv_version","alias_value":"2004.04228v1","created_at":"2026-07-05T00:54:05.916984+00:00"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2004.04228","created_at":"2026-07-05T00:54:05.916984+00:00"},{"alias_kind":"pith_short_12","alias_value":"UVCGA2JRDG6G","created_at":"2026-07-05T00:54:05.916984+00:00"},{"alias_kind":"pith_short_16","alias_value":"UVCGA2JRDG6GXTH6","created_at":"2026-07-05T00:54:05.916984+00:00"},{"alias_kind":"pith_short_8","alias_value":"UVCGA2JR","created_at":"2026-07-05T00:54:05.916984+00:00"}],"events":[],"event_summary":{},"paper_claims":[],"inbound_citations":{"count":4,"internal_anchor_count":0,"sample":[{"citing_arxiv_id":"2605.28073","citing_title":"StoryLens: Preference-Aligned Story Rewriting via Context-Aware Narrative Enrichment","ref_index":5,"is_internal_anchor":false},{"citing_arxiv_id":"2606.27446","citing_title":"Causal Connections: Leveraging Multilingual Fine-Tuning for Financial QA@FinCausal 2026","ref_index":47,"is_internal_anchor":false},{"citing_arxiv_id":"2606.27316","citing_title":"LLM-Based Examination of Eligibility Criteria from Securities Prospectuses at the German Central Bank","ref_index":17,"is_internal_anchor":false},{"citing_arxiv_id":"2604.24052","citing_title":"QEVA: A Reference-Free Evaluation Metric for Narrative Video Summarization with Multimodal Question Answering","ref_index":3,"is_internal_anchor":false}]},"formal_canon":{"evidence_count":0,"sample":[],"anchors":[]},"links":{"html":"https://pith.science/pith/UVCGA2JRDG6GXTH6EIEGYJJKMT","json":"https://pith.science/pith/UVCGA2JRDG6GXTH6EIEGYJJKMT.json","graph_json":"https://pith.science/api/pith-number/UVCGA2JRDG6GXTH6EIEGYJJKMT/graph.json","events_json":"https://pith.science/api/pith-number/UVCGA2JRDG6GXTH6EIEGYJJKMT/events.json","paper":"https://pith.science/paper/UVCGA2JR"},"agent_actions":{"view_html":"https://pith.science/pith/UVCGA2JRDG6GXTH6EIEGYJJKMT","download_json":"https://pith.science/pith/UVCGA2JRDG6GXTH6EIEGYJJKMT.json","view_paper":"https://pith.science/paper/UVCGA2JR","resolve_alias":"https://pith.science/api/pith-number/resolve?arxiv=2004.04228&json=true","fetch_graph":"https://pith.science/api/pith-number/UVCGA2JRDG6GXTH6EIEGYJJKMT/graph.json","fetch_events":"https://pith.science/api/pith-number/UVCGA2JRDG6GXTH6EIEGYJJKMT/events.json","actions":{"anchor_timestamp":"https://pith.science/pith/UVCGA2JRDG6GXTH6EIEGYJJKMT/action/timestamp_anchor","attest_storage":"https://pith.science/pith/UVCGA2JRDG6GXTH6EIEGYJJKMT/action/storage_attestation","attest_author":"https://pith.science/pith/UVCGA2JRDG6GXTH6EIEGYJJKMT/action/author_attestation","sign_citation":"https://pith.science/pith/UVCGA2JRDG6GXTH6EIEGYJJKMT/action/citation_signature","submit_replication":"https://pith.science/pith/UVCGA2JRDG6GXTH6EIEGYJJKMT/action/replication_record"}},"created_at":"2026-07-05T00:54:05.916984+00:00","updated_at":"2026-07-05T00:54:05.916984+00:00"}