{"record_type":"pith_number_record","schema_url":"https://pith.science/schemas/pith-number/v1.json","pith_number":"pith:2023:N6CLQI7ROWV6IVKB2Q7EE2XTZD","short_pith_number":"pith:N6CLQI7R","schema_version":"1.0","canonical_sha256":"6f84b823f175abe45541d43e426af3c8c29f88624b89f1b5cc380eb25f3cd78e","source":{"kind":"arxiv","id":"2305.14450","version":2},"attestation_state":"computed","paper":{"title":"An Empirical Study on Information Extraction using Large Language Models","license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","headline":"","cross_cats":[],"primary_cat":"cs.CL","authors_text":"Benyou Wang, Chaohao Yang, Lu Liu, Prayag Tiwari, Ridong Han, Tao Peng, Xiang Wan","submitted_at":"2023-05-23T18:17:43Z","abstract_excerpt":"Human-like large language models (LLMs), especially the most powerful and popular ones in OpenAI's GPT family, have proven to be very helpful for many natural language processing (NLP) related tasks. Therefore, various attempts have been made to apply LLMs to information extraction (IE), which is a fundamental NLP task that involves extracting information from unstructured plain text. To demonstrate the latest representative progress in LLMs' information extraction ability, we assess the information extraction ability of GPT-4 (the latest version of GPT at the time of writing this paper) from "},"verification_status":{"content_addressed":true,"pith_receipt":true,"author_attested":false,"weak_author_claims":0,"strong_author_claims":0,"externally_anchored":false,"storage_verified":false,"citation_signatures":0,"replication_records":0,"graph_snapshot":true,"references_resolved":false,"formal_links_present":false},"canonical_record":{"source":{"id":"2305.14450","kind":"arxiv","version":2},"metadata":{"license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","primary_cat":"cs.CL","submitted_at":"2023-05-23T18:17:43Z","cross_cats_sorted":[],"title_canon_sha256":"b6f70a1983954aa12e90b4cef93ed1060f2dce14ce570d7b2a97259cd67d8ab7","abstract_canon_sha256":"2bf114a0afc9dce44e260b6e92fc32bce3cf103ab86b1d85468958fc0ac5ce74"},"schema_version":"1.0"},"receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-07-05T09:05:30.704682Z","signature_b64":"3jrdyJWi/MrXvdbhjKXFRI/wIVYnLfMfeL1qAfjiJk/+UAoirQnJrBMpF/5CdUr48Ws065y+NH6NGPW3o2XzBw==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"6f84b823f175abe45541d43e426af3c8c29f88624b89f1b5cc380eb25f3cd78e","last_reissued_at":"2026-07-05T09:05:30.704219Z","signature_status":"signed_v1","first_computed_at":"2026-07-05T09:05:30.704219Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"graph_snapshot":{"paper":{"title":"An Empirical Study on Information Extraction using Large Language Models","license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","headline":"","cross_cats":[],"primary_cat":"cs.CL","authors_text":"Benyou Wang, Chaohao Yang, Lu Liu, Prayag Tiwari, Ridong Han, Tao Peng, Xiang Wan","submitted_at":"2023-05-23T18:17:43Z","abstract_excerpt":"Human-like large language models (LLMs), especially the most powerful and popular ones in OpenAI's GPT family, have proven to be very helpful for many natural language processing (NLP) related tasks. Therefore, various attempts have been made to apply LLMs to information extraction (IE), which is a fundamental NLP task that involves extracting information from unstructured plain text. To demonstrate the latest representative progress in LLMs' information extraction ability, we assess the information extraction ability of GPT-4 (the latest version of GPT at the time of writing this paper) from "},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2305.14450","kind":"arxiv","version":2},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2305.14450/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"aliases":[{"alias_kind":"arxiv","alias_value":"2305.14450","created_at":"2026-07-05T09:05:30.704276+00:00"},{"alias_kind":"arxiv_version","alias_value":"2305.14450v2","created_at":"2026-07-05T09:05:30.704276+00:00"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2305.14450","created_at":"2026-07-05T09:05:30.704276+00:00"},{"alias_kind":"pith_short_12","alias_value":"N6CLQI7ROWV6","created_at":"2026-07-05T09:05:30.704276+00:00"},{"alias_kind":"pith_short_16","alias_value":"N6CLQI7ROWV6IVKB","created_at":"2026-07-05T09:05:30.704276+00:00"},{"alias_kind":"pith_short_8","alias_value":"N6CLQI7R","created_at":"2026-07-05T09:05:30.704276+00:00"}],"events":[],"event_summary":{},"paper_claims":[],"inbound_citations":{"count":5,"internal_anchor_count":0,"sample":[{"citing_arxiv_id":"2606.24734","citing_title":"Task Decomposition for Efficient Annotation","ref_index":169,"is_internal_anchor":false},{"citing_arxiv_id":"2606.28070","citing_title":"JD Oxygen AI Item Center (Oxygen AIIC) V1: An Industrial-Scale LLM/VLM-Centric Solution for Item Understanding, Management, and Applications","ref_index":84,"is_internal_anchor":false},{"citing_arxiv_id":"2606.28070","citing_title":"JD Oxygen AI Item Center (Oxygen AIIC) V1: An Industrial-Scale LLM/VLM-Centric Solution for Item Understanding, Management, and Applications","ref_index":84,"is_internal_anchor":false},{"citing_arxiv_id":"2606.01323","citing_title":"DiffuSent: Towards a Unified Diffusion Framework for Aspect-Based Sentiment Analysis","ref_index":11,"is_internal_anchor":false},{"citing_arxiv_id":"2605.18007","citing_title":"Semantic Reranking at Inference Time for Hard Examples in Rhetorical Role Labeling","ref_index":109,"is_internal_anchor":false}]},"formal_canon":{"evidence_count":0,"sample":[],"anchors":[]},"links":{"html":"https://pith.science/pith/N6CLQI7ROWV6IVKB2Q7EE2XTZD","json":"https://pith.science/pith/N6CLQI7ROWV6IVKB2Q7EE2XTZD.json","graph_json":"https://pith.science/api/pith-number/N6CLQI7ROWV6IVKB2Q7EE2XTZD/graph.json","events_json":"https://pith.science/api/pith-number/N6CLQI7ROWV6IVKB2Q7EE2XTZD/events.json","paper":"https://pith.science/paper/N6CLQI7R"},"agent_actions":{"view_html":"https://pith.science/pith/N6CLQI7ROWV6IVKB2Q7EE2XTZD","download_json":"https://pith.science/pith/N6CLQI7ROWV6IVKB2Q7EE2XTZD.json","view_paper":"https://pith.science/paper/N6CLQI7R","resolve_alias":"https://pith.science/api/pith-number/resolve?arxiv=2305.14450&json=true","fetch_graph":"https://pith.science/api/pith-number/N6CLQI7ROWV6IVKB2Q7EE2XTZD/graph.json","fetch_events":"https://pith.science/api/pith-number/N6CLQI7ROWV6IVKB2Q7EE2XTZD/events.json","actions":{"anchor_timestamp":"https://pith.science/pith/N6CLQI7ROWV6IVKB2Q7EE2XTZD/action/timestamp_anchor","attest_storage":"https://pith.science/pith/N6CLQI7ROWV6IVKB2Q7EE2XTZD/action/storage_attestation","attest_author":"https://pith.science/pith/N6CLQI7ROWV6IVKB2Q7EE2XTZD/action/author_attestation","sign_citation":"https://pith.science/pith/N6CLQI7ROWV6IVKB2Q7EE2XTZD/action/citation_signature","submit_replication":"https://pith.science/pith/N6CLQI7ROWV6IVKB2Q7EE2XTZD/action/replication_record"}},"created_at":"2026-07-05T09:05:30.704276+00:00","updated_at":"2026-07-05T09:05:30.704276+00:00"}