{"record_type":"pith_number_record","schema_url":"https://pith.science/schemas/pith-number/v1.json","pith_number":"pith:2024:GVOEBWS2VTA5RSQAG2G7ZYUSPU","short_pith_number":"pith:GVOEBWS2","schema_version":"1.0","canonical_sha256":"355c40da5aacc1d8ca00368dfce2927d0b3e77711f03d0bbd44a26374b0c9d16","source":{"kind":"arxiv","id":"2404.07376","version":2},"attestation_state":"computed","paper":{"title":"LLMs in Biomedicine: A study on clinical Named Entity Recognition","license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","headline":"","cross_cats":[],"primary_cat":"cs.CL","authors_text":"Fazlolah Mohaghegh, Jiaxin Yang, Joel Stremmel, Kai-Wei Chang, Masoud Monajatipoor, Melika Emami, Mozhdeh Rouhsedaghat","submitted_at":"2024-04-10T22:26:26Z","abstract_excerpt":"Large Language Models (LLMs) demonstrate remarkable versatility in various NLP tasks but encounter distinct challenges in biomedical due to the complexities of language and data scarcity. This paper investigates LLMs application in the biomedical domain by exploring strategies to enhance their performance for the NER task. Our study reveals the importance of meticulously designed prompts in the biomedical. Strategic selection of in-context examples yields a marked improvement, offering ~15-20\\% increase in F1 score across all benchmark datasets for biomedical few-shot NER. Additionally, our re"},"verification_status":{"content_addressed":true,"pith_receipt":true,"author_attested":false,"weak_author_claims":0,"strong_author_claims":0,"externally_anchored":false,"storage_verified":false,"citation_signatures":0,"replication_records":0,"graph_snapshot":true,"references_resolved":false,"formal_links_present":false},"canonical_record":{"source":{"id":"2404.07376","kind":"arxiv","version":2},"metadata":{"license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","primary_cat":"cs.CL","submitted_at":"2024-04-10T22:26:26Z","cross_cats_sorted":[],"title_canon_sha256":"c16f042ce4cf9d54eec5c3ac4215fa0cf564e6185a67555fae699f5826ac5483","abstract_canon_sha256":"555a683bd4fd39ea68272f0ef41b7787695620916545b372308ed90ce58b2326"},"schema_version":"1.0"},"receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-07-05T08:42:35.108387Z","signature_b64":"HLcNLMGW66EFzTGBkycHU38Kj8KDq2W2EANjn4j37hzgPDQb/3G/mkYRBD3KoEPPccictMn2dyG2XhA4xJvpAw==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"355c40da5aacc1d8ca00368dfce2927d0b3e77711f03d0bbd44a26374b0c9d16","last_reissued_at":"2026-07-05T08:42:35.107900Z","signature_status":"signed_v1","first_computed_at":"2026-07-05T08:42:35.107900Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"graph_snapshot":{"paper":{"title":"LLMs in Biomedicine: A study on clinical Named Entity Recognition","license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","headline":"","cross_cats":[],"primary_cat":"cs.CL","authors_text":"Fazlolah Mohaghegh, Jiaxin Yang, Joel Stremmel, Kai-Wei Chang, Masoud Monajatipoor, Melika Emami, Mozhdeh Rouhsedaghat","submitted_at":"2024-04-10T22:26:26Z","abstract_excerpt":"Large Language Models (LLMs) demonstrate remarkable versatility in various NLP tasks but encounter distinct challenges in biomedical due to the complexities of language and data scarcity. This paper investigates LLMs application in the biomedical domain by exploring strategies to enhance their performance for the NER task. Our study reveals the importance of meticulously designed prompts in the biomedical. Strategic selection of in-context examples yields a marked improvement, offering ~15-20\\% increase in F1 score across all benchmark datasets for biomedical few-shot NER. Additionally, our re"},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2404.07376","kind":"arxiv","version":2},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2404.07376/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"aliases":[{"alias_kind":"arxiv","alias_value":"2404.07376","created_at":"2026-07-05T08:42:35.107957+00:00"},{"alias_kind":"arxiv_version","alias_value":"2404.07376v2","created_at":"2026-07-05T08:42:35.107957+00:00"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2404.07376","created_at":"2026-07-05T08:42:35.107957+00:00"},{"alias_kind":"pith_short_12","alias_value":"GVOEBWS2VTA5","created_at":"2026-07-05T08:42:35.107957+00:00"},{"alias_kind":"pith_short_16","alias_value":"GVOEBWS2VTA5RSQA","created_at":"2026-07-05T08:42:35.107957+00:00"},{"alias_kind":"pith_short_8","alias_value":"GVOEBWS2","created_at":"2026-07-05T08:42:35.107957+00:00"}],"events":[],"event_summary":{},"paper_claims":[],"inbound_citations":{"count":5,"internal_anchor_count":0,"sample":[{"citing_arxiv_id":"2606.24734","citing_title":"Task Decomposition for Efficient Annotation","ref_index":34,"is_internal_anchor":false},{"citing_arxiv_id":"2502.16022","citing_title":"Enhancing LLMs for Identifying and Prioritizing Important Medical Jargons from Electronic Health Record Notes Utilizing Data Augmentation: A Comparative Study","ref_index":55,"is_internal_anchor":false},{"citing_arxiv_id":"2604.23059","citing_title":"Implicit Framing in Obstetric Counseling Notes: A Grounded LLM Pipeline on a VBAC-Eligible Cohort","ref_index":26,"is_internal_anchor":false},{"citing_arxiv_id":"2604.06403","citing_title":"FMI@SU ToxHabits: Evaluating LLMs Performance on Toxic Habit Extraction in Spanish Clinical Texts","ref_index":6,"is_internal_anchor":false},{"citing_arxiv_id":"2604.18944","citing_title":"A Mechanism and Optimization Study on the Impact of Information Density on User-Generated Content Named Entity Recognition","ref_index":21,"is_internal_anchor":false}]},"formal_canon":{"evidence_count":0,"sample":[],"anchors":[]},"links":{"html":"https://pith.science/pith/GVOEBWS2VTA5RSQAG2G7ZYUSPU","json":"https://pith.science/pith/GVOEBWS2VTA5RSQAG2G7ZYUSPU.json","graph_json":"https://pith.science/api/pith-number/GVOEBWS2VTA5RSQAG2G7ZYUSPU/graph.json","events_json":"https://pith.science/api/pith-number/GVOEBWS2VTA5RSQAG2G7ZYUSPU/events.json","paper":"https://pith.science/paper/GVOEBWS2"},"agent_actions":{"view_html":"https://pith.science/pith/GVOEBWS2VTA5RSQAG2G7ZYUSPU","download_json":"https://pith.science/pith/GVOEBWS2VTA5RSQAG2G7ZYUSPU.json","view_paper":"https://pith.science/paper/GVOEBWS2","resolve_alias":"https://pith.science/api/pith-number/resolve?arxiv=2404.07376&json=true","fetch_graph":"https://pith.science/api/pith-number/GVOEBWS2VTA5RSQAG2G7ZYUSPU/graph.json","fetch_events":"https://pith.science/api/pith-number/GVOEBWS2VTA5RSQAG2G7ZYUSPU/events.json","actions":{"anchor_timestamp":"https://pith.science/pith/GVOEBWS2VTA5RSQAG2G7ZYUSPU/action/timestamp_anchor","attest_storage":"https://pith.science/pith/GVOEBWS2VTA5RSQAG2G7ZYUSPU/action/storage_attestation","attest_author":"https://pith.science/pith/GVOEBWS2VTA5RSQAG2G7ZYUSPU/action/author_attestation","sign_citation":"https://pith.science/pith/GVOEBWS2VTA5RSQAG2G7ZYUSPU/action/citation_signature","submit_replication":"https://pith.science/pith/GVOEBWS2VTA5RSQAG2G7ZYUSPU/action/replication_record"}},"created_at":"2026-07-05T08:42:35.107957+00:00","updated_at":"2026-07-05T08:42:35.107957+00:00"}