{"record_type":"pith_number_record","schema_url":"https://pith.science/schemas/pith-number/v1.json","pith_number":"pith:2023:BYXUGBTDKH34F7TLSS5JGFT5IY","short_pith_number":"pith:BYXUGBTD","schema_version":"1.0","canonical_sha256":"0e2f43066351f7c2fe6b94ba93167d4622b6b174209e0e39830990b9777798b7","source":{"kind":"arxiv","id":"2306.14924","version":1},"attestation_state":"computed","paper":{"title":"LLM-Assisted Content Analysis: Using Large Language Models to Support Deductive Coding","license":"http://creativecommons.org/licenses/by/4.0/","headline":"","cross_cats":["cs.AI","cs.LG","stat.AP"],"primary_cat":"cs.CL","authors_text":"Annice Kim, Jessica Speer, John Bollenbacher, Michael Wenger, Robert Chew","submitted_at":"2023-06-23T20:57:32Z","abstract_excerpt":"Deductive coding is a widely used qualitative research method for determining the prevalence of themes across documents. While useful, deductive coding is often burdensome and time consuming since it requires researchers to read, interpret, and reliably categorize a large body of unstructured text documents. Large language models (LLMs), like ChatGPT, are a class of quickly evolving AI tools that can perform a range of natural language processing and reasoning tasks. In this study, we explore the use of LLMs to reduce the time it takes for deductive coding while retaining the flexibility of a "},"verification_status":{"content_addressed":true,"pith_receipt":true,"author_attested":false,"weak_author_claims":0,"strong_author_claims":0,"externally_anchored":false,"storage_verified":false,"citation_signatures":0,"replication_records":0,"graph_snapshot":true,"references_resolved":false,"formal_links_present":false},"canonical_record":{"source":{"id":"2306.14924","kind":"arxiv","version":1},"metadata":{"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.CL","submitted_at":"2023-06-23T20:57:32Z","cross_cats_sorted":["cs.AI","cs.LG","stat.AP"],"title_canon_sha256":"619152041370b0526c2e1971bade51af076c6f071a94411c54e7ce2e95f0f10e","abstract_canon_sha256":"2d07a90c69f9cfb6aab59f3d92d593dc479ac25aa054d2b6f0b07689dbb5794f"},"schema_version":"1.0"},"receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-07-05T06:24:51.874637Z","signature_b64":"v6zem/YqgiUAqqZ/qYe3HoDV6d4hLgm10TwfEoF29F9hWH/z3StGCeeRjaEjqM42spJWCK/CqA3TkFsE/lenCw==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"0e2f43066351f7c2fe6b94ba93167d4622b6b174209e0e39830990b9777798b7","last_reissued_at":"2026-07-05T06:24:51.874198Z","signature_status":"signed_v1","first_computed_at":"2026-07-05T06:24:51.874198Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"graph_snapshot":{"paper":{"title":"LLM-Assisted Content Analysis: Using Large Language Models to Support Deductive Coding","license":"http://creativecommons.org/licenses/by/4.0/","headline":"","cross_cats":["cs.AI","cs.LG","stat.AP"],"primary_cat":"cs.CL","authors_text":"Annice Kim, Jessica Speer, John Bollenbacher, Michael Wenger, Robert Chew","submitted_at":"2023-06-23T20:57:32Z","abstract_excerpt":"Deductive coding is a widely used qualitative research method for determining the prevalence of themes across documents. While useful, deductive coding is often burdensome and time consuming since it requires researchers to read, interpret, and reliably categorize a large body of unstructured text documents. Large language models (LLMs), like ChatGPT, are a class of quickly evolving AI tools that can perform a range of natural language processing and reasoning tasks. In this study, we explore the use of LLMs to reduce the time it takes for deductive coding while retaining the flexibility of a "},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2306.14924","kind":"arxiv","version":1},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2306.14924/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"aliases":[{"alias_kind":"arxiv","alias_value":"2306.14924","created_at":"2026-07-05T06:24:51.874264+00:00"},{"alias_kind":"arxiv_version","alias_value":"2306.14924v1","created_at":"2026-07-05T06:24:51.874264+00:00"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2306.14924","created_at":"2026-07-05T06:24:51.874264+00:00"},{"alias_kind":"pith_short_12","alias_value":"BYXUGBTDKH34","created_at":"2026-07-05T06:24:51.874264+00:00"},{"alias_kind":"pith_short_16","alias_value":"BYXUGBTDKH34F7TL","created_at":"2026-07-05T06:24:51.874264+00:00"},{"alias_kind":"pith_short_8","alias_value":"BYXUGBTD","created_at":"2026-07-05T06:24:51.874264+00:00"}],"events":[],"event_summary":{},"paper_claims":[],"inbound_citations":{"count":3,"internal_anchor_count":0,"sample":[{"citing_arxiv_id":"2408.09030","citing_title":"Effects of Collaboration on the Performance of Interactive Theme Discovery Systems","ref_index":5,"is_internal_anchor":false},{"citing_arxiv_id":"2410.13036","citing_title":"Uncovering the Internet's Hidden Values: An Empirical Study of Desirable Behavior Using Highly-Upvoted Content on Reddit","ref_index":16,"is_internal_anchor":false},{"citing_arxiv_id":"2507.11198","citing_title":"Temperature and Persona Shape LLM Agent Consensus With Minimal Accuracy Gains in Qualitative Coding","ref_index":19,"is_internal_anchor":false}]},"formal_canon":{"evidence_count":0,"sample":[],"anchors":[]},"links":{"html":"https://pith.science/pith/BYXUGBTDKH34F7TLSS5JGFT5IY","json":"https://pith.science/pith/BYXUGBTDKH34F7TLSS5JGFT5IY.json","graph_json":"https://pith.science/api/pith-number/BYXUGBTDKH34F7TLSS5JGFT5IY/graph.json","events_json":"https://pith.science/api/pith-number/BYXUGBTDKH34F7TLSS5JGFT5IY/events.json","paper":"https://pith.science/paper/BYXUGBTD"},"agent_actions":{"view_html":"https://pith.science/pith/BYXUGBTDKH34F7TLSS5JGFT5IY","download_json":"https://pith.science/pith/BYXUGBTDKH34F7TLSS5JGFT5IY.json","view_paper":"https://pith.science/paper/BYXUGBTD","resolve_alias":"https://pith.science/api/pith-number/resolve?arxiv=2306.14924&json=true","fetch_graph":"https://pith.science/api/pith-number/BYXUGBTDKH34F7TLSS5JGFT5IY/graph.json","fetch_events":"https://pith.science/api/pith-number/BYXUGBTDKH34F7TLSS5JGFT5IY/events.json","actions":{"anchor_timestamp":"https://pith.science/pith/BYXUGBTDKH34F7TLSS5JGFT5IY/action/timestamp_anchor","attest_storage":"https://pith.science/pith/BYXUGBTDKH34F7TLSS5JGFT5IY/action/storage_attestation","attest_author":"https://pith.science/pith/BYXUGBTDKH34F7TLSS5JGFT5IY/action/author_attestation","sign_citation":"https://pith.science/pith/BYXUGBTDKH34F7TLSS5JGFT5IY/action/citation_signature","submit_replication":"https://pith.science/pith/BYXUGBTDKH34F7TLSS5JGFT5IY/action/replication_record"}},"created_at":"2026-07-05T06:24:51.874264+00:00","updated_at":"2026-07-05T06:24:51.874264+00:00"}