{"record_type":"pith_number_record","schema_url":"https://pith.science/schemas/pith-number/v1.json","pith_number":"pith:2025:L7QRVYYKN3FIS66Y7QSNUD647Z","short_pith_number":"pith:L7QRVYYK","schema_version":"1.0","canonical_sha256":"5fe11ae30a6eca897bd8fc24da0fdcfe475e6b72367a872135374a2ab6136cba","source":{"kind":"arxiv","id":"2504.08961","version":2},"attestation_state":"computed","paper":{"title":"A Fully Automated Pipeline for Conversational Discourse Annotation: Tree Scheme Generation and Labeling with Large Language Models","license":"http://creativecommons.org/licenses/by/4.0/","headline":"","cross_cats":[],"primary_cat":"cs.CL","authors_text":"Ekaterina Kochmar, Kseniia Petukhova","submitted_at":"2025-04-11T20:36:57Z","abstract_excerpt":"Recent advances in Large Language Models (LLMs) have shown promise in automating discourse annotation for conversations. While manually designing tree annotation schemes significantly improves annotation quality for humans and models, their creation remains time-consuming and requires expert knowledge. We propose a fully automated pipeline that uses LLMs to construct such schemes and perform annotation. We evaluate our approach on speech functions (SFs) and the Switchboard-DAMSL (SWBD-DAMSL) taxonomies. Our experiments compare various design choices, and we show that frequency-guided decision "},"verification_status":{"content_addressed":true,"pith_receipt":true,"author_attested":false,"weak_author_claims":0,"strong_author_claims":0,"externally_anchored":false,"storage_verified":false,"citation_signatures":0,"replication_records":0,"graph_snapshot":true,"references_resolved":false,"formal_links_present":false},"canonical_record":{"source":{"id":"2504.08961","kind":"arxiv","version":2},"metadata":{"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.CL","submitted_at":"2025-04-11T20:36:57Z","cross_cats_sorted":[],"title_canon_sha256":"58e37f0d000848f83d9ca1d97fd73a3238e76fbd0fcb3d132808dab30eece732","abstract_canon_sha256":"537090a67c0f40251d035e15911ceb7a6bf86c799a364e6e3e64897ddfe93244"},"schema_version":"1.0"},"receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-07-05T11:14:58.108082Z","signature_b64":"IAiYky22rhhfGLgZGOp0f46Saky3Irlia7mgPzlHDSZJjy2tokhxTHh5Dho8ZQ5htTNOGiaiJWDIdv/WLGbpBg==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"5fe11ae30a6eca897bd8fc24da0fdcfe475e6b72367a872135374a2ab6136cba","last_reissued_at":"2026-07-05T11:14:58.107540Z","signature_status":"signed_v1","first_computed_at":"2026-07-05T11:14:58.107540Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"graph_snapshot":{"paper":{"title":"A Fully Automated Pipeline for Conversational Discourse Annotation: Tree Scheme Generation and Labeling with Large Language Models","license":"http://creativecommons.org/licenses/by/4.0/","headline":"","cross_cats":[],"primary_cat":"cs.CL","authors_text":"Ekaterina Kochmar, Kseniia Petukhova","submitted_at":"2025-04-11T20:36:57Z","abstract_excerpt":"Recent advances in Large Language Models (LLMs) have shown promise in automating discourse annotation for conversations. While manually designing tree annotation schemes significantly improves annotation quality for humans and models, their creation remains time-consuming and requires expert knowledge. We propose a fully automated pipeline that uses LLMs to construct such schemes and perform annotation. We evaluate our approach on speech functions (SFs) and the Switchboard-DAMSL (SWBD-DAMSL) taxonomies. Our experiments compare various design choices, and we show that frequency-guided decision "},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2504.08961","kind":"arxiv","version":2},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2504.08961/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"aliases":[{"alias_kind":"arxiv","alias_value":"2504.08961","created_at":"2026-07-05T11:14:58.107609+00:00"},{"alias_kind":"arxiv_version","alias_value":"2504.08961v2","created_at":"2026-07-05T11:14:58.107609+00:00"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2504.08961","created_at":"2026-07-05T11:14:58.107609+00:00"},{"alias_kind":"pith_short_12","alias_value":"L7QRVYYKN3FI","created_at":"2026-07-05T11:14:58.107609+00:00"},{"alias_kind":"pith_short_16","alias_value":"L7QRVYYKN3FIS66Y","created_at":"2026-07-05T11:14:58.107609+00:00"},{"alias_kind":"pith_short_8","alias_value":"L7QRVYYK","created_at":"2026-07-05T11:14:58.107609+00:00"}],"events":[],"event_summary":{},"paper_claims":[],"inbound_citations":{"count":1,"internal_anchor_count":1,"sample":[{"citing_arxiv_id":"2506.07626","citing_title":"Intent Matters: Enhancing AI Tutoring with Fine-Grained Pedagogical Intent Annotation","ref_index":25,"is_internal_anchor":true}]},"formal_canon":{"evidence_count":0,"sample":[],"anchors":[]},"links":{"html":"https://pith.science/pith/L7QRVYYKN3FIS66Y7QSNUD647Z","json":"https://pith.science/pith/L7QRVYYKN3FIS66Y7QSNUD647Z.json","graph_json":"https://pith.science/api/pith-number/L7QRVYYKN3FIS66Y7QSNUD647Z/graph.json","events_json":"https://pith.science/api/pith-number/L7QRVYYKN3FIS66Y7QSNUD647Z/events.json","paper":"https://pith.science/paper/L7QRVYYK"},"agent_actions":{"view_html":"https://pith.science/pith/L7QRVYYKN3FIS66Y7QSNUD647Z","download_json":"https://pith.science/pith/L7QRVYYKN3FIS66Y7QSNUD647Z.json","view_paper":"https://pith.science/paper/L7QRVYYK","resolve_alias":"https://pith.science/api/pith-number/resolve?arxiv=2504.08961&json=true","fetch_graph":"https://pith.science/api/pith-number/L7QRVYYKN3FIS66Y7QSNUD647Z/graph.json","fetch_events":"https://pith.science/api/pith-number/L7QRVYYKN3FIS66Y7QSNUD647Z/events.json","actions":{"anchor_timestamp":"https://pith.science/pith/L7QRVYYKN3FIS66Y7QSNUD647Z/action/timestamp_anchor","attest_storage":"https://pith.science/pith/L7QRVYYKN3FIS66Y7QSNUD647Z/action/storage_attestation","attest_author":"https://pith.science/pith/L7QRVYYKN3FIS66Y7QSNUD647Z/action/author_attestation","sign_citation":"https://pith.science/pith/L7QRVYYKN3FIS66Y7QSNUD647Z/action/citation_signature","submit_replication":"https://pith.science/pith/L7QRVYYKN3FIS66Y7QSNUD647Z/action/replication_record"}},"created_at":"2026-07-05T11:14:58.107609+00:00","updated_at":"2026-07-05T11:14:58.107609+00:00"}