{"record_type":"pith_number_record","schema_url":"https://pith.science/schemas/pith-number/v1.json","pith_number":"pith:2023:CKDRSSTME37GN3YJDDY2UPYLEH","short_pith_number":"pith:CKDRSSTM","schema_version":"1.0","canonical_sha256":"1287194a6c26fe66ef0918f1aa3f0b21e94b54d48f4240f65c2fd97f8acc095e","source":{"kind":"arxiv","id":"2307.02313","version":2},"attestation_state":"computed","paper":{"title":"Utilizing ChatGPT Generated Data to Retrieve Depression Symptoms from Social Media","license":"http://creativecommons.org/licenses/by-nc-nd/4.0/","headline":"","cross_cats":[],"primary_cat":"cs.CL","authors_text":"Ana-Maria Bucur","submitted_at":"2023-07-05T14:15:15Z","abstract_excerpt":"In this work, we present the contribution of the BLUE team in the eRisk Lab task on searching for symptoms of depression. The task consists of retrieving and ranking Reddit social media sentences that convey symptoms of depression from the BDI-II questionnaire. Given that synthetic data provided by LLMs have been proven to be a reliable method for augmenting data and fine-tuning downstream models, we chose to generate synthetic data using ChatGPT for each of the symptoms of the BDI-II questionnaire. We designed a prompt such that the generated data contains more richness and semantic diversity"},"verification_status":{"content_addressed":true,"pith_receipt":true,"author_attested":false,"weak_author_claims":0,"strong_author_claims":0,"externally_anchored":false,"storage_verified":false,"citation_signatures":0,"replication_records":0,"graph_snapshot":true,"references_resolved":false,"formal_links_present":false},"canonical_record":{"source":{"id":"2307.02313","kind":"arxiv","version":2},"metadata":{"license":"http://creativecommons.org/licenses/by-nc-nd/4.0/","primary_cat":"cs.CL","submitted_at":"2023-07-05T14:15:15Z","cross_cats_sorted":[],"title_canon_sha256":"589e280a5c0ccddcd2c1e49ee791d43d5d2ff571b4ab00873a69f1bdc75b554b","abstract_canon_sha256":"f6f5543d928491074a34788ce0fdd4c36e11603b8b95478e9ec0815eb7c68b39"},"schema_version":"1.0"},"receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-07-05T06:28:20.492026Z","signature_b64":"aoOOC/P67f5as3pIeHP6uEekm7VfU6QlCZWVDjffZMapjHFdafGeE4h2Nzg7fVkLnM/HA07bGDxIR4r6tEhJDQ==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"1287194a6c26fe66ef0918f1aa3f0b21e94b54d48f4240f65c2fd97f8acc095e","last_reissued_at":"2026-07-05T06:28:20.491593Z","signature_status":"signed_v1","first_computed_at":"2026-07-05T06:28:20.491593Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"graph_snapshot":{"paper":{"title":"Utilizing ChatGPT Generated Data to Retrieve Depression Symptoms from Social Media","license":"http://creativecommons.org/licenses/by-nc-nd/4.0/","headline":"","cross_cats":[],"primary_cat":"cs.CL","authors_text":"Ana-Maria Bucur","submitted_at":"2023-07-05T14:15:15Z","abstract_excerpt":"In this work, we present the contribution of the BLUE team in the eRisk Lab task on searching for symptoms of depression. The task consists of retrieving and ranking Reddit social media sentences that convey symptoms of depression from the BDI-II questionnaire. Given that synthetic data provided by LLMs have been proven to be a reliable method for augmenting data and fine-tuning downstream models, we chose to generate synthetic data using ChatGPT for each of the symptoms of the BDI-II questionnaire. We designed a prompt such that the generated data contains more richness and semantic diversity"},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2307.02313","kind":"arxiv","version":2},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2307.02313/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"aliases":[{"alias_kind":"arxiv","alias_value":"2307.02313","created_at":"2026-07-05T06:28:20.491651+00:00"},{"alias_kind":"arxiv_version","alias_value":"2307.02313v2","created_at":"2026-07-05T06:28:20.491651+00:00"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2307.02313","created_at":"2026-07-05T06:28:20.491651+00:00"},{"alias_kind":"pith_short_12","alias_value":"CKDRSSTME37G","created_at":"2026-07-05T06:28:20.491651+00:00"},{"alias_kind":"pith_short_16","alias_value":"CKDRSSTME37GN3YJ","created_at":"2026-07-05T06:28:20.491651+00:00"},{"alias_kind":"pith_short_8","alias_value":"CKDRSSTM","created_at":"2026-07-05T06:28:20.491651+00:00"}],"events":[],"event_summary":{},"paper_claims":[],"inbound_citations":{"count":0,"internal_anchor_count":0,"sample":[]},"formal_canon":{"evidence_count":0,"sample":[],"anchors":[]},"links":{"html":"https://pith.science/pith/CKDRSSTME37GN3YJDDY2UPYLEH","json":"https://pith.science/pith/CKDRSSTME37GN3YJDDY2UPYLEH.json","graph_json":"https://pith.science/api/pith-number/CKDRSSTME37GN3YJDDY2UPYLEH/graph.json","events_json":"https://pith.science/api/pith-number/CKDRSSTME37GN3YJDDY2UPYLEH/events.json","paper":"https://pith.science/paper/CKDRSSTM"},"agent_actions":{"view_html":"https://pith.science/pith/CKDRSSTME37GN3YJDDY2UPYLEH","download_json":"https://pith.science/pith/CKDRSSTME37GN3YJDDY2UPYLEH.json","view_paper":"https://pith.science/paper/CKDRSSTM","resolve_alias":"https://pith.science/api/pith-number/resolve?arxiv=2307.02313&json=true","fetch_graph":"https://pith.science/api/pith-number/CKDRSSTME37GN3YJDDY2UPYLEH/graph.json","fetch_events":"https://pith.science/api/pith-number/CKDRSSTME37GN3YJDDY2UPYLEH/events.json","actions":{"anchor_timestamp":"https://pith.science/pith/CKDRSSTME37GN3YJDDY2UPYLEH/action/timestamp_anchor","attest_storage":"https://pith.science/pith/CKDRSSTME37GN3YJDDY2UPYLEH/action/storage_attestation","attest_author":"https://pith.science/pith/CKDRSSTME37GN3YJDDY2UPYLEH/action/author_attestation","sign_citation":"https://pith.science/pith/CKDRSSTME37GN3YJDDY2UPYLEH/action/citation_signature","submit_replication":"https://pith.science/pith/CKDRSSTME37GN3YJDDY2UPYLEH/action/replication_record"}},"created_at":"2026-07-05T06:28:20.491651+00:00","updated_at":"2026-07-05T06:28:20.491651+00:00"}