{"record_type":"pith_number_record","schema_url":"https://pith.science/schemas/pith-number/v1.json","pith_number":"pith:2024:XIRYMIGMQ3YW4XYHMS3PUH3JMG","short_pith_number":"pith:XIRYMIGM","schema_version":"1.0","canonical_sha256":"ba238620cc86f16e5f0764b6fa1f6961978723157d4d5f5fe2dd79bfac927b48","source":{"kind":"arxiv","id":"2405.19519","version":2},"attestation_state":"computed","paper":{"title":"Two-Layer Retrieval-Augmented Generation Framework for Low-Resource Medical Question Answering Using Reddit Data: Proof-of-Concept Study","license":"http://creativecommons.org/licenses/by-nc-nd/4.0/","headline":"","cross_cats":["cs.AI"],"primary_cat":"cs.CL","authors_text":"Abeed Sarker, Anthony Spadaro, Danielle Mowery, Drew Walker, JaMor Hairston, Jeanmarie Perrone, Jeanne Powell, Jennifer Love, Matthew Reyna, Natalie Hernandez, Rachel Wightman, Rasheeta Chandler, Reza Sameni, Sahithi Lakamana, Sangmi Kim, Selen Bozkurt, Snigdha Peddireddy, Sudeshna Das, Swati Rajwal, Yao Ge, Yunyu Xiao, Yuting Guo","submitted_at":"2024-05-29T20:56:52Z","abstract_excerpt":"The increasing use of social media to share lived and living experiences of substance use presents a unique opportunity to obtain information on side effects, use patterns, and opinions on novel psychoactive substances. However, due to the large volume of data, obtaining useful insights through natural language processing technologies such as large language models is challenging. This paper aims to develop a retrieval-augmented generation (RAG) architecture for medical question answering pertaining to clinicians' queries on emerging issues associated with health-related topics, using user-gene"},"verification_status":{"content_addressed":true,"pith_receipt":true,"author_attested":false,"weak_author_claims":0,"strong_author_claims":0,"externally_anchored":false,"storage_verified":false,"citation_signatures":0,"replication_records":0,"graph_snapshot":true,"references_resolved":false,"formal_links_present":false},"canonical_record":{"source":{"id":"2405.19519","kind":"arxiv","version":2},"metadata":{"license":"http://creativecommons.org/licenses/by-nc-nd/4.0/","primary_cat":"cs.CL","submitted_at":"2024-05-29T20:56:52Z","cross_cats_sorted":["cs.AI"],"title_canon_sha256":"e23de255934252ead23b91ef5df79228dd6a983527effb981015a53cc3fee228","abstract_canon_sha256":"cecb482500cc372eac37d3b4c716f99b885d412b02a913d3d8cae754adb7ea1a"},"schema_version":"1.0"},"receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-07-05T09:57:43.096112Z","signature_b64":"NY5gjrb+0Gl0H3D3qeA1Sti+rsVNdTKshbazgySJnCC3FLHc7eRj66PrlTxudX+CW0dbdIb+So77AVd7rjKiBg==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"ba238620cc86f16e5f0764b6fa1f6961978723157d4d5f5fe2dd79bfac927b48","last_reissued_at":"2026-07-05T09:57:43.095570Z","signature_status":"signed_v1","first_computed_at":"2026-07-05T09:57:43.095570Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"graph_snapshot":{"paper":{"title":"Two-Layer Retrieval-Augmented Generation Framework for Low-Resource Medical Question Answering Using Reddit Data: Proof-of-Concept Study","license":"http://creativecommons.org/licenses/by-nc-nd/4.0/","headline":"","cross_cats":["cs.AI"],"primary_cat":"cs.CL","authors_text":"Abeed Sarker, Anthony Spadaro, Danielle Mowery, Drew Walker, JaMor Hairston, Jeanmarie Perrone, Jeanne Powell, Jennifer Love, Matthew Reyna, Natalie Hernandez, Rachel Wightman, Rasheeta Chandler, Reza Sameni, Sahithi Lakamana, Sangmi Kim, Selen Bozkurt, Snigdha Peddireddy, Sudeshna Das, Swati Rajwal, Yao Ge, Yunyu Xiao, Yuting Guo","submitted_at":"2024-05-29T20:56:52Z","abstract_excerpt":"The increasing use of social media to share lived and living experiences of substance use presents a unique opportunity to obtain information on side effects, use patterns, and opinions on novel psychoactive substances. However, due to the large volume of data, obtaining useful insights through natural language processing technologies such as large language models is challenging. This paper aims to develop a retrieval-augmented generation (RAG) architecture for medical question answering pertaining to clinicians' queries on emerging issues associated with health-related topics, using user-gene"},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2405.19519","kind":"arxiv","version":2},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2405.19519/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"aliases":[{"alias_kind":"arxiv","alias_value":"2405.19519","created_at":"2026-07-05T09:57:43.095629+00:00"},{"alias_kind":"arxiv_version","alias_value":"2405.19519v2","created_at":"2026-07-05T09:57:43.095629+00:00"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2405.19519","created_at":"2026-07-05T09:57:43.095629+00:00"},{"alias_kind":"pith_short_12","alias_value":"XIRYMIGMQ3YW","created_at":"2026-07-05T09:57:43.095629+00:00"},{"alias_kind":"pith_short_16","alias_value":"XIRYMIGMQ3YW4XYH","created_at":"2026-07-05T09:57:43.095629+00:00"},{"alias_kind":"pith_short_8","alias_value":"XIRYMIGM","created_at":"2026-07-05T09:57:43.095629+00:00"}],"events":[],"event_summary":{},"paper_claims":[],"inbound_citations":{"count":0,"internal_anchor_count":0,"sample":[]},"formal_canon":{"evidence_count":0,"sample":[],"anchors":[]},"links":{"html":"https://pith.science/pith/XIRYMIGMQ3YW4XYHMS3PUH3JMG","json":"https://pith.science/pith/XIRYMIGMQ3YW4XYHMS3PUH3JMG.json","graph_json":"https://pith.science/api/pith-number/XIRYMIGMQ3YW4XYHMS3PUH3JMG/graph.json","events_json":"https://pith.science/api/pith-number/XIRYMIGMQ3YW4XYHMS3PUH3JMG/events.json","paper":"https://pith.science/paper/XIRYMIGM"},"agent_actions":{"view_html":"https://pith.science/pith/XIRYMIGMQ3YW4XYHMS3PUH3JMG","download_json":"https://pith.science/pith/XIRYMIGMQ3YW4XYHMS3PUH3JMG.json","view_paper":"https://pith.science/paper/XIRYMIGM","resolve_alias":"https://pith.science/api/pith-number/resolve?arxiv=2405.19519&json=true","fetch_graph":"https://pith.science/api/pith-number/XIRYMIGMQ3YW4XYHMS3PUH3JMG/graph.json","fetch_events":"https://pith.science/api/pith-number/XIRYMIGMQ3YW4XYHMS3PUH3JMG/events.json","actions":{"anchor_timestamp":"https://pith.science/pith/XIRYMIGMQ3YW4XYHMS3PUH3JMG/action/timestamp_anchor","attest_storage":"https://pith.science/pith/XIRYMIGMQ3YW4XYHMS3PUH3JMG/action/storage_attestation","attest_author":"https://pith.science/pith/XIRYMIGMQ3YW4XYHMS3PUH3JMG/action/author_attestation","sign_citation":"https://pith.science/pith/XIRYMIGMQ3YW4XYHMS3PUH3JMG/action/citation_signature","submit_replication":"https://pith.science/pith/XIRYMIGMQ3YW4XYHMS3PUH3JMG/action/replication_record"}},"created_at":"2026-07-05T09:57:43.095629+00:00","updated_at":"2026-07-05T09:57:43.095629+00:00"}