{"record_type":"pith_number_record","schema_url":"https://pith.science/schemas/pith-number/v1.json","pith_number":"pith:2025:B5RFF6OXWFX4OUS62VUJ472OMJ","short_pith_number":"pith:B5RFF6OX","schema_version":"1.0","canonical_sha256":"0f6252f9d7b16fc7525ed5689e7f4e6276d7b7b6514fd49d09c6f0ca4cc4f503","source":{"kind":"arxiv","id":"2501.12356","version":1},"attestation_state":"computed","paper":{"title":"Vision-Language Models for Automated Chest X-ray Interpretation: Leveraging ViT and GPT-2","license":"http://creativecommons.org/licenses/by-nc-nd/4.0/","headline":"","cross_cats":[],"primary_cat":"cs.CV","authors_text":"Md. Rakibul Islam, Md. Zahid Hossain, Most. Sharmin Sultana Samu, Mustofa Ahmed","submitted_at":"2025-01-21T18:36:18Z","abstract_excerpt":"Radiology plays a pivotal role in modern medicine due to its non-invasive diagnostic capabilities. However, the manual generation of unstructured medical reports is time consuming and prone to errors. It creates a significant bottleneck in clinical workflows. Despite advancements in AI-generated radiology reports, challenges remain in achieving detailed and accurate report generation. In this study we have evaluated different combinations of multimodal models that integrate Computer Vision and Natural Language Processing to generate comprehensive radiology reports. We employed a pretrained Vis"},"verification_status":{"content_addressed":true,"pith_receipt":true,"author_attested":false,"weak_author_claims":0,"strong_author_claims":0,"externally_anchored":false,"storage_verified":false,"citation_signatures":0,"replication_records":0,"graph_snapshot":true,"references_resolved":false,"formal_links_present":false},"canonical_record":{"source":{"id":"2501.12356","kind":"arxiv","version":1},"metadata":{"license":"http://creativecommons.org/licenses/by-nc-nd/4.0/","primary_cat":"cs.CV","submitted_at":"2025-01-21T18:36:18Z","cross_cats_sorted":[],"title_canon_sha256":"be88807074f39a10d18c33a9122b00d7dada7507d4e713d2c7f80140c2399882","abstract_canon_sha256":"5a50e7efb5bd646ef50c03f2814541c45b06702a59e50d8c4d14702d3ac52ded"},"schema_version":"1.0"},"receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-07-05T10:03:35.051518Z","signature_b64":"ZEMKnGoDXdq5m7JIT/5I90DkKmKpuvGDxsOlI1KXQDBw7lGKYdpdiBy253HNan3Vy1lgAXUSYGVqf+YqtMgbDQ==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"0f6252f9d7b16fc7525ed5689e7f4e6276d7b7b6514fd49d09c6f0ca4cc4f503","last_reissued_at":"2026-07-05T10:03:35.051100Z","signature_status":"signed_v1","first_computed_at":"2026-07-05T10:03:35.051100Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"graph_snapshot":{"paper":{"title":"Vision-Language Models for Automated Chest X-ray Interpretation: Leveraging ViT and GPT-2","license":"http://creativecommons.org/licenses/by-nc-nd/4.0/","headline":"","cross_cats":[],"primary_cat":"cs.CV","authors_text":"Md. Rakibul Islam, Md. Zahid Hossain, Most. Sharmin Sultana Samu, Mustofa Ahmed","submitted_at":"2025-01-21T18:36:18Z","abstract_excerpt":"Radiology plays a pivotal role in modern medicine due to its non-invasive diagnostic capabilities. However, the manual generation of unstructured medical reports is time consuming and prone to errors. It creates a significant bottleneck in clinical workflows. Despite advancements in AI-generated radiology reports, challenges remain in achieving detailed and accurate report generation. In this study we have evaluated different combinations of multimodal models that integrate Computer Vision and Natural Language Processing to generate comprehensive radiology reports. We employed a pretrained Vis"},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2501.12356","kind":"arxiv","version":1},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2501.12356/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"aliases":[{"alias_kind":"arxiv","alias_value":"2501.12356","created_at":"2026-07-05T10:03:35.051156+00:00"},{"alias_kind":"arxiv_version","alias_value":"2501.12356v1","created_at":"2026-07-05T10:03:35.051156+00:00"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2501.12356","created_at":"2026-07-05T10:03:35.051156+00:00"},{"alias_kind":"pith_short_12","alias_value":"B5RFF6OXWFX4","created_at":"2026-07-05T10:03:35.051156+00:00"},{"alias_kind":"pith_short_16","alias_value":"B5RFF6OXWFX4OUS6","created_at":"2026-07-05T10:03:35.051156+00:00"},{"alias_kind":"pith_short_8","alias_value":"B5RFF6OX","created_at":"2026-07-05T10:03:35.051156+00:00"}],"events":[],"event_summary":{},"paper_claims":[],"inbound_citations":{"count":1,"internal_anchor_count":1,"sample":[{"citing_arxiv_id":"2505.21715","citing_title":"Privacy-Preserving Chest X-ray Report Generation via Multimodal Federated Learning with ViT and GPT-2","ref_index":51,"is_internal_anchor":true}]},"formal_canon":{"evidence_count":0,"sample":[],"anchors":[]},"links":{"html":"https://pith.science/pith/B5RFF6OXWFX4OUS62VUJ472OMJ","json":"https://pith.science/pith/B5RFF6OXWFX4OUS62VUJ472OMJ.json","graph_json":"https://pith.science/api/pith-number/B5RFF6OXWFX4OUS62VUJ472OMJ/graph.json","events_json":"https://pith.science/api/pith-number/B5RFF6OXWFX4OUS62VUJ472OMJ/events.json","paper":"https://pith.science/paper/B5RFF6OX"},"agent_actions":{"view_html":"https://pith.science/pith/B5RFF6OXWFX4OUS62VUJ472OMJ","download_json":"https://pith.science/pith/B5RFF6OXWFX4OUS62VUJ472OMJ.json","view_paper":"https://pith.science/paper/B5RFF6OX","resolve_alias":"https://pith.science/api/pith-number/resolve?arxiv=2501.12356&json=true","fetch_graph":"https://pith.science/api/pith-number/B5RFF6OXWFX4OUS62VUJ472OMJ/graph.json","fetch_events":"https://pith.science/api/pith-number/B5RFF6OXWFX4OUS62VUJ472OMJ/events.json","actions":{"anchor_timestamp":"https://pith.science/pith/B5RFF6OXWFX4OUS62VUJ472OMJ/action/timestamp_anchor","attest_storage":"https://pith.science/pith/B5RFF6OXWFX4OUS62VUJ472OMJ/action/storage_attestation","attest_author":"https://pith.science/pith/B5RFF6OXWFX4OUS62VUJ472OMJ/action/author_attestation","sign_citation":"https://pith.science/pith/B5RFF6OXWFX4OUS62VUJ472OMJ/action/citation_signature","submit_replication":"https://pith.science/pith/B5RFF6OXWFX4OUS62VUJ472OMJ/action/replication_record"}},"created_at":"2026-07-05T10:03:35.051156+00:00","updated_at":"2026-07-05T10:03:35.051156+00:00"}