{"record_type":"pith_number_record","schema_url":"https://pith.science/schemas/pith-number/v1.json","pith_number":"pith:2024:SYF4FQ4QWE4VAM564BEMI5WJ7S","short_pith_number":"pith:SYF4FQ4Q","schema_version":"1.0","canonical_sha256":"960bc2c390b1395033bee048c476c9fcae50f84f3ad74c3310d4df3ec740e2d3","source":{"kind":"arxiv","id":"2406.03776","version":2},"attestation_state":"computed","paper":{"title":"XL-HeadTags: Leveraging Multimodal Retrieval Augmentation for the Multilingual Generation of News Headlines and Tags","license":"http://creativecommons.org/licenses/by/4.0/","headline":"","cross_cats":["cs.AI","cs.CV","cs.IR"],"primary_cat":"cs.CL","authors_text":"Abu Ubaida Akash, Faisal Tareque Shohan, Mir Tafseer Nayeem, Samsul Islam, Shafiq Joty","submitted_at":"2024-06-06T06:40:19Z","abstract_excerpt":"Millions of news articles published online daily can overwhelm readers. Headlines and entity (topic) tags are essential for guiding readers to decide if the content is worth their time. While headline generation has been extensively studied, tag generation remains largely unexplored, yet it offers readers better access to topics of interest. The need for conciseness in capturing readers' attention necessitates improved content selection strategies for identifying salient and relevant segments within lengthy articles, thereby guiding language models effectively. To address this, we propose to l"},"verification_status":{"content_addressed":true,"pith_receipt":true,"author_attested":false,"weak_author_claims":0,"strong_author_claims":0,"externally_anchored":false,"storage_verified":false,"citation_signatures":0,"replication_records":0,"graph_snapshot":true,"references_resolved":false,"formal_links_present":false},"canonical_record":{"source":{"id":"2406.03776","kind":"arxiv","version":2},"metadata":{"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.CL","submitted_at":"2024-06-06T06:40:19Z","cross_cats_sorted":["cs.AI","cs.CV","cs.IR"],"title_canon_sha256":"d1321cb22e9dfebf6bb7b159f2f3cbed5c0e7522450ce752bec60b95f33bf69d","abstract_canon_sha256":"e561055754fef7a3e313a502a28488616b5d568b0006c58e4251d26ed4f07ec8"},"schema_version":"1.0"},"receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-07-05T08:28:42.382385Z","signature_b64":"ggW/Qc9OYy/kDUYOcb6cGvGtaxYs+kU+hHqjzX+wTk8cf/DcTxTCUAwz5Gal+SFewlQJeSpNND1wqmtFMjU9Dg==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"960bc2c390b1395033bee048c476c9fcae50f84f3ad74c3310d4df3ec740e2d3","last_reissued_at":"2026-07-05T08:28:42.381938Z","signature_status":"signed_v1","first_computed_at":"2026-07-05T08:28:42.381938Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"graph_snapshot":{"paper":{"title":"XL-HeadTags: Leveraging Multimodal Retrieval Augmentation for the Multilingual Generation of News Headlines and Tags","license":"http://creativecommons.org/licenses/by/4.0/","headline":"","cross_cats":["cs.AI","cs.CV","cs.IR"],"primary_cat":"cs.CL","authors_text":"Abu Ubaida Akash, Faisal Tareque Shohan, Mir Tafseer Nayeem, Samsul Islam, Shafiq Joty","submitted_at":"2024-06-06T06:40:19Z","abstract_excerpt":"Millions of news articles published online daily can overwhelm readers. Headlines and entity (topic) tags are essential for guiding readers to decide if the content is worth their time. While headline generation has been extensively studied, tag generation remains largely unexplored, yet it offers readers better access to topics of interest. The need for conciseness in capturing readers' attention necessitates improved content selection strategies for identifying salient and relevant segments within lengthy articles, thereby guiding language models effectively. To address this, we propose to l"},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2406.03776","kind":"arxiv","version":2},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2406.03776/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"aliases":[{"alias_kind":"arxiv","alias_value":"2406.03776","created_at":"2026-07-05T08:28:42.381997+00:00"},{"alias_kind":"arxiv_version","alias_value":"2406.03776v2","created_at":"2026-07-05T08:28:42.381997+00:00"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2406.03776","created_at":"2026-07-05T08:28:42.381997+00:00"},{"alias_kind":"pith_short_12","alias_value":"SYF4FQ4QWE4V","created_at":"2026-07-05T08:28:42.381997+00:00"},{"alias_kind":"pith_short_16","alias_value":"SYF4FQ4QWE4VAM56","created_at":"2026-07-05T08:28:42.381997+00:00"},{"alias_kind":"pith_short_8","alias_value":"SYF4FQ4Q","created_at":"2026-07-05T08:28:42.381997+00:00"}],"events":[],"event_summary":{},"paper_claims":[],"inbound_citations":{"count":1,"internal_anchor_count":0,"sample":[{"citing_arxiv_id":"2604.09552","citing_title":"MCERF: Advancing Multimodal LLM Evaluation of Engineering Documentation with Enhanced Retrieval","ref_index":32,"is_internal_anchor":false}]},"formal_canon":{"evidence_count":0,"sample":[],"anchors":[]},"links":{"html":"https://pith.science/pith/SYF4FQ4QWE4VAM564BEMI5WJ7S","json":"https://pith.science/pith/SYF4FQ4QWE4VAM564BEMI5WJ7S.json","graph_json":"https://pith.science/api/pith-number/SYF4FQ4QWE4VAM564BEMI5WJ7S/graph.json","events_json":"https://pith.science/api/pith-number/SYF4FQ4QWE4VAM564BEMI5WJ7S/events.json","paper":"https://pith.science/paper/SYF4FQ4Q"},"agent_actions":{"view_html":"https://pith.science/pith/SYF4FQ4QWE4VAM564BEMI5WJ7S","download_json":"https://pith.science/pith/SYF4FQ4QWE4VAM564BEMI5WJ7S.json","view_paper":"https://pith.science/paper/SYF4FQ4Q","resolve_alias":"https://pith.science/api/pith-number/resolve?arxiv=2406.03776&json=true","fetch_graph":"https://pith.science/api/pith-number/SYF4FQ4QWE4VAM564BEMI5WJ7S/graph.json","fetch_events":"https://pith.science/api/pith-number/SYF4FQ4QWE4VAM564BEMI5WJ7S/events.json","actions":{"anchor_timestamp":"https://pith.science/pith/SYF4FQ4QWE4VAM564BEMI5WJ7S/action/timestamp_anchor","attest_storage":"https://pith.science/pith/SYF4FQ4QWE4VAM564BEMI5WJ7S/action/storage_attestation","attest_author":"https://pith.science/pith/SYF4FQ4QWE4VAM564BEMI5WJ7S/action/author_attestation","sign_citation":"https://pith.science/pith/SYF4FQ4QWE4VAM564BEMI5WJ7S/action/citation_signature","submit_replication":"https://pith.science/pith/SYF4FQ4QWE4VAM564BEMI5WJ7S/action/replication_record"}},"created_at":"2026-07-05T08:28:42.381997+00:00","updated_at":"2026-07-05T08:28:42.381997+00:00"}