{"record_type":"pith_number_record","schema_url":"https://pith.science/schemas/pith-number/v1.json","pith_number":"pith:2025:GQHWLMG2DLQ2VGPUF7XVWJUQP5","short_pith_number":"pith:GQHWLMG2","schema_version":"1.0","canonical_sha256":"340f65b0da1ae1aa99f42fef5b26907f6f8f2758070b4431015f1550253b1361","source":{"kind":"arxiv","id":"2501.02167","version":1},"attestation_state":"computed","paper":{"title":"Generating Multimodal Images with GAN: Integrating Text, Image, and Style","license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","headline":"","cross_cats":[],"primary_cat":"cs.CV","authors_text":"Ao Xiang, Chaoyi Tan, Kowei Shih, Wenqing Zhang, Xinshi Li, Zhen Qi","submitted_at":"2025-01-04T02:51:28Z","abstract_excerpt":"In the field of computer vision, multimodal image generation has become a research hotspot, especially the task of integrating text, image, and style. In this study, we propose a multimodal image generation method based on Generative Adversarial Networks (GAN), capable of effectively combining text descriptions, reference images, and style information to generate images that meet multimodal requirements. This method involves the design of a text encoder, an image feature extractor, and a style integration module, ensuring that the generated images maintain high quality in terms of visual conte"},"verification_status":{"content_addressed":true,"pith_receipt":true,"author_attested":false,"weak_author_claims":0,"strong_author_claims":0,"externally_anchored":false,"storage_verified":false,"citation_signatures":0,"replication_records":0,"graph_snapshot":true,"references_resolved":false,"formal_links_present":false},"canonical_record":{"source":{"id":"2501.02167","kind":"arxiv","version":1},"metadata":{"license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","primary_cat":"cs.CV","submitted_at":"2025-01-04T02:51:28Z","cross_cats_sorted":[],"title_canon_sha256":"703f043edd9aef8aff30a9800759841486699cc06dae1eb5c88d5e2efef2e8d7","abstract_canon_sha256":"8e72d5ff6752d013ef61a854150f7a7f796c21103c127a16708ad8975daede3b"},"schema_version":"1.0"},"receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-07-05T09:57:03.241879Z","signature_b64":"rMc1Tt46X5rNvpSiZnj2dOpLYnWI2Su58AQ8nDM9fEFkCyuGCe6zaGlWakV1/X723lY/cjLUtoTKyf8lj4rLDA==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"340f65b0da1ae1aa99f42fef5b26907f6f8f2758070b4431015f1550253b1361","last_reissued_at":"2026-07-05T09:57:03.241400Z","signature_status":"signed_v1","first_computed_at":"2026-07-05T09:57:03.241400Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"graph_snapshot":{"paper":{"title":"Generating Multimodal Images with GAN: Integrating Text, Image, and Style","license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","headline":"","cross_cats":[],"primary_cat":"cs.CV","authors_text":"Ao Xiang, Chaoyi Tan, Kowei Shih, Wenqing Zhang, Xinshi Li, Zhen Qi","submitted_at":"2025-01-04T02:51:28Z","abstract_excerpt":"In the field of computer vision, multimodal image generation has become a research hotspot, especially the task of integrating text, image, and style. In this study, we propose a multimodal image generation method based on Generative Adversarial Networks (GAN), capable of effectively combining text descriptions, reference images, and style information to generate images that meet multimodal requirements. This method involves the design of a text encoder, an image feature extractor, and a style integration module, ensuring that the generated images maintain high quality in terms of visual conte"},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2501.02167","kind":"arxiv","version":1},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2501.02167/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"aliases":[{"alias_kind":"arxiv","alias_value":"2501.02167","created_at":"2026-07-05T09:57:03.241456+00:00"},{"alias_kind":"arxiv_version","alias_value":"2501.02167v1","created_at":"2026-07-05T09:57:03.241456+00:00"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2501.02167","created_at":"2026-07-05T09:57:03.241456+00:00"},{"alias_kind":"pith_short_12","alias_value":"GQHWLMG2DLQ2","created_at":"2026-07-05T09:57:03.241456+00:00"},{"alias_kind":"pith_short_16","alias_value":"GQHWLMG2DLQ2VGPU","created_at":"2026-07-05T09:57:03.241456+00:00"},{"alias_kind":"pith_short_8","alias_value":"GQHWLMG2","created_at":"2026-07-05T09:57:03.241456+00:00"}],"events":[],"event_summary":{},"paper_claims":[],"inbound_citations":{"count":1,"internal_anchor_count":1,"sample":[{"citing_arxiv_id":"2507.16672","citing_title":"Meta-Learning for Cold-Start Personalization in Prompt-Tuned LLMs","ref_index":34,"is_internal_anchor":true}]},"formal_canon":{"evidence_count":0,"sample":[],"anchors":[]},"links":{"html":"https://pith.science/pith/GQHWLMG2DLQ2VGPUF7XVWJUQP5","json":"https://pith.science/pith/GQHWLMG2DLQ2VGPUF7XVWJUQP5.json","graph_json":"https://pith.science/api/pith-number/GQHWLMG2DLQ2VGPUF7XVWJUQP5/graph.json","events_json":"https://pith.science/api/pith-number/GQHWLMG2DLQ2VGPUF7XVWJUQP5/events.json","paper":"https://pith.science/paper/GQHWLMG2"},"agent_actions":{"view_html":"https://pith.science/pith/GQHWLMG2DLQ2VGPUF7XVWJUQP5","download_json":"https://pith.science/pith/GQHWLMG2DLQ2VGPUF7XVWJUQP5.json","view_paper":"https://pith.science/paper/GQHWLMG2","resolve_alias":"https://pith.science/api/pith-number/resolve?arxiv=2501.02167&json=true","fetch_graph":"https://pith.science/api/pith-number/GQHWLMG2DLQ2VGPUF7XVWJUQP5/graph.json","fetch_events":"https://pith.science/api/pith-number/GQHWLMG2DLQ2VGPUF7XVWJUQP5/events.json","actions":{"anchor_timestamp":"https://pith.science/pith/GQHWLMG2DLQ2VGPUF7XVWJUQP5/action/timestamp_anchor","attest_storage":"https://pith.science/pith/GQHWLMG2DLQ2VGPUF7XVWJUQP5/action/storage_attestation","attest_author":"https://pith.science/pith/GQHWLMG2DLQ2VGPUF7XVWJUQP5/action/author_attestation","sign_citation":"https://pith.science/pith/GQHWLMG2DLQ2VGPUF7XVWJUQP5/action/citation_signature","submit_replication":"https://pith.science/pith/GQHWLMG2DLQ2VGPUF7XVWJUQP5/action/replication_record"}},"created_at":"2026-07-05T09:57:03.241456+00:00","updated_at":"2026-07-05T09:57:03.241456+00:00"}