{"record_type":"pith_number_record","schema_url":"https://pith.science/schemas/pith-number/v1.json","pith_number":"pith:2025:3KBGFHAOWJSBGIY7VLXW26GC4V","short_pith_number":"pith:3KBGFHAO","schema_version":"1.0","canonical_sha256":"da82629c0eb26413231faaef6d78c2e5526dae31d784b2e70ea728652c675938","source":{"kind":"arxiv","id":"2507.11842","version":1},"attestation_state":"computed","paper":{"title":"CosmoFlow: Scale-Aware Representation Learning for Cosmology with Flow Matching","license":"http://creativecommons.org/licenses/by/4.0/","headline":"","cross_cats":["cs.LG"],"primary_cat":"astro-ph.CO","authors_text":"Carolina Cuesta-Lazaro, Haewon Jeong, Sidharth Kannan, Tian Qiu","submitted_at":"2025-07-16T02:15:31Z","abstract_excerpt":"Generative machine learning models have been demonstrated to be able to learn low dimensional representations of data that preserve information required for downstream tasks. In this work, we demonstrate that flow matching based generative models can learn compact, semantically rich latent representations of field level cold dark matter (CDM) simulation data without supervision. Our model, CosmoFlow, learns representations 32x smaller than the raw field data, usable for field level reconstruction, synthetic data generation, and parameter inference. Our model also learns interpretable represent"},"verification_status":{"content_addressed":true,"pith_receipt":true,"author_attested":false,"weak_author_claims":0,"strong_author_claims":0,"externally_anchored":false,"storage_verified":false,"citation_signatures":0,"replication_records":0,"graph_snapshot":true,"references_resolved":false,"formal_links_present":false},"canonical_record":{"source":{"id":"2507.11842","kind":"arxiv","version":1},"metadata":{"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"astro-ph.CO","submitted_at":"2025-07-16T02:15:31Z","cross_cats_sorted":["cs.LG"],"title_canon_sha256":"b6153a4a26cf59a47c95dbb0766d0137f3fe45ca372825945e74e6886380d70e","abstract_canon_sha256":"43940e66d26545827d4c01deefcf74aecd4f02b37ff3d07469e5dd714905d03d"},"schema_version":"1.0"},"receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-07-05T11:37:51.882468Z","signature_b64":"i6QxB2mEAsDdY8cwWtn3Qynt3pvtwTBle5P5W3coy1In88XmDCP4gJ1J1NwOcJgPED0+BpBy2P1oADHzFiemBQ==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"da82629c0eb26413231faaef6d78c2e5526dae31d784b2e70ea728652c675938","last_reissued_at":"2026-07-05T11:37:51.881919Z","signature_status":"signed_v1","first_computed_at":"2026-07-05T11:37:51.881919Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"graph_snapshot":{"paper":{"title":"CosmoFlow: Scale-Aware Representation Learning for Cosmology with Flow Matching","license":"http://creativecommons.org/licenses/by/4.0/","headline":"","cross_cats":["cs.LG"],"primary_cat":"astro-ph.CO","authors_text":"Carolina Cuesta-Lazaro, Haewon Jeong, Sidharth Kannan, Tian Qiu","submitted_at":"2025-07-16T02:15:31Z","abstract_excerpt":"Generative machine learning models have been demonstrated to be able to learn low dimensional representations of data that preserve information required for downstream tasks. In this work, we demonstrate that flow matching based generative models can learn compact, semantically rich latent representations of field level cold dark matter (CDM) simulation data without supervision. Our model, CosmoFlow, learns representations 32x smaller than the raw field data, usable for field level reconstruction, synthetic data generation, and parameter inference. Our model also learns interpretable represent"},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2507.11842","kind":"arxiv","version":1},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2507.11842/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"aliases":[{"alias_kind":"arxiv","alias_value":"2507.11842","created_at":"2026-07-05T11:37:51.881980+00:00"},{"alias_kind":"arxiv_version","alias_value":"2507.11842v1","created_at":"2026-07-05T11:37:51.881980+00:00"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2507.11842","created_at":"2026-07-05T11:37:51.881980+00:00"},{"alias_kind":"pith_short_12","alias_value":"3KBGFHAOWJSB","created_at":"2026-07-05T11:37:51.881980+00:00"},{"alias_kind":"pith_short_16","alias_value":"3KBGFHAOWJSBGIY7","created_at":"2026-07-05T11:37:51.881980+00:00"},{"alias_kind":"pith_short_8","alias_value":"3KBGFHAO","created_at":"2026-07-05T11:37:51.881980+00:00"}],"events":[],"event_summary":{},"paper_claims":[],"inbound_citations":{"count":2,"internal_anchor_count":0,"sample":[{"citing_arxiv_id":"2606.10023","citing_title":"Learning the Universe: Posterior Reliability of Neural Generative Models in High-Dimensional Field-Level Inference of Cosmic Initial Conditions","ref_index":61,"is_internal_anchor":false},{"citing_arxiv_id":"2605.23114","citing_title":"Increasing the Precision of Surrogate Models for Weak Lensing Mass Maps with Flow Matching","ref_index":30,"is_internal_anchor":false}]},"formal_canon":{"evidence_count":0,"sample":[],"anchors":[]},"links":{"html":"https://pith.science/pith/3KBGFHAOWJSBGIY7VLXW26GC4V","json":"https://pith.science/pith/3KBGFHAOWJSBGIY7VLXW26GC4V.json","graph_json":"https://pith.science/api/pith-number/3KBGFHAOWJSBGIY7VLXW26GC4V/graph.json","events_json":"https://pith.science/api/pith-number/3KBGFHAOWJSBGIY7VLXW26GC4V/events.json","paper":"https://pith.science/paper/3KBGFHAO"},"agent_actions":{"view_html":"https://pith.science/pith/3KBGFHAOWJSBGIY7VLXW26GC4V","download_json":"https://pith.science/pith/3KBGFHAOWJSBGIY7VLXW26GC4V.json","view_paper":"https://pith.science/paper/3KBGFHAO","resolve_alias":"https://pith.science/api/pith-number/resolve?arxiv=2507.11842&json=true","fetch_graph":"https://pith.science/api/pith-number/3KBGFHAOWJSBGIY7VLXW26GC4V/graph.json","fetch_events":"https://pith.science/api/pith-number/3KBGFHAOWJSBGIY7VLXW26GC4V/events.json","actions":{"anchor_timestamp":"https://pith.science/pith/3KBGFHAOWJSBGIY7VLXW26GC4V/action/timestamp_anchor","attest_storage":"https://pith.science/pith/3KBGFHAOWJSBGIY7VLXW26GC4V/action/storage_attestation","attest_author":"https://pith.science/pith/3KBGFHAOWJSBGIY7VLXW26GC4V/action/author_attestation","sign_citation":"https://pith.science/pith/3KBGFHAOWJSBGIY7VLXW26GC4V/action/citation_signature","submit_replication":"https://pith.science/pith/3KBGFHAOWJSBGIY7VLXW26GC4V/action/replication_record"}},"created_at":"2026-07-05T11:37:51.881980+00:00","updated_at":"2026-07-05T11:37:51.881980+00:00"}