{"state_type":"pith_open_graph_state","state_version":"1.0","pith_number":"pith:2026:VUPYIABJ5G4KIF3EPAQR2LLB3G","merge_version":"pith-open-graph-merge-v1","event_count":2,"valid_event_count":2,"invalid_event_count":0,"equivocation_count":0,"current":{"canonical_record":{"metadata":{"abstract_canon_sha256":"ae0a5272f863ec00a0916dff34b121a5f9ac9651c5a3443aafbdf9c41fae95d6","cross_cats_sorted":[],"license":"http://creativecommons.org/licenses/by-nc-sa/4.0/","primary_cat":"cs.CL","submitted_at":"2026-05-26T17:47:14Z","title_canon_sha256":"929417e6842cc3f218205a362b695386c6a278f2f6060a2d63f63bfe5bfdf524"},"schema_version":"1.0","source":{"id":"2605.27345","kind":"arxiv","version":1}},"source_aliases":[{"alias_kind":"arxiv","alias_value":"2605.27345","created_at":"2026-05-27T02:06:19Z"},{"alias_kind":"arxiv_version","alias_value":"2605.27345v1","created_at":"2026-05-27T02:06:19Z"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2605.27345","created_at":"2026-05-27T02:06:19Z"},{"alias_kind":"pith_short_12","alias_value":"VUPYIABJ5G4K","created_at":"2026-05-27T02:06:19Z"},{"alias_kind":"pith_short_16","alias_value":"VUPYIABJ5G4KIF3E","created_at":"2026-05-27T02:06:19Z"},{"alias_kind":"pith_short_8","alias_value":"VUPYIABJ","created_at":"2026-05-27T02:06:19Z"}],"graph_snapshots":[{"event_id":"sha256:6fa0f8e40b3742fedd3e44f24ee15e3ef9f339a3bc185c27cf14c4e495b3829b","target":"graph","created_at":"2026-05-27T02:06:19Z","signer":{"key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signer_id":"pith.science","signer_type":"pith_registry"},"payload":{"graph_snapshot":{"author_claims":{"count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","strong_count":0},"builder_version":"pith-number-builder-2026-05-17-v1","claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"integrity":{"available":true,"clean":true,"detectors_run":[],"endpoint":"/pith/2605.27345/integrity.json","findings":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938","summary":{"advisory":0,"by_detector":{},"critical":0,"informational":0}},"paper":{"abstract_excerpt":"Reliable evaluation is essential for understanding large language model (LLM) performance, yet today's go-to metrics, namely token-overlap scores (e.g., ROUGE) and embedding-based measures (e.g., BERTScore), often misjudge semantic similarity of documents. Our study shows that both token-overlap metrics and embedding-based metrics routinely assign nearly identical scores to texts that directly contradict each other, thereby potentially masking fundamental errors. We introduce MATCHA, an automatic metric that jointly rewards semantic agreement with a reference and penalizes contradictions. MATC","authors_text":"Carsten Eickhoff, Ece Sena Etoglu, Seyed Ali Bahrainian, Siran Li","cross_cats":[],"headline":"","license":"http://creativecommons.org/licenses/by-nc-sa/4.0/","primary_cat":"cs.CL","submitted_at":"2026-05-26T17:47:14Z","title":"MATCHA: Matching Text via Contrastive Semantic Alignment"},"references":{"count":0,"internal_anchors":0,"resolved_work":0,"sample":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2605.27345","kind":"arxiv","version":1},"verdict":{"created_at":null,"id":null,"model_set":{},"one_line_summary":"","pipeline_version":null,"pith_extraction_headline":"","strongest_claim":"","weakest_assumption":""}},"verdict_id":null}}],"author_attestations":[],"timestamp_anchors":[],"storage_attestations":[],"citation_signatures":[],"replication_records":[],"corrections":[],"mirror_hints":[],"record_created":{"event_id":"sha256:3fea51366dd090024258997a54f036e9e702fd95767665204df36abdd3d862fd","target":"record","created_at":"2026-05-27T02:06:19Z","signer":{"key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signer_id":"pith.science","signer_type":"pith_registry"},"payload":{"attestation_state":"computed","canonical_record":{"metadata":{"abstract_canon_sha256":"ae0a5272f863ec00a0916dff34b121a5f9ac9651c5a3443aafbdf9c41fae95d6","cross_cats_sorted":[],"license":"http://creativecommons.org/licenses/by-nc-sa/4.0/","primary_cat":"cs.CL","submitted_at":"2026-05-26T17:47:14Z","title_canon_sha256":"929417e6842cc3f218205a362b695386c6a278f2f6060a2d63f63bfe5bfdf524"},"schema_version":"1.0","source":{"id":"2605.27345","kind":"arxiv","version":1}},"canonical_sha256":"ad1f840029e9b8a4176478211d2d61d9ab7eab502ce73c1a09c3d9caacec9e97","receipt":{"algorithm":"ed25519","builder_version":"pith-number-builder-2026-05-17-v1","canonical_sha256":"ad1f840029e9b8a4176478211d2d61d9ab7eab502ce73c1a09c3d9caacec9e97","first_computed_at":"2026-05-27T02:06:19.012511Z","key_id":"pith-v1-2026-05","kind":"pith_receipt","last_reissued_at":"2026-05-27T02:06:19.012511Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","receipt_version":"0.3","signature_b64":"lm+AqgmGL4ixSabapZGgrE7R2o9Elo+b6AfABUMRNO724ItC2S0+UhCe7nDAOEGXcL8hILy4jeA2ZmDUoMsuDA==","signature_status":"signed_v1","signed_at":"2026-05-27T02:06:19.013213Z","signed_message":"canonical_sha256_bytes"},"source_id":"2605.27345","source_kind":"arxiv","source_version":1}}},"equivocations":[],"invalid_events":[],"applied_event_ids":["sha256:3fea51366dd090024258997a54f036e9e702fd95767665204df36abdd3d862fd","sha256:6fa0f8e40b3742fedd3e44f24ee15e3ef9f339a3bc185c27cf14c4e495b3829b"],"state_sha256":"da9ac8411d12c0328612d48a1a54a03bd8debed690ab1bc2712e8e5c2529da1f"}