{"bundle_type":"pith_open_graph_bundle","bundle_version":"1.0","pith_number":"pith:2025:MWPBOOMLK3XGTVHDGHEQFCVGDN","short_pith_number":"pith:MWPBOOML","canonical_record":{"source":{"id":"2510.00492","kind":"arxiv","version":3},"metadata":{"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.AI","submitted_at":"2025-10-01T04:21:14Z","cross_cats_sorted":[],"title_canon_sha256":"09729f63a62eaeea58d269bc27b03deed6efd725bae462158451ed900da57e5d","abstract_canon_sha256":"496b494e03389a1e856128de68e1297249a275a3d83478bfe20f8653d2d53074"},"schema_version":"1.0"},"canonical_sha256":"659e17398b56ee69d4e331c9028aa61b44c91114ada8428da7f54f6cf5892aa1","source":{"kind":"arxiv","id":"2510.00492","version":3},"source_aliases":[{"alias_kind":"arxiv","alias_value":"2510.00492","created_at":"2026-07-15T00:21:12Z"},{"alias_kind":"arxiv_version","alias_value":"2510.00492v3","created_at":"2026-07-15T00:21:12Z"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2510.00492","created_at":"2026-07-15T00:21:12Z"},{"alias_kind":"pith_short_12","alias_value":"MWPBOOMLK3XG","created_at":"2026-07-15T00:21:12Z"},{"alias_kind":"pith_short_16","alias_value":"MWPBOOMLK3XGTVHD","created_at":"2026-07-15T00:21:12Z"},{"alias_kind":"pith_short_8","alias_value":"MWPBOOML","created_at":"2026-07-15T00:21:12Z"}],"events":[{"event_type":"record_created","subject_pith_number":"pith:2025:MWPBOOMLK3XGTVHDGHEQFCVGDN","target":"record","payload":{"canonical_record":{"source":{"id":"2510.00492","kind":"arxiv","version":3},"metadata":{"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.AI","submitted_at":"2025-10-01T04:21:14Z","cross_cats_sorted":[],"title_canon_sha256":"09729f63a62eaeea58d269bc27b03deed6efd725bae462158451ed900da57e5d","abstract_canon_sha256":"496b494e03389a1e856128de68e1297249a275a3d83478bfe20f8653d2d53074"},"schema_version":"1.0"},"canonical_sha256":"659e17398b56ee69d4e331c9028aa61b44c91114ada8428da7f54f6cf5892aa1","receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-07-15T00:21:12.982625Z","signature_b64":"2YCS+N1RxOQEtwrFxOebCnfWnc0fJ0KkorR1jftO37Igztr/HBGNtyokoKYsnctRVus+cOjigP414Rb4AMfkAQ==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"659e17398b56ee69d4e331c9028aa61b44c91114ada8428da7f54f6cf5892aa1","last_reissued_at":"2026-07-15T00:21:12.981700Z","signature_status":"signed_v1","first_computed_at":"2026-07-15T00:21:12.981700Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"source_kind":"arxiv","source_id":"2510.00492","source_version":3,"attestation_state":"computed"},"signer":{"signer_id":"pith.science","signer_type":"pith_registry","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"created_at":"2026-07-15T00:21:12Z","supersedes":[],"prev_event":null,"signature":{"signature_status":"signed_v1","algorithm":"ed25519","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signature_b64":"nfF2H0XL1UfU5oSLNS0v+9sr5djW07IpTSF+QAJUWD3T6xgOCXsify/rVI3biVM2scAm3P9OxhsSJR1MtqgZBg==","signed_message":"open_graph_event_sha256_bytes","signed_at":"2026-08-02T12:29:44.433100Z"},"content_sha256":"b426b7f723f91ddef12d8b9e0e9d6d6d16c2303de1b10d67d8c24e7ff986ff07","schema_version":"1.0","event_id":"sha256:b426b7f723f91ddef12d8b9e0e9d6d6d16c2303de1b10d67d8c24e7ff986ff07"},{"event_type":"graph_snapshot","subject_pith_number":"pith:2025:MWPBOOMLK3XGTVHDGHEQFCVGDN","target":"graph","payload":{"graph_snapshot":{"paper":{"title":"Rethinking Reward Models for Multi-Domain Test-Time Scaling","license":"http://creativecommons.org/licenses/by/4.0/","headline":"","cross_cats":[],"primary_cat":"cs.AI","authors_text":"Dominik Wagner, Dong Bok Lee, Dongki Kim, Heejun Lee, Jiang Bian, Jingjing Fu, Jinheon Baek, Jinyu Wang, Jiongdao Jin, Lei Song, Minki Kang, Sangwoo Park, Seanie Lee, Sung Ju Hwang, Tobias Bocklet","submitted_at":"2025-10-01T04:21:14Z","abstract_excerpt":"The reliability of large language models (LLMs) during test-time scaling is often assessed with \\emph{external verifiers} or \\emph{reward models} that distinguish correct reasoning from flawed logic. Prior work has studied both outcome reward models (ORMs), which assess only the final answer, and process reward models (PRMs), which score intermediate reasoning steps. Although PRMs are often viewed as advantageous due to their finer-grained supervision, much of the supporting evidence comes from math-adjacent settings, and their relative benefits across broader domains remain unclear. We presen"},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2510.00492","kind":"arxiv","version":3},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2510.00492/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"verdict_id":null},"signer":{"signer_id":"pith.science","signer_type":"pith_registry","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"created_at":"2026-07-15T00:21:12Z","supersedes":[],"prev_event":null,"signature":{"signature_status":"signed_v1","algorithm":"ed25519","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signature_b64":"KhDz94hktxhddweThqSMFEqdQx86GdoR3Jz7k9SwLPnw0fLzlCQVUjy0yjcHAiwLrgEsqBB9roGyHQj6QeMNBQ==","signed_message":"open_graph_event_sha256_bytes","signed_at":"2026-08-02T12:29:44.433694Z"},"content_sha256":"075830f14538a7c2ed34499d3bbf534cf952f41bf896e955aafd7577b9d4a351","schema_version":"1.0","event_id":"sha256:075830f14538a7c2ed34499d3bbf534cf952f41bf896e955aafd7577b9d4a351"}],"timestamp_proofs":[],"mirror_hints":[{"mirror_type":"https","name":"Pith Resolver","base_url":"https://pith.science","bundle_url":"https://pith.science/pith/MWPBOOMLK3XGTVHDGHEQFCVGDN/bundle.json","state_url":"https://pith.science/pith/MWPBOOMLK3XGTVHDGHEQFCVGDN/state.json","well_known_bundle_url":"https://pith.science/.well-known/pith/MWPBOOMLK3XGTVHDGHEQFCVGDN/bundle.json","status":"primary"}],"public_keys":[{"key_id":"pith-v1-2026-05","algorithm":"ed25519","format":"raw","public_key_b64":"stVStoiQhXFxp4s2pdzPNoqVNBMojDU/fJ2db5S3CbM=","public_key_hex":"b2d552b68890857171a78b36a5dccf368a953413288c353f7c9d9d6f94b709b3","fingerprint_sha256_b32_first128bits":"RVFV5Z2OI2J3ZUO7ERDEBCYNKS","fingerprint_sha256_hex":"8d4b5ee74e4693bcd1df2446408b0d54","rotates_at":null,"url":"https://pith.science/pith-signing-key.json","notes":"Pith uses this Ed25519 key to sign canonical record SHA-256 digests. Verify with: ed25519_verify(public_key, message=canonical_sha256_bytes, signature=base64decode(signature_b64))."}],"merge_version":"pith-open-graph-merge-v1","built_at":"2026-08-02T12:29:44Z","links":{"resolver":"https://pith.science/pith/MWPBOOMLK3XGTVHDGHEQFCVGDN","bundle":"https://pith.science/pith/MWPBOOMLK3XGTVHDGHEQFCVGDN/bundle.json","state":"https://pith.science/pith/MWPBOOMLK3XGTVHDGHEQFCVGDN/state.json","well_known_bundle":"https://pith.science/.well-known/pith/MWPBOOMLK3XGTVHDGHEQFCVGDN/bundle.json"},"state":{"state_type":"pith_open_graph_state","state_version":"1.0","pith_number":"pith:2025:MWPBOOMLK3XGTVHDGHEQFCVGDN","merge_version":"pith-open-graph-merge-v1","event_count":2,"valid_event_count":2,"invalid_event_count":0,"equivocation_count":0,"current":{"canonical_record":{"metadata":{"abstract_canon_sha256":"496b494e03389a1e856128de68e1297249a275a3d83478bfe20f8653d2d53074","cross_cats_sorted":[],"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.AI","submitted_at":"2025-10-01T04:21:14Z","title_canon_sha256":"09729f63a62eaeea58d269bc27b03deed6efd725bae462158451ed900da57e5d"},"schema_version":"1.0","source":{"id":"2510.00492","kind":"arxiv","version":3}},"source_aliases":[{"alias_kind":"arxiv","alias_value":"2510.00492","created_at":"2026-07-15T00:21:12Z"},{"alias_kind":"arxiv_version","alias_value":"2510.00492v3","created_at":"2026-07-15T00:21:12Z"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2510.00492","created_at":"2026-07-15T00:21:12Z"},{"alias_kind":"pith_short_12","alias_value":"MWPBOOMLK3XG","created_at":"2026-07-15T00:21:12Z"},{"alias_kind":"pith_short_16","alias_value":"MWPBOOMLK3XGTVHD","created_at":"2026-07-15T00:21:12Z"},{"alias_kind":"pith_short_8","alias_value":"MWPBOOML","created_at":"2026-07-15T00:21:12Z"}],"graph_snapshots":[{"event_id":"sha256:075830f14538a7c2ed34499d3bbf534cf952f41bf896e955aafd7577b9d4a351","target":"graph","created_at":"2026-07-15T00:21:12Z","signer":{"key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signer_id":"pith.science","signer_type":"pith_registry"},"payload":{"graph_snapshot":{"author_claims":{"count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","strong_count":0},"builder_version":"pith-number-builder-2026-05-17-v1","claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"integrity":{"available":true,"clean":true,"detectors_run":[],"endpoint":"/pith/2510.00492/integrity.json","findings":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938","summary":{"advisory":0,"by_detector":{},"critical":0,"informational":0}},"paper":{"abstract_excerpt":"The reliability of large language models (LLMs) during test-time scaling is often assessed with \\emph{external verifiers} or \\emph{reward models} that distinguish correct reasoning from flawed logic. Prior work has studied both outcome reward models (ORMs), which assess only the final answer, and process reward models (PRMs), which score intermediate reasoning steps. Although PRMs are often viewed as advantageous due to their finer-grained supervision, much of the supporting evidence comes from math-adjacent settings, and their relative benefits across broader domains remain unclear. We presen","authors_text":"Dominik Wagner, Dong Bok Lee, Dongki Kim, Heejun Lee, Jiang Bian, Jingjing Fu, Jinheon Baek, Jinyu Wang, Jiongdao Jin, Lei Song, Minki Kang, Sangwoo Park, Seanie Lee, Sung Ju Hwang, Tobias Bocklet","cross_cats":[],"headline":"","license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.AI","submitted_at":"2025-10-01T04:21:14Z","title":"Rethinking Reward Models for Multi-Domain Test-Time Scaling"},"references":{"count":0,"internal_anchors":0,"resolved_work":0,"sample":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2510.00492","kind":"arxiv","version":3},"verdict":{"created_at":null,"id":null,"model_set":{},"one_line_summary":"","pipeline_version":null,"pith_extraction_headline":"","strongest_claim":"","weakest_assumption":""}},"verdict_id":null}}],"author_attestations":[],"timestamp_anchors":[],"storage_attestations":[],"citation_signatures":[],"replication_records":[],"corrections":[],"mirror_hints":[],"record_created":{"event_id":"sha256:b426b7f723f91ddef12d8b9e0e9d6d6d16c2303de1b10d67d8c24e7ff986ff07","target":"record","created_at":"2026-07-15T00:21:12Z","signer":{"key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signer_id":"pith.science","signer_type":"pith_registry"},"payload":{"attestation_state":"computed","canonical_record":{"metadata":{"abstract_canon_sha256":"496b494e03389a1e856128de68e1297249a275a3d83478bfe20f8653d2d53074","cross_cats_sorted":[],"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.AI","submitted_at":"2025-10-01T04:21:14Z","title_canon_sha256":"09729f63a62eaeea58d269bc27b03deed6efd725bae462158451ed900da57e5d"},"schema_version":"1.0","source":{"id":"2510.00492","kind":"arxiv","version":3}},"canonical_sha256":"659e17398b56ee69d4e331c9028aa61b44c91114ada8428da7f54f6cf5892aa1","receipt":{"algorithm":"ed25519","builder_version":"pith-number-builder-2026-05-17-v1","canonical_sha256":"659e17398b56ee69d4e331c9028aa61b44c91114ada8428da7f54f6cf5892aa1","first_computed_at":"2026-07-15T00:21:12.981700Z","key_id":"pith-v1-2026-05","kind":"pith_receipt","last_reissued_at":"2026-07-15T00:21:12.981700Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","receipt_version":"0.3","signature_b64":"2YCS+N1RxOQEtwrFxOebCnfWnc0fJ0KkorR1jftO37Igztr/HBGNtyokoKYsnctRVus+cOjigP414Rb4AMfkAQ==","signature_status":"signed_v1","signed_at":"2026-07-15T00:21:12.982625Z","signed_message":"canonical_sha256_bytes"},"source_id":"2510.00492","source_kind":"arxiv","source_version":3}}},"equivocations":[],"invalid_events":[],"applied_event_ids":["sha256:b426b7f723f91ddef12d8b9e0e9d6d6d16c2303de1b10d67d8c24e7ff986ff07","sha256:075830f14538a7c2ed34499d3bbf534cf952f41bf896e955aafd7577b9d4a351"],"state_sha256":"c415098f9d8143b089d65bd690cffbedd0f1295261ddd2a03e9217c55d38e3ca"},"bundle_signature":{"signature_status":"signed_v1","algorithm":"ed25519","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signature_b64":"nVpQLV/Xrj4gDD6kqPkTY4TG9yica5hlD6aFM0KvPKT0OYcy9YY27jXn32feIthQTW3gUF9SwhRZHacu/31iAg==","signed_message":"bundle_sha256_bytes","signed_at":"2026-08-02T12:29:44.439383Z","bundle_sha256":"1446e74afb31907420530b47d6d4c26082751f5c2eb02b1280c8b9afa048bbe1"}}