{"state_type":"pith_open_graph_state","state_version":"1.0","pith_number":"pith:2026:LP4YQRV6M7E3L2NQIIYI3KDNS4","merge_version":"pith-open-graph-merge-v1","event_count":2,"valid_event_count":2,"invalid_event_count":0,"equivocation_count":0,"current":{"canonical_record":{"metadata":{"abstract_canon_sha256":"76a7376fb39fdf1e50cbc1ca2d15fb496d48f93e24cbc6a2a29247ead4d361fc","cross_cats_sorted":[],"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.LG","submitted_at":"2026-07-30T02:55:02Z","title_canon_sha256":"34e56cf46f22c09a87f501eb86dfa77d21e348349566dc07fa3b4f63e4fd6543"},"schema_version":"1.0","source":{"id":"2607.27610","kind":"arxiv","version":1}},"source_aliases":[{"alias_kind":"arxiv","alias_value":"2607.27610","created_at":"2026-07-31T01:27:25Z"},{"alias_kind":"arxiv_version","alias_value":"2607.27610v1","created_at":"2026-07-31T01:27:25Z"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2607.27610","created_at":"2026-07-31T01:27:25Z"},{"alias_kind":"pith_short_12","alias_value":"LP4YQRV6M7E3","created_at":"2026-07-31T01:27:25Z"},{"alias_kind":"pith_short_16","alias_value":"LP4YQRV6M7E3L2NQ","created_at":"2026-07-31T01:27:25Z"},{"alias_kind":"pith_short_8","alias_value":"LP4YQRV6","created_at":"2026-07-31T01:27:25Z"}],"graph_snapshots":[{"event_id":"sha256:771d46cd2f46e5d89883ac71482d509a0f08907c3b7c7cc49ae5d7d661250bce","target":"graph","created_at":"2026-07-31T01:27:25Z","signer":{"key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signer_id":"pith.science","signer_type":"pith_registry"},"payload":{"graph_snapshot":{"author_claims":{"count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","strong_count":0},"builder_version":"pith-number-builder-2026-05-17-v1","claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"integrity":{"available":true,"clean":true,"detectors_run":[],"endpoint":"/pith/2607.27610/integrity.json","findings":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938","summary":{"advisory":0,"by_detector":{},"critical":0,"informational":0}},"paper":{"abstract_excerpt":"Reinforcement learning (RL) finetuning significantly enhances the reasoning capabilities of large language models (LLMs), yet its effectiveness critically depends on selecting prompts of appropriate difficulty for the current policy. This is challenging because prompt difficulty evolves throughout training. Existing online methods therefore face a trade-off: evaluation-based approaches are accurate but expensive, while prediction-based approaches are efficient but typically assume stationary difficulty, making them ill-suited to RL's non-stationary training dynamics. To address these issues, w","authors_text":"Baochang Zhang, Haiguang Liu, Haodong Zhu, Linlin Yang, Sheng Xu, Yangyang Ren, Yanjing Li","cross_cats":[],"headline":"","license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.LG","submitted_at":"2026-07-30T02:55:02Z","title":"Kalman Meets Curriculum: Efficient Dynamic Prompt Selection for Adaptive RL Finetuning"},"references":{"count":0,"internal_anchors":0,"resolved_work":0,"sample":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2607.27610","kind":"arxiv","version":1},"verdict":{"created_at":null,"id":null,"model_set":{},"one_line_summary":"","pipeline_version":null,"pith_extraction_headline":"","strongest_claim":"","weakest_assumption":""}},"verdict_id":null}}],"author_attestations":[],"timestamp_anchors":[],"storage_attestations":[],"citation_signatures":[],"replication_records":[],"corrections":[],"mirror_hints":[],"record_created":{"event_id":"sha256:7f5f77c85e60326189e799a5b293a8783953ce924faa0adce6d4f2694b6dd8c4","target":"record","created_at":"2026-07-31T01:27:25Z","signer":{"key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signer_id":"pith.science","signer_type":"pith_registry"},"payload":{"attestation_state":"computed","canonical_record":{"metadata":{"abstract_canon_sha256":"76a7376fb39fdf1e50cbc1ca2d15fb496d48f93e24cbc6a2a29247ead4d361fc","cross_cats_sorted":[],"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.LG","submitted_at":"2026-07-30T02:55:02Z","title_canon_sha256":"34e56cf46f22c09a87f501eb86dfa77d21e348349566dc07fa3b4f63e4fd6543"},"schema_version":"1.0","source":{"id":"2607.27610","kind":"arxiv","version":1}},"canonical_sha256":"5bf98846be67c9b5e9b042308da86d97010064daf950c0d6bbb55222f7f2ad5f","receipt":{"builder_version":"pith-number-builder-2026-05-17-v1","canonical_sha256":"5bf98846be67c9b5e9b042308da86d97010064daf950c0d6bbb55222f7f2ad5f","first_computed_at":"2026-07-31T01:27:25.845894Z","kind":"pith_receipt","last_reissued_at":"2026-07-31T01:27:25.845894Z","receipt_version":"0.3","signature_status":"unsigned_v0"},"source_id":"2607.27610","source_kind":"arxiv","source_version":1}}},"equivocations":[],"invalid_events":[],"applied_event_ids":["sha256:7f5f77c85e60326189e799a5b293a8783953ce924faa0adce6d4f2694b6dd8c4","sha256:771d46cd2f46e5d89883ac71482d509a0f08907c3b7c7cc49ae5d7d661250bce"],"state_sha256":"b0e1bbd906693c437b07bfcd4964ff2308934c4d3376e674d04166f016fe05c3"}