{"bundle_type":"pith_open_graph_bundle","bundle_version":"1.0","pith_number":"pith:2026:LP4YQRV6M7E3L2NQIIYI3KDNS4","short_pith_number":"pith:LP4YQRV6","canonical_record":{"source":{"id":"2607.27610","kind":"arxiv","version":1},"metadata":{"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.LG","submitted_at":"2026-07-30T02:55:02Z","cross_cats_sorted":[],"title_canon_sha256":"34e56cf46f22c09a87f501eb86dfa77d21e348349566dc07fa3b4f63e4fd6543","abstract_canon_sha256":"76a7376fb39fdf1e50cbc1ca2d15fb496d48f93e24cbc6a2a29247ead4d361fc"},"schema_version":"1.0"},"canonical_sha256":"5bf98846be67c9b5e9b042308da86d97010064daf950c0d6bbb55222f7f2ad5f","source":{"kind":"arxiv","id":"2607.27610","version":1},"source_aliases":[{"alias_kind":"arxiv","alias_value":"2607.27610","created_at":"2026-07-31T01:27:25Z"},{"alias_kind":"arxiv_version","alias_value":"2607.27610v1","created_at":"2026-07-31T01:27:25Z"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2607.27610","created_at":"2026-07-31T01:27:25Z"},{"alias_kind":"pith_short_12","alias_value":"LP4YQRV6M7E3","created_at":"2026-07-31T01:27:25Z"},{"alias_kind":"pith_short_16","alias_value":"LP4YQRV6M7E3L2NQ","created_at":"2026-07-31T01:27:25Z"},{"alias_kind":"pith_short_8","alias_value":"LP4YQRV6","created_at":"2026-07-31T01:27:25Z"}],"events":[{"event_type":"record_created","subject_pith_number":"pith:2026:LP4YQRV6M7E3L2NQIIYI3KDNS4","target":"record","payload":{"canonical_record":{"source":{"id":"2607.27610","kind":"arxiv","version":1},"metadata":{"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.LG","submitted_at":"2026-07-30T02:55:02Z","cross_cats_sorted":[],"title_canon_sha256":"34e56cf46f22c09a87f501eb86dfa77d21e348349566dc07fa3b4f63e4fd6543","abstract_canon_sha256":"76a7376fb39fdf1e50cbc1ca2d15fb496d48f93e24cbc6a2a29247ead4d361fc"},"schema_version":"1.0"},"canonical_sha256":"5bf98846be67c9b5e9b042308da86d97010064daf950c0d6bbb55222f7f2ad5f","receipt":{"kind":"pith_receipt","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"5bf98846be67c9b5e9b042308da86d97010064daf950c0d6bbb55222f7f2ad5f","last_reissued_at":"2026-07-31T01:27:25.845894Z","signature_status":"unsigned_v0","first_computed_at":"2026-07-31T01:27:25.845894Z"},"source_kind":"arxiv","source_id":"2607.27610","source_version":1,"attestation_state":"computed"},"signer":{"signer_id":"pith.science","signer_type":"pith_registry","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"created_at":"2026-07-31T01:27:25Z","supersedes":[],"prev_event":null,"signature":{"signature_status":"signed_v1","algorithm":"ed25519","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signature_b64":"+RvYcl9RtmcXW86CDm2oejyrRUoLf12y2Z3xZAuvq6PIGbYUsNf2aBef9y5GRu2Evt9vYiySdMGq/hw9csJhCQ==","signed_message":"open_graph_event_sha256_bytes","signed_at":"2026-08-04T14:50:45.511461Z"},"content_sha256":"7f5f77c85e60326189e799a5b293a8783953ce924faa0adce6d4f2694b6dd8c4","schema_version":"1.0","event_id":"sha256:7f5f77c85e60326189e799a5b293a8783953ce924faa0adce6d4f2694b6dd8c4"},{"event_type":"graph_snapshot","subject_pith_number":"pith:2026:LP4YQRV6M7E3L2NQIIYI3KDNS4","target":"graph","payload":{"graph_snapshot":{"paper":{"title":"Kalman Meets Curriculum: Efficient Dynamic Prompt Selection for Adaptive RL Finetuning","license":"http://creativecommons.org/licenses/by/4.0/","headline":"","cross_cats":[],"primary_cat":"cs.LG","authors_text":"Baochang Zhang, Haiguang Liu, Haodong Zhu, Linlin Yang, Sheng Xu, Yangyang Ren, Yanjing Li","submitted_at":"2026-07-30T02:55:02Z","abstract_excerpt":"Reinforcement learning (RL) finetuning significantly enhances the reasoning capabilities of large language models (LLMs), yet its effectiveness critically depends on selecting prompts of appropriate difficulty for the current policy. This is challenging because prompt difficulty evolves throughout training. Existing online methods therefore face a trade-off: evaluation-based approaches are accurate but expensive, while prediction-based approaches are efficient but typically assume stationary difficulty, making them ill-suited to RL's non-stationary training dynamics. To address these issues, w"},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2607.27610","kind":"arxiv","version":1},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2607.27610/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"verdict_id":null},"signer":{"signer_id":"pith.science","signer_type":"pith_registry","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"created_at":"2026-07-31T01:27:25Z","supersedes":[],"prev_event":null,"signature":{"signature_status":"signed_v1","algorithm":"ed25519","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signature_b64":"2TLWhWoj9xe1DbpOHLlGG5MJHuE38T/pM9+Q5GHeV/NW3NaOPzwNRW8xuXlk0QL6SxLupeFwKuxm8KD5vgW1Bw==","signed_message":"open_graph_event_sha256_bytes","signed_at":"2026-08-04T14:50:45.512399Z"},"content_sha256":"771d46cd2f46e5d89883ac71482d509a0f08907c3b7c7cc49ae5d7d661250bce","schema_version":"1.0","event_id":"sha256:771d46cd2f46e5d89883ac71482d509a0f08907c3b7c7cc49ae5d7d661250bce"}],"timestamp_proofs":[],"mirror_hints":[{"mirror_type":"https","name":"Pith Resolver","base_url":"https://pith.science","bundle_url":"https://pith.science/pith/LP4YQRV6M7E3L2NQIIYI3KDNS4/bundle.json","state_url":"https://pith.science/pith/LP4YQRV6M7E3L2NQIIYI3KDNS4/state.json","well_known_bundle_url":"https://pith.science/.well-known/pith/LP4YQRV6M7E3L2NQIIYI3KDNS4/bundle.json","status":"primary"}],"public_keys":[{"key_id":"pith-v1-2026-05","algorithm":"ed25519","format":"raw","public_key_b64":"stVStoiQhXFxp4s2pdzPNoqVNBMojDU/fJ2db5S3CbM=","public_key_hex":"b2d552b68890857171a78b36a5dccf368a953413288c353f7c9d9d6f94b709b3","fingerprint_sha256_b32_first128bits":"RVFV5Z2OI2J3ZUO7ERDEBCYNKS","fingerprint_sha256_hex":"8d4b5ee74e4693bcd1df2446408b0d54","rotates_at":null,"url":"https://pith.science/pith-signing-key.json","notes":"Pith uses this Ed25519 key to sign canonical record SHA-256 digests. Verify with: ed25519_verify(public_key, message=canonical_sha256_bytes, signature=base64decode(signature_b64))."}],"merge_version":"pith-open-graph-merge-v1","built_at":"2026-08-04T14:50:45Z","links":{"resolver":"https://pith.science/pith/LP4YQRV6M7E3L2NQIIYI3KDNS4","bundle":"https://pith.science/pith/LP4YQRV6M7E3L2NQIIYI3KDNS4/bundle.json","state":"https://pith.science/pith/LP4YQRV6M7E3L2NQIIYI3KDNS4/state.json","well_known_bundle":"https://pith.science/.well-known/pith/LP4YQRV6M7E3L2NQIIYI3KDNS4/bundle.json"},"state":{"state_type":"pith_open_graph_state","state_version":"1.0","pith_number":"pith:2026:LP4YQRV6M7E3L2NQIIYI3KDNS4","merge_version":"pith-open-graph-merge-v1","event_count":2,"valid_event_count":2,"invalid_event_count":0,"equivocation_count":0,"current":{"canonical_record":{"metadata":{"abstract_canon_sha256":"76a7376fb39fdf1e50cbc1ca2d15fb496d48f93e24cbc6a2a29247ead4d361fc","cross_cats_sorted":[],"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.LG","submitted_at":"2026-07-30T02:55:02Z","title_canon_sha256":"34e56cf46f22c09a87f501eb86dfa77d21e348349566dc07fa3b4f63e4fd6543"},"schema_version":"1.0","source":{"id":"2607.27610","kind":"arxiv","version":1}},"source_aliases":[{"alias_kind":"arxiv","alias_value":"2607.27610","created_at":"2026-07-31T01:27:25Z"},{"alias_kind":"arxiv_version","alias_value":"2607.27610v1","created_at":"2026-07-31T01:27:25Z"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2607.27610","created_at":"2026-07-31T01:27:25Z"},{"alias_kind":"pith_short_12","alias_value":"LP4YQRV6M7E3","created_at":"2026-07-31T01:27:25Z"},{"alias_kind":"pith_short_16","alias_value":"LP4YQRV6M7E3L2NQ","created_at":"2026-07-31T01:27:25Z"},{"alias_kind":"pith_short_8","alias_value":"LP4YQRV6","created_at":"2026-07-31T01:27:25Z"}],"graph_snapshots":[{"event_id":"sha256:771d46cd2f46e5d89883ac71482d509a0f08907c3b7c7cc49ae5d7d661250bce","target":"graph","created_at":"2026-07-31T01:27:25Z","signer":{"key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signer_id":"pith.science","signer_type":"pith_registry"},"payload":{"graph_snapshot":{"author_claims":{"count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","strong_count":0},"builder_version":"pith-number-builder-2026-05-17-v1","claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"integrity":{"available":true,"clean":true,"detectors_run":[],"endpoint":"/pith/2607.27610/integrity.json","findings":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938","summary":{"advisory":0,"by_detector":{},"critical":0,"informational":0}},"paper":{"abstract_excerpt":"Reinforcement learning (RL) finetuning significantly enhances the reasoning capabilities of large language models (LLMs), yet its effectiveness critically depends on selecting prompts of appropriate difficulty for the current policy. This is challenging because prompt difficulty evolves throughout training. Existing online methods therefore face a trade-off: evaluation-based approaches are accurate but expensive, while prediction-based approaches are efficient but typically assume stationary difficulty, making them ill-suited to RL's non-stationary training dynamics. To address these issues, w","authors_text":"Baochang Zhang, Haiguang Liu, Haodong Zhu, Linlin Yang, Sheng Xu, Yangyang Ren, Yanjing Li","cross_cats":[],"headline":"","license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.LG","submitted_at":"2026-07-30T02:55:02Z","title":"Kalman Meets Curriculum: Efficient Dynamic Prompt Selection for Adaptive RL Finetuning"},"references":{"count":0,"internal_anchors":0,"resolved_work":0,"sample":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2607.27610","kind":"arxiv","version":1},"verdict":{"created_at":null,"id":null,"model_set":{},"one_line_summary":"","pipeline_version":null,"pith_extraction_headline":"","strongest_claim":"","weakest_assumption":""}},"verdict_id":null}}],"author_attestations":[],"timestamp_anchors":[],"storage_attestations":[],"citation_signatures":[],"replication_records":[],"corrections":[],"mirror_hints":[],"record_created":{"event_id":"sha256:7f5f77c85e60326189e799a5b293a8783953ce924faa0adce6d4f2694b6dd8c4","target":"record","created_at":"2026-07-31T01:27:25Z","signer":{"key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signer_id":"pith.science","signer_type":"pith_registry"},"payload":{"attestation_state":"computed","canonical_record":{"metadata":{"abstract_canon_sha256":"76a7376fb39fdf1e50cbc1ca2d15fb496d48f93e24cbc6a2a29247ead4d361fc","cross_cats_sorted":[],"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.LG","submitted_at":"2026-07-30T02:55:02Z","title_canon_sha256":"34e56cf46f22c09a87f501eb86dfa77d21e348349566dc07fa3b4f63e4fd6543"},"schema_version":"1.0","source":{"id":"2607.27610","kind":"arxiv","version":1}},"canonical_sha256":"5bf98846be67c9b5e9b042308da86d97010064daf950c0d6bbb55222f7f2ad5f","receipt":{"builder_version":"pith-number-builder-2026-05-17-v1","canonical_sha256":"5bf98846be67c9b5e9b042308da86d97010064daf950c0d6bbb55222f7f2ad5f","first_computed_at":"2026-07-31T01:27:25.845894Z","kind":"pith_receipt","last_reissued_at":"2026-07-31T01:27:25.845894Z","receipt_version":"0.3","signature_status":"unsigned_v0"},"source_id":"2607.27610","source_kind":"arxiv","source_version":1}}},"equivocations":[],"invalid_events":[],"applied_event_ids":["sha256:7f5f77c85e60326189e799a5b293a8783953ce924faa0adce6d4f2694b6dd8c4","sha256:771d46cd2f46e5d89883ac71482d509a0f08907c3b7c7cc49ae5d7d661250bce"],"state_sha256":"b0e1bbd906693c437b07bfcd4964ff2308934c4d3376e674d04166f016fe05c3"},"bundle_signature":{"signature_status":"signed_v1","algorithm":"ed25519","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signature_b64":"wSiziabxvcezg9LaK9J77y6jNKHr67zy3gLZnKD9LvztlDC9DCt6bvcQdUqCTri9br5/amHFVBqC8pi9PfSrCA==","signed_message":"bundle_sha256_bytes","signed_at":"2026-08-04T14:50:45.522436Z","bundle_sha256":"e55875e9166599120e62286ca49707f12234d088b3f6859a1325b9e56ad5a289"}}