{"bundle_type":"pith_open_graph_bundle","bundle_version":"1.0","pith_number":"pith:2023:TNJ6H6Y4KIK55WXLWKEOWM2OUV","short_pith_number":"pith:TNJ6H6Y4","canonical_record":{"source":{"id":"2310.16487","kind":"arxiv","version":1},"metadata":{"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.LG","submitted_at":"2023-10-25T09:17:25Z","cross_cats_sorted":[],"title_canon_sha256":"26763ec6326aae20a782665ea81d49d5fe39f983a049c3c59871b4248ed79519","abstract_canon_sha256":"72e4043061d07a203852063da6c3358250fa483a10fb4f3fcaa9f33c9368d2aa"},"schema_version":"1.0"},"canonical_sha256":"9b53e3fb1c5215dedaebb288eb334ea54f746e4386a6446d65a0c54dd0e29e26","source":{"kind":"arxiv","id":"2310.16487","version":1},"source_aliases":[{"alias_kind":"arxiv","alias_value":"2310.16487","created_at":"2026-07-05T07:04:58Z"},{"alias_kind":"arxiv_version","alias_value":"2310.16487v1","created_at":"2026-07-05T07:04:58Z"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2310.16487","created_at":"2026-07-05T07:04:58Z"},{"alias_kind":"pith_short_12","alias_value":"TNJ6H6Y4KIK5","created_at":"2026-07-05T07:04:58Z"},{"alias_kind":"pith_short_16","alias_value":"TNJ6H6Y4KIK55WXL","created_at":"2026-07-05T07:04:58Z"},{"alias_kind":"pith_short_8","alias_value":"TNJ6H6Y4","created_at":"2026-07-05T07:04:58Z"}],"events":[{"event_type":"record_created","subject_pith_number":"pith:2023:TNJ6H6Y4KIK55WXLWKEOWM2OUV","target":"record","payload":{"canonical_record":{"source":{"id":"2310.16487","kind":"arxiv","version":1},"metadata":{"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.LG","submitted_at":"2023-10-25T09:17:25Z","cross_cats_sorted":[],"title_canon_sha256":"26763ec6326aae20a782665ea81d49d5fe39f983a049c3c59871b4248ed79519","abstract_canon_sha256":"72e4043061d07a203852063da6c3358250fa483a10fb4f3fcaa9f33c9368d2aa"},"schema_version":"1.0"},"canonical_sha256":"9b53e3fb1c5215dedaebb288eb334ea54f746e4386a6446d65a0c54dd0e29e26","receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-07-05T07:04:58.239027Z","signature_b64":"R8RiNb79B4I8mAUYi9aclMifHGQUq/Ummn0xrx9JI2dORdiUb+qxU+VddlSfyHpP26Mkz06Yel3JuYnuZqHMAA==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"9b53e3fb1c5215dedaebb288eb334ea54f746e4386a6446d65a0c54dd0e29e26","last_reissued_at":"2026-07-05T07:04:58.238575Z","signature_status":"signed_v1","first_computed_at":"2026-07-05T07:04:58.238575Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"source_kind":"arxiv","source_id":"2310.16487","source_version":1,"attestation_state":"computed"},"signer":{"signer_id":"pith.science","signer_type":"pith_registry","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"created_at":"2026-07-05T07:04:58Z","supersedes":[],"prev_event":null,"signature":{"signature_status":"signed_v1","algorithm":"ed25519","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signature_b64":"0xZioUBa3WHln1o3Ubb7eZ4GHGIOzgXkNLGsnM0HYnRN9oNpW6u3kZwCtLj/CGv67XqPX4BQWOWY49hq+iRLBQ==","signed_message":"open_graph_event_sha256_bytes","signed_at":"2026-08-10T13:27:36.755993Z"},"content_sha256":"2261b4e20ca913444dbb4a34c18246b4da30c931e8afc813bd261c093b2f638c","schema_version":"1.0","event_id":"sha256:2261b4e20ca913444dbb4a34c18246b4da30c931e8afc813bd261c093b2f638c"},{"event_type":"graph_snapshot","subject_pith_number":"pith:2023:TNJ6H6Y4KIK55WXLWKEOWM2OUV","target":"graph","payload":{"graph_snapshot":{"paper":{"title":"Hyperparameter Optimization for Multi-Objective Reinforcement Learning","license":"http://creativecommons.org/licenses/by/4.0/","headline":"","cross_cats":[],"primary_cat":"cs.LG","authors_text":"Daniel Gareev, El-Ghazali Talbi, Florian Felten, Gr\\'egoire Danoy","submitted_at":"2023-10-25T09:17:25Z","abstract_excerpt":"Reinforcement learning (RL) has emerged as a powerful approach for tackling complex problems. The recent introduction of multi-objective reinforcement learning (MORL) has further expanded the scope of RL by enabling agents to make trade-offs among multiple objectives. This advancement not only has broadened the range of problems that can be tackled but also created numerous opportunities for exploration and advancement. Yet, the effectiveness of RL agents heavily relies on appropriately setting their hyperparameters. In practice, this task often proves to be challenging, leading to unsuccessfu"},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2310.16487","kind":"arxiv","version":1},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2310.16487/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"verdict_id":null},"signer":{"signer_id":"pith.science","signer_type":"pith_registry","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"created_at":"2026-07-05T07:04:58Z","supersedes":[],"prev_event":null,"signature":{"signature_status":"signed_v1","algorithm":"ed25519","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signature_b64":"iFq7kgrkQte+Q8Qk0yzimtbLZPdse46uRcoNgaOeAF9bCoi2+f2uUo8qxvI+FnNxAU0jrpB465UZ0D0FhBfvAw==","signed_message":"open_graph_event_sha256_bytes","signed_at":"2026-08-10T13:27:36.756548Z"},"content_sha256":"9f0e362a5e4ffc26c7ca339bd7915721ce182d01f9ece1fbeab5f61b4e361545","schema_version":"1.0","event_id":"sha256:9f0e362a5e4ffc26c7ca339bd7915721ce182d01f9ece1fbeab5f61b4e361545"}],"timestamp_proofs":[],"mirror_hints":[{"mirror_type":"https","name":"Pith Resolver","base_url":"https://pith.science","bundle_url":"https://pith.science/pith/TNJ6H6Y4KIK55WXLWKEOWM2OUV/bundle.json","state_url":"https://pith.science/pith/TNJ6H6Y4KIK55WXLWKEOWM2OUV/state.json","well_known_bundle_url":"https://pith.science/.well-known/pith/TNJ6H6Y4KIK55WXLWKEOWM2OUV/bundle.json","status":"primary"}],"public_keys":[{"key_id":"pith-v1-2026-05","algorithm":"ed25519","format":"raw","public_key_b64":"stVStoiQhXFxp4s2pdzPNoqVNBMojDU/fJ2db5S3CbM=","public_key_hex":"b2d552b68890857171a78b36a5dccf368a953413288c353f7c9d9d6f94b709b3","fingerprint_sha256_b32_first128bits":"RVFV5Z2OI2J3ZUO7ERDEBCYNKS","fingerprint_sha256_hex":"8d4b5ee74e4693bcd1df2446408b0d54","rotates_at":null,"url":"https://pith.science/pith-signing-key.json","notes":"Pith uses this Ed25519 key to sign canonical record SHA-256 digests. Verify with: ed25519_verify(public_key, message=canonical_sha256_bytes, signature=base64decode(signature_b64))."}],"merge_version":"pith-open-graph-merge-v1","built_at":"2026-08-10T13:27:36Z","links":{"resolver":"https://pith.science/pith/TNJ6H6Y4KIK55WXLWKEOWM2OUV","bundle":"https://pith.science/pith/TNJ6H6Y4KIK55WXLWKEOWM2OUV/bundle.json","state":"https://pith.science/pith/TNJ6H6Y4KIK55WXLWKEOWM2OUV/state.json","well_known_bundle":"https://pith.science/.well-known/pith/TNJ6H6Y4KIK55WXLWKEOWM2OUV/bundle.json"},"state":{"state_type":"pith_open_graph_state","state_version":"1.0","pith_number":"pith:2023:TNJ6H6Y4KIK55WXLWKEOWM2OUV","merge_version":"pith-open-graph-merge-v1","event_count":2,"valid_event_count":2,"invalid_event_count":0,"equivocation_count":0,"current":{"canonical_record":{"metadata":{"abstract_canon_sha256":"72e4043061d07a203852063da6c3358250fa483a10fb4f3fcaa9f33c9368d2aa","cross_cats_sorted":[],"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.LG","submitted_at":"2023-10-25T09:17:25Z","title_canon_sha256":"26763ec6326aae20a782665ea81d49d5fe39f983a049c3c59871b4248ed79519"},"schema_version":"1.0","source":{"id":"2310.16487","kind":"arxiv","version":1}},"source_aliases":[{"alias_kind":"arxiv","alias_value":"2310.16487","created_at":"2026-07-05T07:04:58Z"},{"alias_kind":"arxiv_version","alias_value":"2310.16487v1","created_at":"2026-07-05T07:04:58Z"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2310.16487","created_at":"2026-07-05T07:04:58Z"},{"alias_kind":"pith_short_12","alias_value":"TNJ6H6Y4KIK5","created_at":"2026-07-05T07:04:58Z"},{"alias_kind":"pith_short_16","alias_value":"TNJ6H6Y4KIK55WXL","created_at":"2026-07-05T07:04:58Z"},{"alias_kind":"pith_short_8","alias_value":"TNJ6H6Y4","created_at":"2026-07-05T07:04:58Z"}],"graph_snapshots":[{"event_id":"sha256:9f0e362a5e4ffc26c7ca339bd7915721ce182d01f9ece1fbeab5f61b4e361545","target":"graph","created_at":"2026-07-05T07:04:58Z","signer":{"key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signer_id":"pith.science","signer_type":"pith_registry"},"payload":{"graph_snapshot":{"author_claims":{"count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","strong_count":0},"builder_version":"pith-number-builder-2026-05-17-v1","claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"integrity":{"available":true,"clean":true,"detectors_run":[],"endpoint":"/pith/2310.16487/integrity.json","findings":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938","summary":{"advisory":0,"by_detector":{},"critical":0,"informational":0}},"paper":{"abstract_excerpt":"Reinforcement learning (RL) has emerged as a powerful approach for tackling complex problems. The recent introduction of multi-objective reinforcement learning (MORL) has further expanded the scope of RL by enabling agents to make trade-offs among multiple objectives. This advancement not only has broadened the range of problems that can be tackled but also created numerous opportunities for exploration and advancement. Yet, the effectiveness of RL agents heavily relies on appropriately setting their hyperparameters. In practice, this task often proves to be challenging, leading to unsuccessfu","authors_text":"Daniel Gareev, El-Ghazali Talbi, Florian Felten, Gr\\'egoire Danoy","cross_cats":[],"headline":"","license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.LG","submitted_at":"2023-10-25T09:17:25Z","title":"Hyperparameter Optimization for Multi-Objective Reinforcement Learning"},"references":{"count":0,"internal_anchors":0,"resolved_work":0,"sample":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2310.16487","kind":"arxiv","version":1},"verdict":{"created_at":null,"id":null,"model_set":{},"one_line_summary":"","pipeline_version":null,"pith_extraction_headline":"","strongest_claim":"","weakest_assumption":""}},"verdict_id":null}}],"author_attestations":[],"timestamp_anchors":[],"storage_attestations":[],"citation_signatures":[],"replication_records":[],"corrections":[],"mirror_hints":[],"record_created":{"event_id":"sha256:2261b4e20ca913444dbb4a34c18246b4da30c931e8afc813bd261c093b2f638c","target":"record","created_at":"2026-07-05T07:04:58Z","signer":{"key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signer_id":"pith.science","signer_type":"pith_registry"},"payload":{"attestation_state":"computed","canonical_record":{"metadata":{"abstract_canon_sha256":"72e4043061d07a203852063da6c3358250fa483a10fb4f3fcaa9f33c9368d2aa","cross_cats_sorted":[],"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.LG","submitted_at":"2023-10-25T09:17:25Z","title_canon_sha256":"26763ec6326aae20a782665ea81d49d5fe39f983a049c3c59871b4248ed79519"},"schema_version":"1.0","source":{"id":"2310.16487","kind":"arxiv","version":1}},"canonical_sha256":"9b53e3fb1c5215dedaebb288eb334ea54f746e4386a6446d65a0c54dd0e29e26","receipt":{"algorithm":"ed25519","builder_version":"pith-number-builder-2026-05-17-v1","canonical_sha256":"9b53e3fb1c5215dedaebb288eb334ea54f746e4386a6446d65a0c54dd0e29e26","first_computed_at":"2026-07-05T07:04:58.238575Z","key_id":"pith-v1-2026-05","kind":"pith_receipt","last_reissued_at":"2026-07-05T07:04:58.238575Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","receipt_version":"0.3","signature_b64":"R8RiNb79B4I8mAUYi9aclMifHGQUq/Ummn0xrx9JI2dORdiUb+qxU+VddlSfyHpP26Mkz06Yel3JuYnuZqHMAA==","signature_status":"signed_v1","signed_at":"2026-07-05T07:04:58.239027Z","signed_message":"canonical_sha256_bytes"},"source_id":"2310.16487","source_kind":"arxiv","source_version":1}}},"equivocations":[],"invalid_events":[],"applied_event_ids":["sha256:2261b4e20ca913444dbb4a34c18246b4da30c931e8afc813bd261c093b2f638c","sha256:9f0e362a5e4ffc26c7ca339bd7915721ce182d01f9ece1fbeab5f61b4e361545"],"state_sha256":"b4e4d8df7e3cbae84849642ee887931656b53e48edcb6067569f8ffdd90be6f1"},"bundle_signature":{"signature_status":"signed_v1","algorithm":"ed25519","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signature_b64":"d/y85+7uscYSfiGqSqIyEqFRt6T+p+ueFjj79CIR/sBBKk2HMfype1LTrsgqHGYA7q50WCMshUuJezTveyTSAg==","signed_message":"bundle_sha256_bytes","signed_at":"2026-08-10T13:27:36.763850Z","bundle_sha256":"7692101919a8ed0c7fab96f6b9a95b0a038294f70c59c11de8ed71691efc6289"}}