{"record_type":"pith_number_record","schema_url":"https://pith.science/schemas/pith-number/v1.json","pith_number":"pith:2024:X5DSCMJU4EFU4DQVPLGVLPBIOH","short_pith_number":"pith:X5DSCMJU","schema_version":"1.0","canonical_sha256":"bf47213134e10b4e0e157acd55bc2871e43888430316663ba36689069ee2b95a","source":{"kind":"arxiv","id":"2406.01631","version":2},"attestation_state":"computed","paper":{"title":"SUBER: An RL Environment with Simulated Human Behavior for Recommender Systems","license":"http://creativecommons.org/licenses/by-sa/4.0/","headline":"","cross_cats":["cs.LG"],"primary_cat":"cs.IR","authors_text":"Flint Xiaofeng Fan, Giorgio Piatti, Luca A. Lanzend\\\"orfer, Nathan Corecco, Roger Wattenhofer","submitted_at":"2024-06-01T11:56:08Z","abstract_excerpt":"Reinforcement learning (RL) has gained popularity in the realm of recommender systems due to its ability to optimize long-term rewards and guide users in discovering relevant content. However, the successful implementation of RL in recommender systems is challenging because of several factors, including the limited availability of online data for training on-policy methods. This scarcity requires expensive human interaction for online model training. Furthermore, the development of effective evaluation frameworks that accurately reflect the quality of models remains a fundamental challenge in "},"verification_status":{"content_addressed":true,"pith_receipt":true,"author_attested":false,"weak_author_claims":0,"strong_author_claims":0,"externally_anchored":false,"storage_verified":false,"citation_signatures":0,"replication_records":0,"graph_snapshot":true,"references_resolved":false,"formal_links_present":false},"canonical_record":{"source":{"id":"2406.01631","kind":"arxiv","version":2},"metadata":{"license":"http://creativecommons.org/licenses/by-sa/4.0/","primary_cat":"cs.IR","submitted_at":"2024-06-01T11:56:08Z","cross_cats_sorted":["cs.LG"],"title_canon_sha256":"34e5569acf5c5e7a8835ce26e024073060adc443fb6c5646ded40763aef85fec","abstract_canon_sha256":"bec7a7f4e9f95d75fc9b10f1583db57fd459a484798646369e129052ae3157f1"},"schema_version":"1.0"},"receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-07-05T08:57:04.588098Z","signature_b64":"J1BjGBe1E7ULhNJ9Mm3ts1B6e0Alm30vC6UMV7o6U9T7srV5dVmpHXmVgMvYiJxLepLnBN1E28ZRgYt3vVIbBA==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"bf47213134e10b4e0e157acd55bc2871e43888430316663ba36689069ee2b95a","last_reissued_at":"2026-07-05T08:57:04.587628Z","signature_status":"signed_v1","first_computed_at":"2026-07-05T08:57:04.587628Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"graph_snapshot":{"paper":{"title":"SUBER: An RL Environment with Simulated Human Behavior for Recommender Systems","license":"http://creativecommons.org/licenses/by-sa/4.0/","headline":"","cross_cats":["cs.LG"],"primary_cat":"cs.IR","authors_text":"Flint Xiaofeng Fan, Giorgio Piatti, Luca A. Lanzend\\\"orfer, Nathan Corecco, Roger Wattenhofer","submitted_at":"2024-06-01T11:56:08Z","abstract_excerpt":"Reinforcement learning (RL) has gained popularity in the realm of recommender systems due to its ability to optimize long-term rewards and guide users in discovering relevant content. However, the successful implementation of RL in recommender systems is challenging because of several factors, including the limited availability of online data for training on-policy methods. This scarcity requires expensive human interaction for online model training. Furthermore, the development of effective evaluation frameworks that accurately reflect the quality of models remains a fundamental challenge in "},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2406.01631","kind":"arxiv","version":2},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2406.01631/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"aliases":[{"alias_kind":"arxiv","alias_value":"2406.01631","created_at":"2026-07-05T08:57:04.587688+00:00"},{"alias_kind":"arxiv_version","alias_value":"2406.01631v2","created_at":"2026-07-05T08:57:04.587688+00:00"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2406.01631","created_at":"2026-07-05T08:57:04.587688+00:00"},{"alias_kind":"pith_short_12","alias_value":"X5DSCMJU4EFU","created_at":"2026-07-05T08:57:04.587688+00:00"},{"alias_kind":"pith_short_16","alias_value":"X5DSCMJU4EFU4DQV","created_at":"2026-07-05T08:57:04.587688+00:00"},{"alias_kind":"pith_short_8","alias_value":"X5DSCMJU","created_at":"2026-07-05T08:57:04.587688+00:00"}],"events":[],"event_summary":{},"paper_claims":[],"inbound_citations":{"count":1,"internal_anchor_count":0,"sample":[{"citing_arxiv_id":"2510.27157","citing_title":"A Survey on Generative Recommendation: Data, Model, and Tasks","ref_index":19,"is_internal_anchor":false}]},"formal_canon":{"evidence_count":0,"sample":[],"anchors":[]},"links":{"html":"https://pith.science/pith/X5DSCMJU4EFU4DQVPLGVLPBIOH","json":"https://pith.science/pith/X5DSCMJU4EFU4DQVPLGVLPBIOH.json","graph_json":"https://pith.science/api/pith-number/X5DSCMJU4EFU4DQVPLGVLPBIOH/graph.json","events_json":"https://pith.science/api/pith-number/X5DSCMJU4EFU4DQVPLGVLPBIOH/events.json","paper":"https://pith.science/paper/X5DSCMJU"},"agent_actions":{"view_html":"https://pith.science/pith/X5DSCMJU4EFU4DQVPLGVLPBIOH","download_json":"https://pith.science/pith/X5DSCMJU4EFU4DQVPLGVLPBIOH.json","view_paper":"https://pith.science/paper/X5DSCMJU","resolve_alias":"https://pith.science/api/pith-number/resolve?arxiv=2406.01631&json=true","fetch_graph":"https://pith.science/api/pith-number/X5DSCMJU4EFU4DQVPLGVLPBIOH/graph.json","fetch_events":"https://pith.science/api/pith-number/X5DSCMJU4EFU4DQVPLGVLPBIOH/events.json","actions":{"anchor_timestamp":"https://pith.science/pith/X5DSCMJU4EFU4DQVPLGVLPBIOH/action/timestamp_anchor","attest_storage":"https://pith.science/pith/X5DSCMJU4EFU4DQVPLGVLPBIOH/action/storage_attestation","attest_author":"https://pith.science/pith/X5DSCMJU4EFU4DQVPLGVLPBIOH/action/author_attestation","sign_citation":"https://pith.science/pith/X5DSCMJU4EFU4DQVPLGVLPBIOH/action/citation_signature","submit_replication":"https://pith.science/pith/X5DSCMJU4EFU4DQVPLGVLPBIOH/action/replication_record"}},"created_at":"2026-07-05T08:57:04.587688+00:00","updated_at":"2026-07-05T08:57:04.587688+00:00"}