{"record_type":"pith_number_record","schema_url":"https://pith.science/schemas/pith-number/v1.json","pith_number":"pith:2025:Q7BELYXTG2I2BPKHJ5UJGMAZPR","short_pith_number":"pith:Q7BELYXT","schema_version":"1.0","canonical_sha256":"87c245e2f33691a0bd474f689330197c5d95954c2c0c44ab1e2c68342288dc6a","source":{"kind":"arxiv","id":"2507.02626","version":1},"attestation_state":"computed","paper":{"title":"VRAgent-R1: Boosting Video Recommendation with MLLM-based Agents via Reinforcement Learning","license":"http://creativecommons.org/licenses/by/4.0/","headline":"","cross_cats":[],"primary_cat":"cs.MM","authors_text":"Boyu Chen, Chengxiang Zhuo, Chenyun Yu, Lei Cheng, Ouyang Yi, Siran Chen, Yali Wang, Yuxiao Luo, Zang Li","submitted_at":"2025-07-03T13:52:24Z","abstract_excerpt":"Owing to powerful natural language processing and generative capabilities, large language model (LLM) agents have emerged as a promising solution for enhancing recommendation systems via user simulation. However, in the realm of video recommendation, existing studies predominantly resort to prompt-based simulation using frozen LLMs and encounter the intricate challenge of multimodal content understanding. This frequently results in suboptimal item modeling and user preference learning, thereby ultimately constraining recommendation performance. To address these challenges, we introduce VRAgent"},"verification_status":{"content_addressed":true,"pith_receipt":true,"author_attested":false,"weak_author_claims":0,"strong_author_claims":0,"externally_anchored":false,"storage_verified":false,"citation_signatures":0,"replication_records":0,"graph_snapshot":true,"references_resolved":false,"formal_links_present":false},"canonical_record":{"source":{"id":"2507.02626","kind":"arxiv","version":1},"metadata":{"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.MM","submitted_at":"2025-07-03T13:52:24Z","cross_cats_sorted":[],"title_canon_sha256":"8b6c9ca4062c576118c8f0ea4c392dd2a2c5fc129827c4903439393efc3843ba","abstract_canon_sha256":"a66c27575e0b8de2007e5c6ad1846f5f84a4fdb818d162a1ea4fed22a7e16e00"},"schema_version":"1.0"},"receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-07-05T11:31:31.958790Z","signature_b64":"tP8jY+zdqc4XjnVpdsMTPwJS63uwUcZkPF3FQoHIsZLsyntAOetaXmCraLx4+0D4V+T3mKn4gAGT+naEEWD6BA==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"87c245e2f33691a0bd474f689330197c5d95954c2c0c44ab1e2c68342288dc6a","last_reissued_at":"2026-07-05T11:31:31.958305Z","signature_status":"signed_v1","first_computed_at":"2026-07-05T11:31:31.958305Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"graph_snapshot":{"paper":{"title":"VRAgent-R1: Boosting Video Recommendation with MLLM-based Agents via Reinforcement Learning","license":"http://creativecommons.org/licenses/by/4.0/","headline":"","cross_cats":[],"primary_cat":"cs.MM","authors_text":"Boyu Chen, Chengxiang Zhuo, Chenyun Yu, Lei Cheng, Ouyang Yi, Siran Chen, Yali Wang, Yuxiao Luo, Zang Li","submitted_at":"2025-07-03T13:52:24Z","abstract_excerpt":"Owing to powerful natural language processing and generative capabilities, large language model (LLM) agents have emerged as a promising solution for enhancing recommendation systems via user simulation. However, in the realm of video recommendation, existing studies predominantly resort to prompt-based simulation using frozen LLMs and encounter the intricate challenge of multimodal content understanding. This frequently results in suboptimal item modeling and user preference learning, thereby ultimately constraining recommendation performance. To address these challenges, we introduce VRAgent"},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2507.02626","kind":"arxiv","version":1},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2507.02626/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"aliases":[{"alias_kind":"arxiv","alias_value":"2507.02626","created_at":"2026-07-05T11:31:31.958376+00:00"},{"alias_kind":"arxiv_version","alias_value":"2507.02626v1","created_at":"2026-07-05T11:31:31.958376+00:00"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2507.02626","created_at":"2026-07-05T11:31:31.958376+00:00"},{"alias_kind":"pith_short_12","alias_value":"Q7BELYXTG2I2","created_at":"2026-07-05T11:31:31.958376+00:00"},{"alias_kind":"pith_short_16","alias_value":"Q7BELYXTG2I2BPKH","created_at":"2026-07-05T11:31:31.958376+00:00"},{"alias_kind":"pith_short_8","alias_value":"Q7BELYXT","created_at":"2026-07-05T11:31:31.958376+00:00"}],"events":[],"event_summary":{},"paper_claims":[],"inbound_citations":{"count":1,"internal_anchor_count":0,"sample":[{"citing_arxiv_id":"2606.13289","citing_title":"HYDRA-X: Native Unified Multimodal Models with Holistic Visual Tokenizers","ref_index":125,"is_internal_anchor":false}]},"formal_canon":{"evidence_count":0,"sample":[],"anchors":[]},"links":{"html":"https://pith.science/pith/Q7BELYXTG2I2BPKHJ5UJGMAZPR","json":"https://pith.science/pith/Q7BELYXTG2I2BPKHJ5UJGMAZPR.json","graph_json":"https://pith.science/api/pith-number/Q7BELYXTG2I2BPKHJ5UJGMAZPR/graph.json","events_json":"https://pith.science/api/pith-number/Q7BELYXTG2I2BPKHJ5UJGMAZPR/events.json","paper":"https://pith.science/paper/Q7BELYXT"},"agent_actions":{"view_html":"https://pith.science/pith/Q7BELYXTG2I2BPKHJ5UJGMAZPR","download_json":"https://pith.science/pith/Q7BELYXTG2I2BPKHJ5UJGMAZPR.json","view_paper":"https://pith.science/paper/Q7BELYXT","resolve_alias":"https://pith.science/api/pith-number/resolve?arxiv=2507.02626&json=true","fetch_graph":"https://pith.science/api/pith-number/Q7BELYXTG2I2BPKHJ5UJGMAZPR/graph.json","fetch_events":"https://pith.science/api/pith-number/Q7BELYXTG2I2BPKHJ5UJGMAZPR/events.json","actions":{"anchor_timestamp":"https://pith.science/pith/Q7BELYXTG2I2BPKHJ5UJGMAZPR/action/timestamp_anchor","attest_storage":"https://pith.science/pith/Q7BELYXTG2I2BPKHJ5UJGMAZPR/action/storage_attestation","attest_author":"https://pith.science/pith/Q7BELYXTG2I2BPKHJ5UJGMAZPR/action/author_attestation","sign_citation":"https://pith.science/pith/Q7BELYXTG2I2BPKHJ5UJGMAZPR/action/citation_signature","submit_replication":"https://pith.science/pith/Q7BELYXTG2I2BPKHJ5UJGMAZPR/action/replication_record"}},"created_at":"2026-07-05T11:31:31.958376+00:00","updated_at":"2026-07-05T11:31:31.958376+00:00"}