{"record_type":"pith_number_record","schema_url":"https://pith.science/schemas/pith-number/v1.json","pith_number":"pith:2025:IQS3E5EBWYJYMB5SFFWAGRQQCP","short_pith_number":"pith:IQS3E5EB","schema_version":"1.0","canonical_sha256":"4425b27481b6138607b2296c03461013c82863191b11fbf0abd91a2380b10db5","source":{"kind":"arxiv","id":"2508.05838","version":1},"attestation_state":"computed","paper":{"title":"Integrating Vision Foundation Models with Reinforcement Learning for Enhanced Object Interaction","license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","headline":"","cross_cats":["cs.AI","cs.CV","cs.LG","cs.SY","eess.SY"],"primary_cat":"cs.RO","authors_text":"Ahmad Farooq, Kamran Iqbal","submitted_at":"2025-08-07T20:29:01Z","abstract_excerpt":"This paper presents a novel approach that integrates vision foundation models with reinforcement learning to enhance object interaction capabilities in simulated environments. By combining the Segment Anything Model (SAM) and YOLOv5 with a Proximal Policy Optimization (PPO) agent operating in the AI2-THOR simulation environment, we enable the agent to perceive and interact with objects more effectively. Our comprehensive experiments, conducted across four diverse indoor kitchen settings, demonstrate significant improvements in object interaction success rates and navigation efficiency compared"},"verification_status":{"content_addressed":true,"pith_receipt":true,"author_attested":false,"weak_author_claims":0,"strong_author_claims":0,"externally_anchored":false,"storage_verified":false,"citation_signatures":0,"replication_records":0,"graph_snapshot":true,"references_resolved":false,"formal_links_present":false},"canonical_record":{"source":{"id":"2508.05838","kind":"arxiv","version":1},"metadata":{"license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","primary_cat":"cs.RO","submitted_at":"2025-08-07T20:29:01Z","cross_cats_sorted":["cs.AI","cs.CV","cs.LG","cs.SY","eess.SY"],"title_canon_sha256":"07223d5a02d6f47a5d95133548aefdfe707a04abc917cc81c666b53337ebb651","abstract_canon_sha256":"73ab430515ced67d3be1762004df08b292b7318612c51e43a40c1db00f6d54c8"},"schema_version":"1.0"},"receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-07-05T11:50:44.166947Z","signature_b64":"cPEmrGd/iIMnhas66boMNPuiZU/i3dodDnZvr2XOtyrSkYQfLB0fjUTrYY361lXsdpnQ07tZxy5Fn01ZPa4cAg==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"4425b27481b6138607b2296c03461013c82863191b11fbf0abd91a2380b10db5","last_reissued_at":"2026-07-05T11:50:44.166503Z","signature_status":"signed_v1","first_computed_at":"2026-07-05T11:50:44.166503Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"graph_snapshot":{"paper":{"title":"Integrating Vision Foundation Models with Reinforcement Learning for Enhanced Object Interaction","license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","headline":"","cross_cats":["cs.AI","cs.CV","cs.LG","cs.SY","eess.SY"],"primary_cat":"cs.RO","authors_text":"Ahmad Farooq, Kamran Iqbal","submitted_at":"2025-08-07T20:29:01Z","abstract_excerpt":"This paper presents a novel approach that integrates vision foundation models with reinforcement learning to enhance object interaction capabilities in simulated environments. By combining the Segment Anything Model (SAM) and YOLOv5 with a Proximal Policy Optimization (PPO) agent operating in the AI2-THOR simulation environment, we enable the agent to perceive and interact with objects more effectively. Our comprehensive experiments, conducted across four diverse indoor kitchen settings, demonstrate significant improvements in object interaction success rates and navigation efficiency compared"},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2508.05838","kind":"arxiv","version":1},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2508.05838/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"aliases":[{"alias_kind":"arxiv","alias_value":"2508.05838","created_at":"2026-07-05T11:50:44.166555+00:00"},{"alias_kind":"arxiv_version","alias_value":"2508.05838v1","created_at":"2026-07-05T11:50:44.166555+00:00"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2508.05838","created_at":"2026-07-05T11:50:44.166555+00:00"},{"alias_kind":"pith_short_12","alias_value":"IQS3E5EBWYJY","created_at":"2026-07-05T11:50:44.166555+00:00"},{"alias_kind":"pith_short_16","alias_value":"IQS3E5EBWYJYMB5S","created_at":"2026-07-05T11:50:44.166555+00:00"},{"alias_kind":"pith_short_8","alias_value":"IQS3E5EB","created_at":"2026-07-05T11:50:44.166555+00:00"}],"events":[],"event_summary":{},"paper_claims":[],"inbound_citations":{"count":0,"internal_anchor_count":0,"sample":[]},"formal_canon":{"evidence_count":0,"sample":[],"anchors":[]},"links":{"html":"https://pith.science/pith/IQS3E5EBWYJYMB5SFFWAGRQQCP","json":"https://pith.science/pith/IQS3E5EBWYJYMB5SFFWAGRQQCP.json","graph_json":"https://pith.science/api/pith-number/IQS3E5EBWYJYMB5SFFWAGRQQCP/graph.json","events_json":"https://pith.science/api/pith-number/IQS3E5EBWYJYMB5SFFWAGRQQCP/events.json","paper":"https://pith.science/paper/IQS3E5EB"},"agent_actions":{"view_html":"https://pith.science/pith/IQS3E5EBWYJYMB5SFFWAGRQQCP","download_json":"https://pith.science/pith/IQS3E5EBWYJYMB5SFFWAGRQQCP.json","view_paper":"https://pith.science/paper/IQS3E5EB","resolve_alias":"https://pith.science/api/pith-number/resolve?arxiv=2508.05838&json=true","fetch_graph":"https://pith.science/api/pith-number/IQS3E5EBWYJYMB5SFFWAGRQQCP/graph.json","fetch_events":"https://pith.science/api/pith-number/IQS3E5EBWYJYMB5SFFWAGRQQCP/events.json","actions":{"anchor_timestamp":"https://pith.science/pith/IQS3E5EBWYJYMB5SFFWAGRQQCP/action/timestamp_anchor","attest_storage":"https://pith.science/pith/IQS3E5EBWYJYMB5SFFWAGRQQCP/action/storage_attestation","attest_author":"https://pith.science/pith/IQS3E5EBWYJYMB5SFFWAGRQQCP/action/author_attestation","sign_citation":"https://pith.science/pith/IQS3E5EBWYJYMB5SFFWAGRQQCP/action/citation_signature","submit_replication":"https://pith.science/pith/IQS3E5EBWYJYMB5SFFWAGRQQCP/action/replication_record"}},"created_at":"2026-07-05T11:50:44.166555+00:00","updated_at":"2026-07-05T11:50:44.166555+00:00"}