{"record_type":"pith_number_record","schema_url":"https://pith.science/schemas/pith-number/v1.json","pith_number":"pith:2025:P36NT5Q65KBN3AKSMU4M2ZTO5B","short_pith_number":"pith:P36NT5Q6","schema_version":"1.0","canonical_sha256":"7efcd9f61eea82dd81526538cd666ee850212fb055482e3491c535848e6193a1","source":{"kind":"arxiv","id":"2509.10884","version":1},"attestation_state":"computed","paper":{"title":"Nav-R1: Reasoning and Navigation in Embodied Scenes","license":"http://creativecommons.org/licenses/by-nc-sa/4.0/","headline":"","cross_cats":["cs.CV"],"primary_cat":"cs.RO","authors_text":"Hao Tang, Qingxiang Liu, Ting Huang, Zeyu Zhang","submitted_at":"2025-09-13T16:31:03Z","abstract_excerpt":"Embodied navigation requires agents to integrate perception, reasoning, and action for robust interaction in complex 3D environments. Existing approaches often suffer from incoherent and unstable reasoning traces that hinder generalization across diverse environments, and difficulty balancing long-horizon semantic reasoning with low-latency control for real-time navigation. To address these challenges, we propose Nav-R1, an embodied foundation model that unifies reasoning in embodied environments. We first construct Nav-CoT-110K, a large-scale dataset of step-by-step Chains-of-Thought (CoT) fo"},"verification_status":{"content_addressed":true,"pith_receipt":true,"author_attested":false,"weak_author_claims":0,"strong_author_claims":0,"externally_anchored":false,"storage_verified":false,"citation_signatures":0,"replication_records":0,"graph_snapshot":true,"references_resolved":false,"formal_links_present":false},"canonical_record":{"source":{"id":"2509.10884","kind":"arxiv","version":1},"metadata":{"license":"http://creativecommons.org/licenses/by-nc-sa/4.0/","primary_cat":"cs.RO","submitted_at":"2025-09-13T16:31:03Z","cross_cats_sorted":["cs.CV"],"title_canon_sha256":"29bdb2dd677089727110f06e253d09e31aac60ddd3c3712c227aa28e35847a18","abstract_canon_sha256":"2a852cacc69f0139dd3b5b18043c97b72cfb924d4bd52a02766f5e00453c5209"},"schema_version":"1.0"},"receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-07-05T12:11:29.794069Z","signature_b64":"Eo6ikqzOXQ+iX/WhYEQRgmIYyorQeTCZAFTNVI6CFX4/XrXherjCqZYJnkte1C3O2ticP7AIBSTH2IdDQuAcBQ==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"7efcd9f61eea82dd81526538cd666ee850212fb055482e3491c535848e6193a1","last_reissued_at":"2026-07-05T12:11:29.793591Z","signature_status":"signed_v1","first_computed_at":"2026-07-05T12:11:29.793591Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"graph_snapshot":{"paper":{"title":"Nav-R1: Reasoning and Navigation in Embodied Scenes","license":"http://creativecommons.org/licenses/by-nc-sa/4.0/","headline":"","cross_cats":["cs.CV"],"primary_cat":"cs.RO","authors_text":"Hao Tang, Qingxiang Liu, Ting Huang, Zeyu Zhang","submitted_at":"2025-09-13T16:31:03Z","abstract_excerpt":"Embodied navigation requires agents to integrate perception, reasoning, and action for robust interaction in complex 3D environments. Existing approaches often suffer from incoherent and unstable reasoning traces that hinder generalization across diverse environments, and difficulty balancing long-horizon semantic reasoning with low-latency control for real-time navigation. To address these challenges, we propose Nav-R1, an embodied foundation model that unifies reasoning in embodied environments. We first construct Nav-CoT-110K, a large-scale dataset of step-by-step Chains-of-Thought (CoT) fo"},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2509.10884","kind":"arxiv","version":1},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2509.10884/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"aliases":[{"alias_kind":"arxiv","alias_value":"2509.10884","created_at":"2026-07-05T12:11:29.793649+00:00"},{"alias_kind":"arxiv_version","alias_value":"2509.10884v1","created_at":"2026-07-05T12:11:29.793649+00:00"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2509.10884","created_at":"2026-07-05T12:11:29.793649+00:00"},{"alias_kind":"pith_short_12","alias_value":"P36NT5Q65KBN","created_at":"2026-07-05T12:11:29.793649+00:00"},{"alias_kind":"pith_short_16","alias_value":"P36NT5Q65KBN3AKS","created_at":"2026-07-05T12:11:29.793649+00:00"},{"alias_kind":"pith_short_8","alias_value":"P36NT5Q6","created_at":"2026-07-05T12:11:29.793649+00:00"}],"events":[],"event_summary":{},"paper_claims":[],"inbound_citations":{"count":8,"internal_anchor_count":0,"sample":[{"citing_arxiv_id":"2606.23557","citing_title":"Dense Reward for Multi-View 3D Reasoning with Global Maps and Local Views","ref_index":32,"is_internal_anchor":false},{"citing_arxiv_id":"2606.01788","citing_title":"PlatonicNav: Unveiling Semantic Correspondence in Navigation with Platonic Topological Maps","ref_index":40,"is_internal_anchor":false},{"citing_arxiv_id":"2606.01621","citing_title":"Goal2Pixel: Grounding Goals to Pixels for Vision-Language Navigation","ref_index":17,"is_internal_anchor":false},{"citing_arxiv_id":"2605.22816","citing_title":"AwareVLN: Reasoning with Self-awareness for Vision-Language Navigation","ref_index":28,"is_internal_anchor":false},{"citing_arxiv_id":"2602.23058","citing_title":"GeoWorld: Geometric World Models","ref_index":44,"is_internal_anchor":false},{"citing_arxiv_id":"2511.17097","citing_title":"Progress-Think: Semantic Progress Reasoning for Vision-Language Navigation","ref_index":16,"is_internal_anchor":false},{"citing_arxiv_id":"2604.27620","citing_title":"SpaAct: Spatially-Activated Transition Learning with Curriculum Adaptation for Vision-Language Navigation","ref_index":44,"is_internal_anchor":false},{"citing_arxiv_id":"2604.17472","citing_title":"UniMesh: Unifying 3D Mesh Understanding and Generation","ref_index":25,"is_internal_anchor":false}]},"formal_canon":{"evidence_count":0,"sample":[],"anchors":[]},"links":{"html":"https://pith.science/pith/P36NT5Q65KBN3AKSMU4M2ZTO5B","json":"https://pith.science/pith/P36NT5Q65KBN3AKSMU4M2ZTO5B.json","graph_json":"https://pith.science/api/pith-number/P36NT5Q65KBN3AKSMU4M2ZTO5B/graph.json","events_json":"https://pith.science/api/pith-number/P36NT5Q65KBN3AKSMU4M2ZTO5B/events.json","paper":"https://pith.science/paper/P36NT5Q6"},"agent_actions":{"view_html":"https://pith.science/pith/P36NT5Q65KBN3AKSMU4M2ZTO5B","download_json":"https://pith.science/pith/P36NT5Q65KBN3AKSMU4M2ZTO5B.json","view_paper":"https://pith.science/paper/P36NT5Q6","resolve_alias":"https://pith.science/api/pith-number/resolve?arxiv=2509.10884&json=true","fetch_graph":"https://pith.science/api/pith-number/P36NT5Q65KBN3AKSMU4M2ZTO5B/graph.json","fetch_events":"https://pith.science/api/pith-number/P36NT5Q65KBN3AKSMU4M2ZTO5B/events.json","actions":{"anchor_timestamp":"https://pith.science/pith/P36NT5Q65KBN3AKSMU4M2ZTO5B/action/timestamp_anchor","attest_storage":"https://pith.science/pith/P36NT5Q65KBN3AKSMU4M2ZTO5B/action/storage_attestation","attest_author":"https://pith.science/pith/P36NT5Q65KBN3AKSMU4M2ZTO5B/action/author_attestation","sign_citation":"https://pith.science/pith/P36NT5Q65KBN3AKSMU4M2ZTO5B/action/citation_signature","submit_replication":"https://pith.science/pith/P36NT5Q65KBN3AKSMU4M2ZTO5B/action/replication_record"}},"created_at":"2026-07-05T12:11:29.793649+00:00","updated_at":"2026-07-05T12:11:29.793649+00:00"}