{"record_type":"pith_number_record","schema_url":"https://pith.science/schemas/pith-number/v1.json","pith_number":"pith:2024:3E4SWZW7C3NUOF4FEISKK44H6S","short_pith_number":"pith:3E4SWZW7","schema_version":"1.0","canonical_sha256":"d9392b66df16db4717852224a57387f4a335d4013e66528345c815c6739e409b","source":{"kind":"arxiv","id":"2408.10015","version":2},"attestation_state":"computed","paper":{"title":"Deterministic Policy Gradient Primal-Dual Methods for Continuous-Space Constrained MDPs","license":"http://creativecommons.org/licenses/by/4.0/","headline":"","cross_cats":["math.OC"],"primary_cat":"cs.AI","authors_text":"Alejandro Ribeiro, Antonio G. Marques, Dongsheng Ding, Sergio Rozada","submitted_at":"2024-08-19T14:11:04Z","abstract_excerpt":"We study the problem of computing deterministic optimal policies for constrained Markov decision processes (MDPs) with continuous state and action spaces, which are widely encountered in constrained dynamical systems. Designing deterministic policy gradient methods in continuous state and action spaces is particularly challenging due to the lack of enumerable state-action pairs and the adoption of deterministic policies, hindering the application of existing policy gradient methods. To this end, we develop a deterministic policy gradient primal-dual method to find an optimal deterministic poli"},"verification_status":{"content_addressed":true,"pith_receipt":true,"author_attested":false,"weak_author_claims":0,"strong_author_claims":0,"externally_anchored":false,"storage_verified":false,"citation_signatures":0,"replication_records":0,"graph_snapshot":true,"references_resolved":false,"formal_links_present":false},"canonical_record":{"source":{"id":"2408.10015","kind":"arxiv","version":2},"metadata":{"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.AI","submitted_at":"2024-08-19T14:11:04Z","cross_cats_sorted":["math.OC"],"title_canon_sha256":"f0d16af837ffc7fd8ec6f49141dfea5b34b3ba50feffaf635ea5684f4ee3dc61","abstract_canon_sha256":"8ff41d85a937ae9bc4172b61af7069a3eab4e21cc813bec80afbcc4fa14f144b"},"schema_version":"1.0"},"receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-07-05T10:44:15.562030Z","signature_b64":"dcQ1U+NjMHilslffauF51gJbrd/U94GPy/DSilp1u3U/S74YeF01S0/WLwmebJoP+M/H0KzltWfsrq61Ji8eDA==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"d9392b66df16db4717852224a57387f4a335d4013e66528345c815c6739e409b","last_reissued_at":"2026-07-05T10:44:15.561410Z","signature_status":"signed_v1","first_computed_at":"2026-07-05T10:44:15.561410Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"graph_snapshot":{"paper":{"title":"Deterministic Policy Gradient Primal-Dual Methods for Continuous-Space Constrained MDPs","license":"http://creativecommons.org/licenses/by/4.0/","headline":"","cross_cats":["math.OC"],"primary_cat":"cs.AI","authors_text":"Alejandro Ribeiro, Antonio G. Marques, Dongsheng Ding, Sergio Rozada","submitted_at":"2024-08-19T14:11:04Z","abstract_excerpt":"We study the problem of computing deterministic optimal policies for constrained Markov decision processes (MDPs) with continuous state and action spaces, which are widely encountered in constrained dynamical systems. Designing deterministic policy gradient methods in continuous state and action spaces is particularly challenging due to the lack of enumerable state-action pairs and the adoption of deterministic policies, hindering the application of existing policy gradient methods. To this end, we develop a deterministic policy gradient primal-dual method to find an optimal deterministic poli"},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2408.10015","kind":"arxiv","version":2},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2408.10015/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"aliases":[{"alias_kind":"arxiv","alias_value":"2408.10015","created_at":"2026-07-05T10:44:15.561493+00:00"},{"alias_kind":"arxiv_version","alias_value":"2408.10015v2","created_at":"2026-07-05T10:44:15.561493+00:00"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2408.10015","created_at":"2026-07-05T10:44:15.561493+00:00"},{"alias_kind":"pith_short_12","alias_value":"3E4SWZW7C3NU","created_at":"2026-07-05T10:44:15.561493+00:00"},{"alias_kind":"pith_short_16","alias_value":"3E4SWZW7C3NUOF4F","created_at":"2026-07-05T10:44:15.561493+00:00"},{"alias_kind":"pith_short_8","alias_value":"3E4SWZW7","created_at":"2026-07-05T10:44:15.561493+00:00"}],"events":[],"event_summary":{},"paper_claims":[],"inbound_citations":{"count":0,"internal_anchor_count":0,"sample":[]},"formal_canon":{"evidence_count":0,"sample":[],"anchors":[]},"links":{"html":"https://pith.science/pith/3E4SWZW7C3NUOF4FEISKK44H6S","json":"https://pith.science/pith/3E4SWZW7C3NUOF4FEISKK44H6S.json","graph_json":"https://pith.science/api/pith-number/3E4SWZW7C3NUOF4FEISKK44H6S/graph.json","events_json":"https://pith.science/api/pith-number/3E4SWZW7C3NUOF4FEISKK44H6S/events.json","paper":"https://pith.science/paper/3E4SWZW7"},"agent_actions":{"view_html":"https://pith.science/pith/3E4SWZW7C3NUOF4FEISKK44H6S","download_json":"https://pith.science/pith/3E4SWZW7C3NUOF4FEISKK44H6S.json","view_paper":"https://pith.science/paper/3E4SWZW7","resolve_alias":"https://pith.science/api/pith-number/resolve?arxiv=2408.10015&json=true","fetch_graph":"https://pith.science/api/pith-number/3E4SWZW7C3NUOF4FEISKK44H6S/graph.json","fetch_events":"https://pith.science/api/pith-number/3E4SWZW7C3NUOF4FEISKK44H6S/events.json","actions":{"anchor_timestamp":"https://pith.science/pith/3E4SWZW7C3NUOF4FEISKK44H6S/action/timestamp_anchor","attest_storage":"https://pith.science/pith/3E4SWZW7C3NUOF4FEISKK44H6S/action/storage_attestation","attest_author":"https://pith.science/pith/3E4SWZW7C3NUOF4FEISKK44H6S/action/author_attestation","sign_citation":"https://pith.science/pith/3E4SWZW7C3NUOF4FEISKK44H6S/action/citation_signature","submit_replication":"https://pith.science/pith/3E4SWZW7C3NUOF4FEISKK44H6S/action/replication_record"}},"created_at":"2026-07-05T10:44:15.561493+00:00","updated_at":"2026-07-05T10:44:15.561493+00:00"}