{"bundle_type":"pith_open_graph_bundle","bundle_version":"1.0","pith_number":"pith:2020:LI7RYNUBFFGLDX4N56CF24G7ZG","short_pith_number":"pith:LI7RYNUB","canonical_record":{"source":{"id":"2008.03525","kind":"arxiv","version":1},"metadata":{"license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","primary_cat":"cs.LG","submitted_at":"2020-08-08T13:43:06Z","cross_cats_sorted":["cs.IT","cs.RO","math.IT","stat.ML"],"title_canon_sha256":"2146b2400a76d9a7e0a113c9dd62ddec683f1a87e8c6bdf3cbedae868c9da19c","abstract_canon_sha256":"3d8fec8b0b2db9cb0c482c0d6572c3213c9802f376690b92c8aabd2369790ae9"},"schema_version":"1.0"},"canonical_sha256":"5a3f1c3681294cb1df8def845d70dfc9835cc54a88a67b949e7c265ad96f44b6","source":{"kind":"arxiv","id":"2008.03525","version":1},"source_aliases":[{"alias_kind":"arxiv","alias_value":"2008.03525","created_at":"2026-07-05T01:25:42Z"},{"alias_kind":"arxiv_version","alias_value":"2008.03525v1","created_at":"2026-07-05T01:25:42Z"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2008.03525","created_at":"2026-07-05T01:25:42Z"},{"alias_kind":"pith_short_12","alias_value":"LI7RYNUBFFGL","created_at":"2026-07-05T01:25:42Z"},{"alias_kind":"pith_short_16","alias_value":"LI7RYNUBFFGLDX4N","created_at":"2026-07-05T01:25:42Z"},{"alias_kind":"pith_short_8","alias_value":"LI7RYNUB","created_at":"2026-07-05T01:25:42Z"}],"events":[{"event_type":"record_created","subject_pith_number":"pith:2020:LI7RYNUBFFGLDX4N56CF24G7ZG","target":"record","payload":{"canonical_record":{"source":{"id":"2008.03525","kind":"arxiv","version":1},"metadata":{"license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","primary_cat":"cs.LG","submitted_at":"2020-08-08T13:43:06Z","cross_cats_sorted":["cs.IT","cs.RO","math.IT","stat.ML"],"title_canon_sha256":"2146b2400a76d9a7e0a113c9dd62ddec683f1a87e8c6bdf3cbedae868c9da19c","abstract_canon_sha256":"3d8fec8b0b2db9cb0c482c0d6572c3213c9802f376690b92c8aabd2369790ae9"},"schema_version":"1.0"},"canonical_sha256":"5a3f1c3681294cb1df8def845d70dfc9835cc54a88a67b949e7c265ad96f44b6","receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-07-05T01:25:42.747348Z","signature_b64":"VrBWrDSJp3/nPBvGo2/Cp5IyiZ9NTMNhIJMdtW13DLuSRQuf6wp5ZWZec8prkxAK5iv9UA0wNPI2KxRfEYgUBA==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"5a3f1c3681294cb1df8def845d70dfc9835cc54a88a67b949e7c265ad96f44b6","last_reissued_at":"2026-07-05T01:25:42.746907Z","signature_status":"signed_v1","first_computed_at":"2026-07-05T01:25:42.746907Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"source_kind":"arxiv","source_id":"2008.03525","source_version":1,"attestation_state":"computed"},"signer":{"signer_id":"pith.science","signer_type":"pith_registry","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"created_at":"2026-07-05T01:25:42Z","supersedes":[],"prev_event":null,"signature":{"signature_status":"signed_v1","algorithm":"ed25519","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signature_b64":"dP7Rt9YbxRpo3Msp7h8/XQX3/f0LCzkWIxQlVz+9fStqzxgby5q10T49ijgBUVSJzcp8bDIvoRkViX1MVVKtCg==","signed_message":"open_graph_event_sha256_bytes","signed_at":"2026-08-12T15:08:11.894744Z"},"content_sha256":"00ded035f801fea347fc22d027653dc3a6de058a1db43b40066ceccb568e3611","schema_version":"1.0","event_id":"sha256:00ded035f801fea347fc22d027653dc3a6de058a1db43b40066ceccb568e3611"},{"event_type":"graph_snapshot","subject_pith_number":"pith:2020:LI7RYNUBFFGLDX4N56CF24G7ZG","target":"graph","payload":{"graph_snapshot":{"paper":{"title":"Non-Adversarial Imitation Learning and its Connections to Adversarial Methods","license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","headline":"","cross_cats":["cs.IT","cs.RO","math.IT","stat.ML"],"primary_cat":"cs.LG","authors_text":"Gerhard Neumann, Oleg Arenz","submitted_at":"2020-08-08T13:43:06Z","abstract_excerpt":"Many modern methods for imitation learning and inverse reinforcement learning, such as GAIL or AIRL, are based on an adversarial formulation. These methods apply GANs to match the expert's distribution over states and actions with the implicit state-action distribution induced by the agent's policy. However, by framing imitation learning as a saddle point problem, adversarial methods can suffer from unstable optimization, and convergence can only be shown for small policy updates. We address these problems by proposing a framework for non-adversarial imitation learning. The resulting algorithm"},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2008.03525","kind":"arxiv","version":1},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2008.03525/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"verdict_id":null},"signer":{"signer_id":"pith.science","signer_type":"pith_registry","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"created_at":"2026-07-05T01:25:42Z","supersedes":[],"prev_event":null,"signature":{"signature_status":"signed_v1","algorithm":"ed25519","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signature_b64":"1bfpdgpKTnK4vA1YxNrgmXvIkbQngJBJVvyjFQOwoJWbh2UNdYM8NvCqZblkKZyaX7ufgjiCUCf8K7FKkMnQBQ==","signed_message":"open_graph_event_sha256_bytes","signed_at":"2026-08-12T15:08:11.895362Z"},"content_sha256":"d3e56d8c1585272abb1c4d655f5501c72ca0baf65b7061799fcc2ddb12f69727","schema_version":"1.0","event_id":"sha256:d3e56d8c1585272abb1c4d655f5501c72ca0baf65b7061799fcc2ddb12f69727"}],"timestamp_proofs":[],"mirror_hints":[{"mirror_type":"https","name":"Pith Resolver","base_url":"https://pith.science","bundle_url":"https://pith.science/pith/LI7RYNUBFFGLDX4N56CF24G7ZG/bundle.json","state_url":"https://pith.science/pith/LI7RYNUBFFGLDX4N56CF24G7ZG/state.json","well_known_bundle_url":"https://pith.science/.well-known/pith/LI7RYNUBFFGLDX4N56CF24G7ZG/bundle.json","status":"primary"}],"public_keys":[{"key_id":"pith-v1-2026-05","algorithm":"ed25519","format":"raw","public_key_b64":"stVStoiQhXFxp4s2pdzPNoqVNBMojDU/fJ2db5S3CbM=","public_key_hex":"b2d552b68890857171a78b36a5dccf368a953413288c353f7c9d9d6f94b709b3","fingerprint_sha256_b32_first128bits":"RVFV5Z2OI2J3ZUO7ERDEBCYNKS","fingerprint_sha256_hex":"8d4b5ee74e4693bcd1df2446408b0d54","rotates_at":null,"url":"https://pith.science/pith-signing-key.json","notes":"Pith uses this Ed25519 key to sign canonical record SHA-256 digests. Verify with: ed25519_verify(public_key, message=canonical_sha256_bytes, signature=base64decode(signature_b64))."}],"merge_version":"pith-open-graph-merge-v1","built_at":"2026-08-12T15:08:11Z","links":{"resolver":"https://pith.science/pith/LI7RYNUBFFGLDX4N56CF24G7ZG","bundle":"https://pith.science/pith/LI7RYNUBFFGLDX4N56CF24G7ZG/bundle.json","state":"https://pith.science/pith/LI7RYNUBFFGLDX4N56CF24G7ZG/state.json","well_known_bundle":"https://pith.science/.well-known/pith/LI7RYNUBFFGLDX4N56CF24G7ZG/bundle.json"},"state":{"state_type":"pith_open_graph_state","state_version":"1.0","pith_number":"pith:2020:LI7RYNUBFFGLDX4N56CF24G7ZG","merge_version":"pith-open-graph-merge-v1","event_count":2,"valid_event_count":2,"invalid_event_count":0,"equivocation_count":0,"current":{"canonical_record":{"metadata":{"abstract_canon_sha256":"3d8fec8b0b2db9cb0c482c0d6572c3213c9802f376690b92c8aabd2369790ae9","cross_cats_sorted":["cs.IT","cs.RO","math.IT","stat.ML"],"license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","primary_cat":"cs.LG","submitted_at":"2020-08-08T13:43:06Z","title_canon_sha256":"2146b2400a76d9a7e0a113c9dd62ddec683f1a87e8c6bdf3cbedae868c9da19c"},"schema_version":"1.0","source":{"id":"2008.03525","kind":"arxiv","version":1}},"source_aliases":[{"alias_kind":"arxiv","alias_value":"2008.03525","created_at":"2026-07-05T01:25:42Z"},{"alias_kind":"arxiv_version","alias_value":"2008.03525v1","created_at":"2026-07-05T01:25:42Z"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2008.03525","created_at":"2026-07-05T01:25:42Z"},{"alias_kind":"pith_short_12","alias_value":"LI7RYNUBFFGL","created_at":"2026-07-05T01:25:42Z"},{"alias_kind":"pith_short_16","alias_value":"LI7RYNUBFFGLDX4N","created_at":"2026-07-05T01:25:42Z"},{"alias_kind":"pith_short_8","alias_value":"LI7RYNUB","created_at":"2026-07-05T01:25:42Z"}],"graph_snapshots":[{"event_id":"sha256:d3e56d8c1585272abb1c4d655f5501c72ca0baf65b7061799fcc2ddb12f69727","target":"graph","created_at":"2026-07-05T01:25:42Z","signer":{"key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signer_id":"pith.science","signer_type":"pith_registry"},"payload":{"graph_snapshot":{"author_claims":{"count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","strong_count":0},"builder_version":"pith-number-builder-2026-05-17-v1","claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"integrity":{"available":true,"clean":true,"detectors_run":[],"endpoint":"/pith/2008.03525/integrity.json","findings":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938","summary":{"advisory":0,"by_detector":{},"critical":0,"informational":0}},"paper":{"abstract_excerpt":"Many modern methods for imitation learning and inverse reinforcement learning, such as GAIL or AIRL, are based on an adversarial formulation. These methods apply GANs to match the expert's distribution over states and actions with the implicit state-action distribution induced by the agent's policy. However, by framing imitation learning as a saddle point problem, adversarial methods can suffer from unstable optimization, and convergence can only be shown for small policy updates. We address these problems by proposing a framework for non-adversarial imitation learning. The resulting algorithm","authors_text":"Gerhard Neumann, Oleg Arenz","cross_cats":["cs.IT","cs.RO","math.IT","stat.ML"],"headline":"","license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","primary_cat":"cs.LG","submitted_at":"2020-08-08T13:43:06Z","title":"Non-Adversarial Imitation Learning and its Connections to Adversarial Methods"},"references":{"count":0,"internal_anchors":0,"resolved_work":0,"sample":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2008.03525","kind":"arxiv","version":1},"verdict":{"created_at":null,"id":null,"model_set":{},"one_line_summary":"","pipeline_version":null,"pith_extraction_headline":"","strongest_claim":"","weakest_assumption":""}},"verdict_id":null}}],"author_attestations":[],"timestamp_anchors":[],"storage_attestations":[],"citation_signatures":[],"replication_records":[],"corrections":[],"mirror_hints":[],"record_created":{"event_id":"sha256:00ded035f801fea347fc22d027653dc3a6de058a1db43b40066ceccb568e3611","target":"record","created_at":"2026-07-05T01:25:42Z","signer":{"key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signer_id":"pith.science","signer_type":"pith_registry"},"payload":{"attestation_state":"computed","canonical_record":{"metadata":{"abstract_canon_sha256":"3d8fec8b0b2db9cb0c482c0d6572c3213c9802f376690b92c8aabd2369790ae9","cross_cats_sorted":["cs.IT","cs.RO","math.IT","stat.ML"],"license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","primary_cat":"cs.LG","submitted_at":"2020-08-08T13:43:06Z","title_canon_sha256":"2146b2400a76d9a7e0a113c9dd62ddec683f1a87e8c6bdf3cbedae868c9da19c"},"schema_version":"1.0","source":{"id":"2008.03525","kind":"arxiv","version":1}},"canonical_sha256":"5a3f1c3681294cb1df8def845d70dfc9835cc54a88a67b949e7c265ad96f44b6","receipt":{"algorithm":"ed25519","builder_version":"pith-number-builder-2026-05-17-v1","canonical_sha256":"5a3f1c3681294cb1df8def845d70dfc9835cc54a88a67b949e7c265ad96f44b6","first_computed_at":"2026-07-05T01:25:42.746907Z","key_id":"pith-v1-2026-05","kind":"pith_receipt","last_reissued_at":"2026-07-05T01:25:42.746907Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","receipt_version":"0.3","signature_b64":"VrBWrDSJp3/nPBvGo2/Cp5IyiZ9NTMNhIJMdtW13DLuSRQuf6wp5ZWZec8prkxAK5iv9UA0wNPI2KxRfEYgUBA==","signature_status":"signed_v1","signed_at":"2026-07-05T01:25:42.747348Z","signed_message":"canonical_sha256_bytes"},"source_id":"2008.03525","source_kind":"arxiv","source_version":1}}},"equivocations":[],"invalid_events":[],"applied_event_ids":["sha256:00ded035f801fea347fc22d027653dc3a6de058a1db43b40066ceccb568e3611","sha256:d3e56d8c1585272abb1c4d655f5501c72ca0baf65b7061799fcc2ddb12f69727"],"state_sha256":"63f8784d94b7b333c8247bd5ecaa097052a0323409a48ad44339f916c5e09bb3"},"bundle_signature":{"signature_status":"signed_v1","algorithm":"ed25519","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signature_b64":"kbPM2ZSgQ2X2QKobsbOrtPskwctlE7oksqd7XNcSnDaSx5GOhBd5MGCZ+/4gYebZInmX1kQzG0vbHCO2h10MBQ==","signed_message":"bundle_sha256_bytes","signed_at":"2026-08-12T15:08:11.900383Z","bundle_sha256":"9d8f357d034147f232bc7cb1c8793813500043b08bff63583522fb87de985e98"}}