{"bundle_type":"pith_open_graph_bundle","bundle_version":"1.0","pith_number":"pith:2017:3FSSMXMGX4ZFOIRPDETS4ZP5BU","short_pith_number":"pith:3FSSMXMG","canonical_record":{"source":{"id":"1711.00804","kind":"arxiv","version":2},"metadata":{"license":"http://creativecommons.org/licenses/by-sa/4.0/","primary_cat":"cs.SD","submitted_at":"2017-11-02T16:32:23Z","cross_cats_sorted":["cs.AI","cs.IR","eess.AS"],"title_canon_sha256":"c4ad4808dd52768a3084502c0e0f0944705131684d435f689b5a38e7a71241b0","abstract_canon_sha256":"754b0457dfa656e501eb24bfe2a98d7230b8069a9407fd71979e5f6ae1993657"},"schema_version":"1.0"},"canonical_sha256":"d965265d86bf3257222f19272e65fd0d3c0f961984e446fefaebf16738df53e0","source":{"kind":"arxiv","id":"1711.00804","version":2},"source_aliases":[{"alias_kind":"arxiv","alias_value":"1711.00804","created_at":"2026-05-18T00:19:16Z"},{"alias_kind":"arxiv_version","alias_value":"1711.00804v2","created_at":"2026-05-18T00:19:16Z"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.1711.00804","created_at":"2026-05-18T00:19:16Z"},{"alias_kind":"pith_short_12","alias_value":"3FSSMXMGX4ZF","created_at":"2026-05-18T12:30:58Z"},{"alias_kind":"pith_short_16","alias_value":"3FSSMXMGX4ZFOIRP","created_at":"2026-05-18T12:30:58Z"},{"alias_kind":"pith_short_8","alias_value":"3FSSMXMG","created_at":"2026-05-18T12:30:58Z"}],"events":[{"event_type":"record_created","subject_pith_number":"pith:2017:3FSSMXMGX4ZFOIRPDETS4ZP5BU","target":"record","payload":{"canonical_record":{"source":{"id":"1711.00804","kind":"arxiv","version":2},"metadata":{"license":"http://creativecommons.org/licenses/by-sa/4.0/","primary_cat":"cs.SD","submitted_at":"2017-11-02T16:32:23Z","cross_cats_sorted":["cs.AI","cs.IR","eess.AS"],"title_canon_sha256":"c4ad4808dd52768a3084502c0e0f0944705131684d435f689b5a38e7a71241b0","abstract_canon_sha256":"754b0457dfa656e501eb24bfe2a98d7230b8069a9407fd71979e5f6ae1993657"},"schema_version":"1.0"},"canonical_sha256":"d965265d86bf3257222f19272e65fd0d3c0f961984e446fefaebf16738df53e0","receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-05-18T00:19:16.889999Z","signature_b64":"PGVkvZ1KOx83jV0EF/wyMLMVo1qIHENwp3ellUXeoyuMnOoasb6CDFEPdfNloT6K+OgE0ztmckaRhkwmCNecDQ==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"d965265d86bf3257222f19272e65fd0d3c0f961984e446fefaebf16738df53e0","last_reissued_at":"2026-05-18T00:19:16.889444Z","signature_status":"signed_v1","first_computed_at":"2026-05-18T00:19:16.889444Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"source_kind":"arxiv","source_id":"1711.00804","source_version":2,"attestation_state":"computed"},"signer":{"signer_id":"pith.science","signer_type":"pith_registry","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"created_at":"2026-05-18T00:19:16Z","supersedes":[],"prev_event":null,"signature":{"signature_status":"signed_v1","algorithm":"ed25519","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signature_b64":"tdIzVQ1ZKR7hHQ8WvTlEs9wxyKmpoBLZRJMDIdYTS0uKKoE3Rgg2gcUW6sEv2f2ZqY6oF+MbO1nOV5WR/3WYBw==","signed_message":"open_graph_event_sha256_bytes","signed_at":"2026-05-30T20:21:30.575944Z"},"content_sha256":"1f9f245ef9d0cea3c3301ed95ea5435ba398daae6bc4f532f84472c88c6c71aa","schema_version":"1.0","event_id":"sha256:1f9f245ef9d0cea3c3301ed95ea5435ba398daae6bc4f532f84472c88c6c71aa"},{"event_type":"graph_snapshot","subject_pith_number":"pith:2017:3FSSMXMGX4ZFOIRPDETS4ZP5BU","target":"graph","payload":{"graph_snapshot":{"paper":{"title":"Framework for evaluation of sound event detection in web videos","license":"http://creativecommons.org/licenses/by-sa/4.0/","headline":"","cross_cats":["cs.AI","cs.IR","eess.AS"],"primary_cat":"cs.SD","authors_text":"Ankit Shah, Anurag Kumar, Benjamin Elizalde, Bhiksha Raj, Rohan Badlani","submitted_at":"2017-11-02T16:32:23Z","abstract_excerpt":"The largest source of sound events is web videos. Most videos lack sound event labels at segment level, however, a significant number of them do respond to text queries, from a match found using metadata by search engines. In this paper we explore the extent to which a search query can be used as the true label for detection of sound events in videos. We present a framework for large-scale sound event recognition on web videos. The framework crawls videos using search queries corresponding to 78 sound event labels drawn from three datasets. The datasets are used to train three classifiers, and"},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"1711.00804","kind":"arxiv","version":2},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"verdict_id":null},"signer":{"signer_id":"pith.science","signer_type":"pith_registry","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"created_at":"2026-05-18T00:19:16Z","supersedes":[],"prev_event":null,"signature":{"signature_status":"signed_v1","algorithm":"ed25519","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signature_b64":"i6oOnH8XLth2I7NRrhXmQNMY/E3kVW1wQ/61k49fmeE8wplX6s5sTVoa5js4zTf2iI/id120xrDGpL1z5/CDBA==","signed_message":"open_graph_event_sha256_bytes","signed_at":"2026-05-30T20:21:30.576689Z"},"content_sha256":"48560c1aa1c61713279ff31da3a9f245235e5def997b85e9f9bd8527399ee226","schema_version":"1.0","event_id":"sha256:48560c1aa1c61713279ff31da3a9f245235e5def997b85e9f9bd8527399ee226"}],"timestamp_proofs":[],"mirror_hints":[{"mirror_type":"https","name":"Pith Resolver","base_url":"https://pith.science","bundle_url":"https://pith.science/pith/3FSSMXMGX4ZFOIRPDETS4ZP5BU/bundle.json","state_url":"https://pith.science/pith/3FSSMXMGX4ZFOIRPDETS4ZP5BU/state.json","well_known_bundle_url":"https://pith.science/.well-known/pith/3FSSMXMGX4ZFOIRPDETS4ZP5BU/bundle.json","status":"primary"}],"public_keys":[{"key_id":"pith-v1-2026-05","algorithm":"ed25519","format":"raw","public_key_b64":"stVStoiQhXFxp4s2pdzPNoqVNBMojDU/fJ2db5S3CbM=","public_key_hex":"b2d552b68890857171a78b36a5dccf368a953413288c353f7c9d9d6f94b709b3","fingerprint_sha256_b32_first128bits":"RVFV5Z2OI2J3ZUO7ERDEBCYNKS","fingerprint_sha256_hex":"8d4b5ee74e4693bcd1df2446408b0d54","rotates_at":null,"url":"https://pith.science/pith-signing-key.json","notes":"Pith uses this Ed25519 key to sign canonical record SHA-256 digests. Verify with: ed25519_verify(public_key, message=canonical_sha256_bytes, signature=base64decode(signature_b64))."}],"merge_version":"pith-open-graph-merge-v1","built_at":"2026-05-30T20:21:30Z","links":{"resolver":"https://pith.science/pith/3FSSMXMGX4ZFOIRPDETS4ZP5BU","bundle":"https://pith.science/pith/3FSSMXMGX4ZFOIRPDETS4ZP5BU/bundle.json","state":"https://pith.science/pith/3FSSMXMGX4ZFOIRPDETS4ZP5BU/state.json","well_known_bundle":"https://pith.science/.well-known/pith/3FSSMXMGX4ZFOIRPDETS4ZP5BU/bundle.json"},"state":{"state_type":"pith_open_graph_state","state_version":"1.0","pith_number":"pith:2017:3FSSMXMGX4ZFOIRPDETS4ZP5BU","merge_version":"pith-open-graph-merge-v1","event_count":2,"valid_event_count":2,"invalid_event_count":0,"equivocation_count":0,"current":{"canonical_record":{"metadata":{"abstract_canon_sha256":"754b0457dfa656e501eb24bfe2a98d7230b8069a9407fd71979e5f6ae1993657","cross_cats_sorted":["cs.AI","cs.IR","eess.AS"],"license":"http://creativecommons.org/licenses/by-sa/4.0/","primary_cat":"cs.SD","submitted_at":"2017-11-02T16:32:23Z","title_canon_sha256":"c4ad4808dd52768a3084502c0e0f0944705131684d435f689b5a38e7a71241b0"},"schema_version":"1.0","source":{"id":"1711.00804","kind":"arxiv","version":2}},"source_aliases":[{"alias_kind":"arxiv","alias_value":"1711.00804","created_at":"2026-05-18T00:19:16Z"},{"alias_kind":"arxiv_version","alias_value":"1711.00804v2","created_at":"2026-05-18T00:19:16Z"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.1711.00804","created_at":"2026-05-18T00:19:16Z"},{"alias_kind":"pith_short_12","alias_value":"3FSSMXMGX4ZF","created_at":"2026-05-18T12:30:58Z"},{"alias_kind":"pith_short_16","alias_value":"3FSSMXMGX4ZFOIRP","created_at":"2026-05-18T12:30:58Z"},{"alias_kind":"pith_short_8","alias_value":"3FSSMXMG","created_at":"2026-05-18T12:30:58Z"}],"graph_snapshots":[{"event_id":"sha256:48560c1aa1c61713279ff31da3a9f245235e5def997b85e9f9bd8527399ee226","target":"graph","created_at":"2026-05-18T00:19:16Z","signer":{"key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signer_id":"pith.science","signer_type":"pith_registry"},"payload":{"graph_snapshot":{"author_claims":{"count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","strong_count":0},"builder_version":"pith-number-builder-2026-05-17-v1","claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"paper":{"abstract_excerpt":"The largest source of sound events is web videos. Most videos lack sound event labels at segment level, however, a significant number of them do respond to text queries, from a match found using metadata by search engines. In this paper we explore the extent to which a search query can be used as the true label for detection of sound events in videos. We present a framework for large-scale sound event recognition on web videos. The framework crawls videos using search queries corresponding to 78 sound event labels drawn from three datasets. The datasets are used to train three classifiers, and","authors_text":"Ankit Shah, Anurag Kumar, Benjamin Elizalde, Bhiksha Raj, Rohan Badlani","cross_cats":["cs.AI","cs.IR","eess.AS"],"headline":"","license":"http://creativecommons.org/licenses/by-sa/4.0/","primary_cat":"cs.SD","submitted_at":"2017-11-02T16:32:23Z","title":"Framework for evaluation of sound event detection in web videos"},"references":{"count":0,"internal_anchors":0,"resolved_work":0,"sample":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"1711.00804","kind":"arxiv","version":2},"verdict":{"created_at":null,"id":null,"model_set":{},"one_line_summary":"","pipeline_version":null,"pith_extraction_headline":"","strongest_claim":"","weakest_assumption":""}},"verdict_id":null}}],"author_attestations":[],"timestamp_anchors":[],"storage_attestations":[],"citation_signatures":[],"replication_records":[],"corrections":[],"mirror_hints":[],"record_created":{"event_id":"sha256:1f9f245ef9d0cea3c3301ed95ea5435ba398daae6bc4f532f84472c88c6c71aa","target":"record","created_at":"2026-05-18T00:19:16Z","signer":{"key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signer_id":"pith.science","signer_type":"pith_registry"},"payload":{"attestation_state":"computed","canonical_record":{"metadata":{"abstract_canon_sha256":"754b0457dfa656e501eb24bfe2a98d7230b8069a9407fd71979e5f6ae1993657","cross_cats_sorted":["cs.AI","cs.IR","eess.AS"],"license":"http://creativecommons.org/licenses/by-sa/4.0/","primary_cat":"cs.SD","submitted_at":"2017-11-02T16:32:23Z","title_canon_sha256":"c4ad4808dd52768a3084502c0e0f0944705131684d435f689b5a38e7a71241b0"},"schema_version":"1.0","source":{"id":"1711.00804","kind":"arxiv","version":2}},"canonical_sha256":"d965265d86bf3257222f19272e65fd0d3c0f961984e446fefaebf16738df53e0","receipt":{"algorithm":"ed25519","builder_version":"pith-number-builder-2026-05-17-v1","canonical_sha256":"d965265d86bf3257222f19272e65fd0d3c0f961984e446fefaebf16738df53e0","first_computed_at":"2026-05-18T00:19:16.889444Z","key_id":"pith-v1-2026-05","kind":"pith_receipt","last_reissued_at":"2026-05-18T00:19:16.889444Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","receipt_version":"0.3","signature_b64":"PGVkvZ1KOx83jV0EF/wyMLMVo1qIHENwp3ellUXeoyuMnOoasb6CDFEPdfNloT6K+OgE0ztmckaRhkwmCNecDQ==","signature_status":"signed_v1","signed_at":"2026-05-18T00:19:16.889999Z","signed_message":"canonical_sha256_bytes"},"source_id":"1711.00804","source_kind":"arxiv","source_version":2}}},"equivocations":[],"invalid_events":[],"applied_event_ids":["sha256:1f9f245ef9d0cea3c3301ed95ea5435ba398daae6bc4f532f84472c88c6c71aa","sha256:48560c1aa1c61713279ff31da3a9f245235e5def997b85e9f9bd8527399ee226"],"state_sha256":"19cb3ec756f7c2cc5ff6c5c9b01d8497606cfb00f1f13f4a20c728860e5caf46"},"bundle_signature":{"signature_status":"signed_v1","algorithm":"ed25519","key_id":"pith-v1-2026-05","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54","signature_b64":"RN2D4kCt7O1bPdzXRWmc0AKjjQQ3DGcz/jlnD5nqQiq1iPUoAbZqyeOSpFomLHFZoVjvfhNICi+rlRJxNSAuDw==","signed_message":"bundle_sha256_bytes","signed_at":"2026-05-30T20:21:30.580053Z","bundle_sha256":"dc695e66166e6e4c73688eee38e9072ce54bb154bead709191a9e54d0c5fa237"}}