{"record_type":"pith_number_record","schema_url":"https://pith.science/schemas/pith-number/v1.json","pith_number":"pith:2025:D3QM6DD5DI3QCMIQQEZHXPX5D6","short_pith_number":"pith:D3QM6DD5","schema_version":"1.0","canonical_sha256":"1ee0cf0c7d1a3701311081327bbefd1f9c7383896ffdec493abeb7362e797d0d","source":{"kind":"arxiv","id":"2506.20550","version":1},"attestation_state":"computed","paper":{"title":"Lightweight Multi-Frame Integration for Robust YOLO Object Detection in Videos","license":"http://creativecommons.org/licenses/by/4.0/","headline":"","cross_cats":["cs.RO"],"primary_cat":"cs.CV","authors_text":"Andreas Zell, Benjamin Kiefer, Martin Messmer, Yitong Quan","submitted_at":"2025-06-25T15:49:07Z","abstract_excerpt":"Modern image-based object detection models, such as YOLOv7, primarily process individual frames independently, thus ignoring valuable temporal context naturally present in videos. Meanwhile, existing video-based detection methods often introduce complex temporal modules, significantly increasing model size and computational complexity. In practical applications such as surveillance and autonomous driving, transient challenges including motion blur, occlusions, and abrupt appearance changes can severely degrade single-frame detection performance. To address these issues, we propose a straightfo"},"verification_status":{"content_addressed":true,"pith_receipt":true,"author_attested":false,"weak_author_claims":0,"strong_author_claims":0,"externally_anchored":false,"storage_verified":false,"citation_signatures":0,"replication_records":0,"graph_snapshot":true,"references_resolved":false,"formal_links_present":false},"canonical_record":{"source":{"id":"2506.20550","kind":"arxiv","version":1},"metadata":{"license":"http://creativecommons.org/licenses/by/4.0/","primary_cat":"cs.CV","submitted_at":"2025-06-25T15:49:07Z","cross_cats_sorted":["cs.RO"],"title_canon_sha256":"6072af3b83cd99902198867b1a2bc7b7ee8a4ba27e0f52bbf63f6d147cbb5048","abstract_canon_sha256":"a7fd68de434e7fa8aa50fb64cda3bdf40e814180dcd2bfbdf4aa233605a7afd1"},"schema_version":"1.0"},"receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-07-05T11:27:07.557408Z","signature_b64":"tO92un69v2Ae/r2JNr917magaaHmhpGaI+t+NYAljL/eoNHzPeO9ZeOB79rUl+I+gynCHr/FJOCHB02RgupIAA==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"1ee0cf0c7d1a3701311081327bbefd1f9c7383896ffdec493abeb7362e797d0d","last_reissued_at":"2026-07-05T11:27:07.556898Z","signature_status":"signed_v1","first_computed_at":"2026-07-05T11:27:07.556898Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"graph_snapshot":{"paper":{"title":"Lightweight Multi-Frame Integration for Robust YOLO Object Detection in Videos","license":"http://creativecommons.org/licenses/by/4.0/","headline":"","cross_cats":["cs.RO"],"primary_cat":"cs.CV","authors_text":"Andreas Zell, Benjamin Kiefer, Martin Messmer, Yitong Quan","submitted_at":"2025-06-25T15:49:07Z","abstract_excerpt":"Modern image-based object detection models, such as YOLOv7, primarily process individual frames independently, thus ignoring valuable temporal context naturally present in videos. Meanwhile, existing video-based detection methods often introduce complex temporal modules, significantly increasing model size and computational complexity. In practical applications such as surveillance and autonomous driving, transient challenges including motion blur, occlusions, and abrupt appearance changes can severely degrade single-frame detection performance. To address these issues, we propose a straightfo"},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2506.20550","kind":"arxiv","version":1},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2506.20550/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"aliases":[{"alias_kind":"arxiv","alias_value":"2506.20550","created_at":"2026-07-05T11:27:07.556957+00:00"},{"alias_kind":"arxiv_version","alias_value":"2506.20550v1","created_at":"2026-07-05T11:27:07.556957+00:00"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2506.20550","created_at":"2026-07-05T11:27:07.556957+00:00"},{"alias_kind":"pith_short_12","alias_value":"D3QM6DD5DI3Q","created_at":"2026-07-05T11:27:07.556957+00:00"},{"alias_kind":"pith_short_16","alias_value":"D3QM6DD5DI3QCMIQ","created_at":"2026-07-05T11:27:07.556957+00:00"},{"alias_kind":"pith_short_8","alias_value":"D3QM6DD5","created_at":"2026-07-05T11:27:07.556957+00:00"}],"events":[],"event_summary":{},"paper_claims":[],"inbound_citations":{"count":1,"internal_anchor_count":1,"sample":[{"citing_arxiv_id":"2608.07571","citing_title":"A Review of Vision-Based Vehicle Detection for UAV-Based Traffic Monitoring: Experimental Insights and Future Directions","ref_index":140,"is_internal_anchor":true}]},"formal_canon":{"evidence_count":0,"sample":[],"anchors":[]},"links":{"html":"https://pith.science/pith/D3QM6DD5DI3QCMIQQEZHXPX5D6","json":"https://pith.science/pith/D3QM6DD5DI3QCMIQQEZHXPX5D6.json","graph_json":"https://pith.science/api/pith-number/D3QM6DD5DI3QCMIQQEZHXPX5D6/graph.json","events_json":"https://pith.science/api/pith-number/D3QM6DD5DI3QCMIQQEZHXPX5D6/events.json","paper":"https://pith.science/paper/D3QM6DD5"},"agent_actions":{"view_html":"https://pith.science/pith/D3QM6DD5DI3QCMIQQEZHXPX5D6","download_json":"https://pith.science/pith/D3QM6DD5DI3QCMIQQEZHXPX5D6.json","view_paper":"https://pith.science/paper/D3QM6DD5","resolve_alias":"https://pith.science/api/pith-number/resolve?arxiv=2506.20550&json=true","fetch_graph":"https://pith.science/api/pith-number/D3QM6DD5DI3QCMIQQEZHXPX5D6/graph.json","fetch_events":"https://pith.science/api/pith-number/D3QM6DD5DI3QCMIQQEZHXPX5D6/events.json","actions":{"anchor_timestamp":"https://pith.science/pith/D3QM6DD5DI3QCMIQQEZHXPX5D6/action/timestamp_anchor","attest_storage":"https://pith.science/pith/D3QM6DD5DI3QCMIQQEZHXPX5D6/action/storage_attestation","attest_author":"https://pith.science/pith/D3QM6DD5DI3QCMIQQEZHXPX5D6/action/author_attestation","sign_citation":"https://pith.science/pith/D3QM6DD5DI3QCMIQQEZHXPX5D6/action/citation_signature","submit_replication":"https://pith.science/pith/D3QM6DD5DI3QCMIQQEZHXPX5D6/action/replication_record"}},"created_at":"2026-07-05T11:27:07.556957+00:00","updated_at":"2026-07-05T11:27:07.556957+00:00"}