{"record_type":"pith_number_record","schema_url":"https://pith.science/schemas/pith-number/v1.json","pith_number":"pith:2025:V4MJ4BNZ25WKAF3ASBM2XZTTL6","short_pith_number":"pith:V4MJ4BNZ","schema_version":"1.0","canonical_sha256":"af189e05b9d76ca017609059abe6735fb487f47c0b59ac95c3188bcdab7e30b2","source":{"kind":"arxiv","id":"2502.03212","version":1},"attestation_state":"computed","paper":{"title":"Leveraging Broadcast Media Subtitle Transcripts for Automatic Speech Recognition and Subtitling","license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","headline":"","cross_cats":["cs.SD"],"primary_cat":"eess.AS","authors_text":"Hugo Van hamme, Jakob Poncelet","submitted_at":"2025-02-05T14:26:58Z","abstract_excerpt":"The recent advancement of speech recognition technology has been driven by large-scale datasets and attention-based architectures, but many challenges still remain, especially for low-resource languages and dialects. This paper explores the integration of weakly supervised transcripts from TV subtitles into automatic speech recognition (ASR) systems, aiming to improve both verbatim transcriptions and automatically generated subtitles. To this end, verbatim data and subtitles are regarded as different domains or languages, due to their distinct characteristics. We propose and compare several en"},"verification_status":{"content_addressed":true,"pith_receipt":true,"author_attested":false,"weak_author_claims":0,"strong_author_claims":0,"externally_anchored":false,"storage_verified":false,"citation_signatures":0,"replication_records":0,"graph_snapshot":true,"references_resolved":false,"formal_links_present":false},"canonical_record":{"source":{"id":"2502.03212","kind":"arxiv","version":1},"metadata":{"license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","primary_cat":"eess.AS","submitted_at":"2025-02-05T14:26:58Z","cross_cats_sorted":["cs.SD"],"title_canon_sha256":"089757389ab323da556242cbc2b78fd3d839ca85885752370a22c922da2f50c7","abstract_canon_sha256":"cd5ebe5711d82f8c030a28873daacbd68dad4811291cad149a861c44b70c4d97"},"schema_version":"1.0"},"receipt":{"kind":"pith_receipt","key_id":"pith-v1-2026-05","algorithm":"ed25519","signed_at":"2026-07-05T10:10:01.499890Z","signature_b64":"EpGiYLooJkvSggPu+SVdRB4NZ7xom8jecJjA01tAGERQ5Yfm9zBwyDLKi9N4EhCa5SreC3OeEjYpSJDLLB4eAg==","signed_message":"canonical_sha256_bytes","builder_version":"pith-number-builder-2026-05-17-v1","receipt_version":"0.3","canonical_sha256":"af189e05b9d76ca017609059abe6735fb487f47c0b59ac95c3188bcdab7e30b2","last_reissued_at":"2026-07-05T10:10:01.499479Z","signature_status":"signed_v1","first_computed_at":"2026-07-05T10:10:01.499479Z","public_key_fingerprint":"8d4b5ee74e4693bcd1df2446408b0d54"},"graph_snapshot":{"paper":{"title":"Leveraging Broadcast Media Subtitle Transcripts for Automatic Speech Recognition and Subtitling","license":"http://arxiv.org/licenses/nonexclusive-distrib/1.0/","headline":"","cross_cats":["cs.SD"],"primary_cat":"eess.AS","authors_text":"Hugo Van hamme, Jakob Poncelet","submitted_at":"2025-02-05T14:26:58Z","abstract_excerpt":"The recent advancement of speech recognition technology has been driven by large-scale datasets and attention-based architectures, but many challenges still remain, especially for low-resource languages and dialects. This paper explores the integration of weakly supervised transcripts from TV subtitles into automatic speech recognition (ASR) systems, aiming to improve both verbatim transcriptions and automatically generated subtitles. To this end, verbatim data and subtitles are regarded as different domains or languages, due to their distinct characteristics. We propose and compare several en"},"claims":{"count":0,"items":[],"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"source":{"id":"2502.03212","kind":"arxiv","version":1},"verdict":{"id":null,"model_set":{},"created_at":null,"strongest_claim":"","one_line_summary":"","pipeline_version":null,"weakest_assumption":"","pith_extraction_headline":""},"integrity":{"clean":true,"summary":{"advisory":0,"critical":0,"by_detector":{},"informational":0},"endpoint":"/pith/2502.03212/integrity.json","findings":[],"available":true,"detectors_run":[],"snapshot_sha256":"c28c3603d3b5d939e8dc4c7e95fa8dfce3d595e45f758748cecf8e644a296938"},"references":{"count":0,"sample":[],"resolved_work":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57","internal_anchors":0},"formal_canon":{"evidence_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"author_claims":{"count":0,"strong_count":0,"snapshot_sha256":"258153158e38e3291e3d48162225fcdb2d5a3ed65a07baac614ab91432fd4f57"},"builder_version":"pith-number-builder-2026-05-17-v1"},"aliases":[{"alias_kind":"arxiv","alias_value":"2502.03212","created_at":"2026-07-05T10:10:01.499538+00:00"},{"alias_kind":"arxiv_version","alias_value":"2502.03212v1","created_at":"2026-07-05T10:10:01.499538+00:00"},{"alias_kind":"doi","alias_value":"10.48550/arxiv.2502.03212","created_at":"2026-07-05T10:10:01.499538+00:00"},{"alias_kind":"pith_short_12","alias_value":"V4MJ4BNZ25WK","created_at":"2026-07-05T10:10:01.499538+00:00"},{"alias_kind":"pith_short_16","alias_value":"V4MJ4BNZ25WKAF3A","created_at":"2026-07-05T10:10:01.499538+00:00"},{"alias_kind":"pith_short_8","alias_value":"V4MJ4BNZ","created_at":"2026-07-05T10:10:01.499538+00:00"}],"events":[],"event_summary":{},"paper_claims":[],"inbound_citations":{"count":1,"internal_anchor_count":0,"sample":[{"citing_arxiv_id":"2606.10864","citing_title":"Phoneme-First Prediction for LLM-Based Speech Recognition","ref_index":39,"is_internal_anchor":false}]},"formal_canon":{"evidence_count":0,"sample":[],"anchors":[]},"links":{"html":"https://pith.science/pith/V4MJ4BNZ25WKAF3ASBM2XZTTL6","json":"https://pith.science/pith/V4MJ4BNZ25WKAF3ASBM2XZTTL6.json","graph_json":"https://pith.science/api/pith-number/V4MJ4BNZ25WKAF3ASBM2XZTTL6/graph.json","events_json":"https://pith.science/api/pith-number/V4MJ4BNZ25WKAF3ASBM2XZTTL6/events.json","paper":"https://pith.science/paper/V4MJ4BNZ"},"agent_actions":{"view_html":"https://pith.science/pith/V4MJ4BNZ25WKAF3ASBM2XZTTL6","download_json":"https://pith.science/pith/V4MJ4BNZ25WKAF3ASBM2XZTTL6.json","view_paper":"https://pith.science/paper/V4MJ4BNZ","resolve_alias":"https://pith.science/api/pith-number/resolve?arxiv=2502.03212&json=true","fetch_graph":"https://pith.science/api/pith-number/V4MJ4BNZ25WKAF3ASBM2XZTTL6/graph.json","fetch_events":"https://pith.science/api/pith-number/V4MJ4BNZ25WKAF3ASBM2XZTTL6/events.json","actions":{"anchor_timestamp":"https://pith.science/pith/V4MJ4BNZ25WKAF3ASBM2XZTTL6/action/timestamp_anchor","attest_storage":"https://pith.science/pith/V4MJ4BNZ25WKAF3ASBM2XZTTL6/action/storage_attestation","attest_author":"https://pith.science/pith/V4MJ4BNZ25WKAF3ASBM2XZTTL6/action/author_attestation","sign_citation":"https://pith.science/pith/V4MJ4BNZ25WKAF3ASBM2XZTTL6/action/citation_signature","submit_replication":"https://pith.science/pith/V4MJ4BNZ25WKAF3ASBM2XZTTL6/action/replication_record"}},"created_at":"2026-07-05T10:10:01.499538+00:00","updated_at":"2026-07-05T10:10:01.499538+00:00"}