{"schema":"https://pith.science/schemas/pith-integrity/v1.json","pith_number":"2608.04351","arxiv_id":"2608.04351","integrity":{"available":true,"endpoint":"/pith/2608.04351/integrity.json","summary":{"critical":0,"advisory":5,"informational":0,"by_detector":{"doi_compliance":{"total":5,"critical":0,"advisory":5,"informational":0}}},"clean":false,"detectors_run":[{"name":"doi_title_agreement","version":"1.0.0","status":"completed","ran_at":"2026-08-09T20:08:31.716120Z","findings_count":0},{"name":"doi_compliance","version":"1.1.0","status":"completed","ran_at":"2026-08-09T19:28:32.878108Z","findings_count":5},{"name":"claim_evidence","version":"1.0.0","status":"completed","ran_at":"2026-08-09T04:48:15.125210Z","findings_count":0},{"name":"citation_quote_validity","version":"0.1.0","status":"skipped","ran_at":"2026-08-06T21:08:52.770277Z","findings_count":0},{"name":"cited_work_retraction","version":"1.0.0","status":"completed","ran_at":"2026-08-06T14:38:07.538262Z","findings_count":0},{"name":"ai_meta_artifact","version":"1.0.0","status":"skipped","ran_at":"2026-08-06T03:09:04.849946Z","findings_count":0}],"findings":[{"detector":"doi_compliance","finding_type":"recoverable_identifier","severity":"advisory","verdict_class":"incontrovertible","note":"DOI in the printed bibliography is fragmented by whitespace or line breaks. A longer candidate (10.1007/s10579-008-9076-6) was visible in the surrounding text but could not be confirmed against doi.org as printed.","detected_doi":"10.1007/s10579-008-9076-6","detected_arxiv_id":null,"ref_index":3,"audited_at":"2026-08-08T19:18:26.497914Z"},{"detector":"doi_compliance","finding_type":"recoverable_identifier","severity":"advisory","verdict_class":"incontrovertible","note":"DOI in the printed bibliography is fragmented by whitespace or line breaks. A longer candidate (10.21437/Interspeech.2025-1232) was visible in the surrounding text but could not be confirmed against doi.org as printed.","detected_doi":"10.21437/Interspeech.2025-1232","detected_arxiv_id":null,"ref_index":9,"audited_at":"2026-08-08T19:18:26.497914Z"},{"detector":"doi_compliance","finding_type":"recoverable_identifier","severity":"advisory","verdict_class":"incontrovertible","note":"DOI in the printed bibliography is fragmented by whitespace or line breaks. A longer candidate (10.48550/ARXIV.2409.05566) was visible in the surrounding text but could not be confirmed against doi.org as printed.","detected_doi":"10.48550/ARXIV.2409.05566","detected_arxiv_id":null,"ref_index":10,"audited_at":"2026-08-08T19:18:26.497914Z"},{"detector":"doi_compliance","finding_type":"recoverable_identifier","severity":"advisory","verdict_class":"incontrovertible","note":"DOI is split by whitespace or line breaks in the printed bibliography. Reconstructed DOI 10.1037/0033-2909.99.2.143 resolves to 'Vocal affect expression: A review and a model for future research.'. A reader following the printed text alone cannot reach it.","detected_doi":"10.1037/0033-2909.99.2.143","detected_arxiv_id":null,"ref_index":35,"audited_at":"2026-08-08T19:18:26.497914Z"},{"detector":"doi_compliance","finding_type":"recoverable_identifier","severity":"advisory","verdict_class":"incontrovertible","note":"DOI in the printed bibliography is fragmented by whitespace or line breaks. A longer candidate (10.18653/v1/2025.emnlp-main.514) was visible in the surrounding text but could not be confirmed against doi.org as printed.","detected_doi":"10.18653/v1/2025.emnlp-main.514","detected_arxiv_id":null,"ref_index":41,"audited_at":"2026-08-08T19:18:26.497914Z"}],"snapshot_sha256":"b09726a587fa735ff72497bf06e4a38abbd09678e116f848d7032654f0314732"},"events":[{"event_id":18721,"event_type":"pith.integrity.v1","payload_sha256":"5fe26b8294312a666d04b753be7e450649c3313f8de9d6d95c5b46c37f62acf5","signature_b64":"R1DskoEmLuEttcLa1hr1phjjUd69twL6p8FAJZCcfZXch0NIzhOD5jjyd2v/0Ay2c69PaBsdQgQIXau2H7tEAQ==","signing_key_id":"pith-v1-2026-05","created_at":"2026-08-08T19:18:35.825743+00:00","payload":{"note":"DOI in the printed bibliography is fragmented by whitespace or line breaks. A longer candidate (10.18653/v1/2025.emnlp-main.514) was visible in the surrounding text but could not be confirmed against doi.org as printed.","snippet":"Chih-Kai Yang, Neo S. Ho, and Hung-yi Lee. 2025. Towards Holistic Evaluation of Large Audio-Language Models: A Comprehensive Survey. InProceedings of the 2025 Conference on Empirical Methods in Natural Language Processing. Association for C","arxiv_id":"2608.04351","detector":"doi_compliance","evidence":{"ref_index":41,"verdict_class":"incontrovertible","resolved_title":null,"printed_excerpt":"10.18653/v1/","reconstructed_doi":"10.18653/v1/2025.emnlp-main.514"},"severity":"advisory","ref_index":41,"audited_at":"2026-08-08T19:18:26.497914Z","event_type":"pith.integrity.v1","detected_doi":"10.18653/v1/2025.emnlp-main.514","detector_url":"https://pith.science/pith-integrity-protocol#doi_compliance","external_url":null,"finding_type":"recoverable_identifier","evidence_hash":"25e0e27aeda97f4709d8cbf2b81939154ebdf34957a40be7fdc2094fe6d71b0e","paper_version":1,"verdict_class":"incontrovertible","resolved_title":null,"detector_version":"1.1.0","detected_arxiv_id":null}},{"event_id":18720,"event_type":"pith.integrity.v1","payload_sha256":"cbff081c959b691a724b874b69663b2c6d1e40b4f1fa9a3d24f55b260b00615b","signature_b64":"HrF06SZb+XaCfoowl0UaDDtz2pVL4Vrw1rc8v8digRx73kPfvZtoeTetakjRBksOtYZb6Uv2UAbeDlqAHwshDA==","signing_key_id":"pith-v1-2026-05","created_at":"2026-08-08T19:18:35.817704+00:00","payload":{"note":"DOI is split by whitespace or line breaks in the printed bibliography. Reconstructed DOI 10.1037/0033-2909.99.2.143 resolves to 'Vocal affect expression: A review and a model for future research.'. A reader following the printed text alone cannot reach it.","snippet":"Klaus R. Scherer. 1986. Vocal affect expression: A review and a model for future research.Psychological Bulletin99, 2 (1986), 143–165. doi:10.1037/0033-2909.99.2. 143","arxiv_id":"2608.04351","detector":"doi_compliance","evidence":{"ref_index":35,"verdict_class":"incontrovertible","resolved_title":"Vocal affect expression: A review and a model for future research.","printed_excerpt":"10.1037/0033-2909.99.2","reconstructed_doi":"10.1037/0033-2909.99.2.143"},"severity":"advisory","ref_index":35,"audited_at":"2026-08-08T19:18:26.497914Z","event_type":"pith.integrity.v1","detected_doi":"10.1037/0033-2909.99.2.143","detector_url":"https://pith.science/pith-integrity-protocol#doi_compliance","external_url":null,"finding_type":"recoverable_identifier","evidence_hash":"1ff633d3e1d681f2b5feaed3a14728c7b8419f0986a892bff5919d74b0d64635","paper_version":1,"verdict_class":"incontrovertible","resolved_title":"Vocal affect expression: A review and a model for future research.","detector_version":"1.1.0","detected_arxiv_id":null}},{"event_id":18719,"event_type":"pith.integrity.v1","payload_sha256":"082b20dfcc5257f211e30793483decca6e7a3812cab5476f638e2f1624fc9965","signature_b64":"Tj45vbttUOu+53Vq1+C7vd24hpDVc1E02W3UhC6RmQz5sVWjbWUyRRl1FXcQvYezhp+G0t7iTx4wkV2k0e9eCA==","signing_key_id":"pith-v1-2026-05","created_at":"2026-08-08T19:18:35.809813+00:00","payload":{"note":"DOI in the printed bibliography is fragmented by whitespace or line breaks. A longer candidate (10.48550/ARXIV.2409.05566) was visible in the surrounding text but could not be confirmed against doi.org as printed.","snippet":"Soumya Dutta and Sriram Ganapathy. 2024. Leveraging Content and Acoustic Representations for Speech Emotion Recognition. doi:10.48550/ARXIV.2409. 05566","arxiv_id":"2608.04351","detector":"doi_compliance","evidence":{"ref_index":10,"verdict_class":"incontrovertible","resolved_title":null,"printed_excerpt":"10.48550/arxiv.2409","reconstructed_doi":"10.48550/ARXIV.2409.05566"},"severity":"advisory","ref_index":10,"audited_at":"2026-08-08T19:18:26.497914Z","event_type":"pith.integrity.v1","detected_doi":"10.48550/ARXIV.2409.05566","detector_url":"https://pith.science/pith-integrity-protocol#doi_compliance","external_url":null,"finding_type":"recoverable_identifier","evidence_hash":"18604e03855d6bf035e5e44cfe15ab11272d3c419614fbd3aed56fccf87f5fa8","paper_version":1,"verdict_class":"incontrovertible","resolved_title":null,"detector_version":"1.1.0","detected_arxiv_id":null}},{"event_id":18718,"event_type":"pith.integrity.v1","payload_sha256":"a7045006d8ac37a3d74f19506d0aaea1a72ee8239bd1561e73ab3b10c6b671f4","signature_b64":"1Ra95CUQ84GjfzysIY63svExwSChJFOpaN5wkT59QGCzOAVB6Zi5UezXxiV22U2VY342obglYsruquMNPij0Aw==","signing_key_id":"pith-v1-2026-05","created_at":"2026-08-08T19:18:35.802026+00:00","payload":{"note":"DOI in the printed bibliography is fragmented by whitespace or line breaks. A longer candidate (10.21437/Interspeech.2025-1232) was visible in the surrounding text but could not be confirmed against doi.org as printed.","snippet":"Hongfei Du, Sidi Lu, Gang Zhou, and Ye Gao. 2025. EAA: Emotion-Aware Audio Large Language Models with Dual Cross-Attention and Context-Aware Instruction Tuning. InInterspeech 2025. 5433–5437. doi:10.21437/Interspeech. 2025-1232","arxiv_id":"2608.04351","detector":"doi_compliance","evidence":{"ref_index":9,"verdict_class":"incontrovertible","resolved_title":null,"printed_excerpt":"10.21437/interspeech","reconstructed_doi":"10.21437/Interspeech.2025-1232"},"severity":"advisory","ref_index":9,"audited_at":"2026-08-08T19:18:26.497914Z","event_type":"pith.integrity.v1","detected_doi":"10.21437/Interspeech.2025-1232","detector_url":"https://pith.science/pith-integrity-protocol#doi_compliance","external_url":null,"finding_type":"recoverable_identifier","evidence_hash":"db6abb1da5719eb79430c0cf58a5aee9fc8609ffaad70afa28f0a066e77e6ef6","paper_version":1,"verdict_class":"incontrovertible","resolved_title":null,"detector_version":"1.1.0","detected_arxiv_id":null}},{"event_id":18717,"event_type":"pith.integrity.v1","payload_sha256":"6e93178bb01071def27825cfa7f9b2dfada2d8b7bf3946a5c289dcd788be038f","signature_b64":"x9gMSzcyNj/x+ZX5qq2Uxdx7e5vMGfbC6HVHgDOyYDsOG3T3zaiE1XmDPU75FtGwybBQ4LqycWNiNfHaN9rlCg==","signing_key_id":"pith-v1-2026-05","created_at":"2026-08-08T19:18:35.794534+00:00","payload":{"note":"DOI in the printed bibliography is fragmented by whitespace or line breaks. A longer candidate (10.1007/s10579-008-9076-6) was visible in the surrounding text but could not be confirmed against doi.org as printed.","snippet":"Carlos Busso, Murtaza Bulut, Chi-Chun Lee, Abe Kazemzadeh, Emily Mower, Samuel Kim, Jeannette N. Chang, Sungbok Lee, and Shrikanth S. Narayanan. 2008. IEMOCAP: interactive emotional dyadic motion capture database.Language Resources and Eval","arxiv_id":"2608.04351","detector":"doi_compliance","evidence":{"ref_index":3,"verdict_class":"incontrovertible","resolved_title":null,"printed_excerpt":"10.1007/s10579-008-","reconstructed_doi":"10.1007/s10579-008-9076-6"},"severity":"advisory","ref_index":3,"audited_at":"2026-08-08T19:18:26.497914Z","event_type":"pith.integrity.v1","detected_doi":"10.1007/s10579-008-9076-6","detector_url":"https://pith.science/pith-integrity-protocol#doi_compliance","external_url":null,"finding_type":"recoverable_identifier","evidence_hash":"01634170154144e344c0f3f255bc560a9452bbe2f295ead29ed7d0090e10f92f","paper_version":1,"verdict_class":"incontrovertible","resolved_title":null,"detector_version":"1.1.0","detected_arxiv_id":null}}],"endpoint_self":"/pith/2608.04351/integrity.json","protocol_url":"https://pith.science/pith-integrity-protocol"}