{"as_of":"2026-08-20T09:13:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:2152d695385b41eb8b8b39360d8a041437553d51825829a7b6efb0b868af1968","coverage":[{"denominator":25,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":25,"source":"paper_references, paper_reference_links","source_observed_at":"2026-07-12T06:06:46.343211Z","state":"measured"},{"denominator":25,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":25,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-20T06:33:59.587034+00:00","state":"measured"},{"denominator":0,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"cited_works","source_observed_at":null,"state":"measured"}],"external_citation_measurements":[],"inbound":[],"links":{"evidence":"/evidence","html":"/paper/2607.02920/citation-record","integrity":"/paper/2607.02920/integrity","json":"/paper/2607.02920/citation-record.json","paper":"/paper/2607.02920"},"outbound":[{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-12T06:06:46.343211Z","title":"Depressive disorder (depression),","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2607.02920","last_updated":"2026-07-03T03:24:38Z","snapshot_observed_at":"2026-08-19T17:04:02.832931Z","submitted_at":"2026-07-03T03:24:38Z","title":"Layer-wise Cross-Lingual Depression Detection from Speech: Analysis with Contrastive Alignment","version":1},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-07-12T06:06:46.343211Z"},"links":{"citing_paper":"/paper/2607.02920"},"observation_digest":"sha256:3bb22403beb7144a8e10c89b3e004a5b2b0609d749e62905aed9689074347e32","observation_id":"e451ee52-f1e0-4d81-bd63-3a3c34a78ec3","resolution":{"observed_at":"2026-07-12T06:06:46.343211Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-12T06:06:46.343211Z","title":"A review of depression and suicide risk assessment using speech analysis,","venue":null,"work_id":null,"year":2015},"citing_paper":{"arxiv_id":"2607.02920","last_updated":"2026-07-03T03:24:38Z","snapshot_observed_at":"2026-08-19T17:04:02.832931Z","submitted_at":"2026-07-03T03:24:38Z","title":"Layer-wise Cross-Lingual Depression Detection from Speech: Analysis with Contrastive Alignment","version":1},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-07-12T06:06:46.343211Z"},"links":{"citing_paper":"/paper/2607.02920"},"observation_digest":"sha256:d162ba1157568f721942e8c15795a5e2d0350c821e3a79e61aad03ea691c0209","observation_id":"a7d3bc50-4302-4854-8d60-c66367e7fa73","resolution":{"observed_at":"2026-07-12T06:06:46.343211Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-12T06:06:46.343211Z","title":"Clinical state tracking in serious mental illness through computational analysis of speech,","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2607.02920","last_updated":"2026-07-03T03:24:38Z","snapshot_observed_at":"2026-08-19T17:04:02.832931Z","submitted_at":"2026-07-03T03:24:38Z","title":"Layer-wise Cross-Lingual Depression Detection from Speech: Analysis with Contrastive Alignment","version":1},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-07-12T06:06:46.343211Z"},"links":{"citing_paper":"/paper/2607.02920"},"observation_digest":"sha256:ca03931c9a632c5004a6cae1ae86ae9a84938837c732a2ab6befa8ebeb64e385","observation_id":"897f45b3-225e-46a3-aa50-fc868088f600","resolution":{"observed_at":"2026-07-12T06:06:46.343211Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-12T06:06:46.343211Z","title":"Speech as a promising biosignal in precision psychiatry,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2607.02920","last_updated":"2026-07-03T03:24:38Z","snapshot_observed_at":"2026-08-19T17:04:02.832931Z","submitted_at":"2026-07-03T03:24:38Z","title":"Layer-wise Cross-Lingual Depression Detection from Speech: Analysis with Contrastive Alignment","version":1},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-07-12T06:06:46.343211Z"},"links":{"citing_paper":"/paper/2607.02920"},"observation_digest":"sha256:5627da1784ff572b4a98480a229ff29c8c8203c74e9a495b5ab47811e4e0f692","observation_id":"9ca8d9af-d3f7-4b8a-9b85-61b7116d2a6c","resolution":{"observed_at":"2026-07-12T06:06:46.343211Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2106.07447","last_updated":"2021-06-14T14:14:28Z","snapshot_observed_at":"2026-08-17T11:07:54.315530Z","submitted_at":"2021-06-14T14:14:28Z","title":"HuBERT: Self-Supervised Speech Representation Learning by Masked Prediction of Hidden Units","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2106.07447","snapshot_observed_at":"2026-07-12T06:06:46.343211Z","title":"HuBERT: Self-supervised speech representation learning by masked prediction of hidden units,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2607.02920","last_updated":"2026-07-03T03:24:38Z","snapshot_observed_at":"2026-08-19T17:04:02.832931Z","submitted_at":"2026-07-03T03:24:38Z","title":"Layer-wise Cross-Lingual Depression Detection from Speech: Analysis with Contrastive Alignment","version":1},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-07-12T06:06:46.343211Z"},"links":{"cited_paper":"/paper/2106.07447","citing_paper":"/paper/2607.02920"},"observation_digest":"sha256:4dc598e09508884a161b9ac92250a5d02b4916e30a99fd51831c37198cfacde5","observation_id":"2e838a15-a1a9-403d-9d33-4fc1b0085c1f","resolution":{"observed_at":"2026-07-12T06:06:46.343211Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-12T06:06:46.343211Z","title":"WavLM: Large-scale self-supervised pre-training for full stack speech processing,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2607.02920","last_updated":"2026-07-03T03:24:38Z","snapshot_observed_at":"2026-08-19T17:04:02.832931Z","submitted_at":"2026-07-03T03:24:38Z","title":"Layer-wise Cross-Lingual Depression Detection from Speech: Analysis with Contrastive Alignment","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-07-12T06:06:46.343211Z"},"links":{"citing_paper":"/paper/2607.02920"},"observation_digest":"sha256:27287daae9d8fec7612306c027d8c9968bd53ed47d19f8df18ea40f6e6877227","observation_id":"0065a578-3380-411b-85c0-1ecfcb1eda17","resolution":{"observed_at":"2026-07-12T06:06:46.343211Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-12T06:06:46.343211Z","title":"Inves- tigation of layer-wise speech representations in self-supervised learning models: A cross-lingual study in detecting depression,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.02920","last_updated":"2026-07-03T03:24:38Z","snapshot_observed_at":"2026-08-19T17:04:02.832931Z","submitted_at":"2026-07-03T03:24:38Z","title":"Layer-wise Cross-Lingual Depression Detection from Speech: Analysis with Contrastive Alignment","version":1},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-07-12T06:06:46.343211Z"},"links":{"citing_paper":"/paper/2607.02920"},"observation_digest":"sha256:c51eebfcb7e3ad9e39e8e8149ec37ed32cdfc162f3816cff44987f1283ef0925","observation_id":"6565990b-9fe7-44b3-834a-21ead81d5500","resolution":{"observed_at":"2026-07-12T06:06:46.343211Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2409.16920","last_updated":"2025-04-30T13:16:09Z","snapshot_observed_at":"2026-08-18T05:16:23.820488Z","submitted_at":"2024-09-25T13:27:17Z","title":"Cross-Lingual Speech Emotion Recognition: Humans vs. Self-Supervised Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2409.16920","snapshot_observed_at":"2026-07-12T06:06:46.343211Z","title":"Cross- lingual speech emotion recognition: Humans vs. self-supervised mod- els,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2607.02920","last_updated":"2026-07-03T03:24:38Z","snapshot_observed_at":"2026-08-19T17:04:02.832931Z","submitted_at":"2026-07-03T03:24:38Z","title":"Layer-wise Cross-Lingual Depression Detection from Speech: Analysis with Contrastive Alignment","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-07-12T06:06:46.343211Z"},"links":{"cited_paper":"/paper/2409.16920","citing_paper":"/paper/2607.02920"},"observation_digest":"sha256:906b4f111e4134ec8e2f95c487a6c14eeaf539f2c3fbe8bedd7b6c79ef816a14","observation_id":"c9a4d282-efc8-4e3a-9c59-61378e3f44b7","resolution":{"observed_at":"2026-07-12T06:06:46.343211Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-12T06:06:46.343211Z","title":"Decoupled weight decay regularization,","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2607.02920","last_updated":"2026-07-03T03:24:38Z","snapshot_observed_at":"2026-08-19T17:04:02.832931Z","submitted_at":"2026-07-03T03:24:38Z","title":"Layer-wise Cross-Lingual Depression Detection from Speech: Analysis with Contrastive Alignment","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-07-12T06:06:46.343211Z"},"links":{"citing_paper":"/paper/2607.02920"},"observation_digest":"sha256:0e8634f76a9700a67d4e5cbbc209919f146d6af3f80fa1ed8eb43f2a1fb99545","observation_id":"fc186c6c-d216-4c52-97d6-6dc7019893ee","resolution":{"observed_at":"2026-07-12T06:06:46.343211Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-12T06:06:46.343211Z","title":"ML-SUPERB: Multilingual speech universal performance benchmark,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2607.02920","last_updated":"2026-07-03T03:24:38Z","snapshot_observed_at":"2026-08-19T17:04:02.832931Z","submitted_at":"2026-07-03T03:24:38Z","title":"Layer-wise Cross-Lingual Depression Detection from Speech: Analysis with Contrastive Alignment","version":1},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-07-12T06:06:46.343211Z"},"links":{"citing_paper":"/paper/2607.02920"},"observation_digest":"sha256:52476e041c0567fef90aca3ce4b3681ee43584e661e13dd80f86f51ef2fcae25","observation_id":"d9eff775-018c-46d0-b401-88f3f224e06c","resolution":{"observed_at":"2026-07-12T06:06:46.343211Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-12T06:06:46.343211Z","title":"SUPERB: Speech processing universal performance benchmark,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2607.02920","last_updated":"2026-07-03T03:24:38Z","snapshot_observed_at":"2026-08-19T17:04:02.832931Z","submitted_at":"2026-07-03T03:24:38Z","title":"Layer-wise Cross-Lingual Depression Detection from Speech: Analysis with Contrastive Alignment","version":1},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-07-12T06:06:46.343211Z"},"links":{"citing_paper":"/paper/2607.02920"},"observation_digest":"sha256:ba91702b89df467e1ca89fb4012f6c8ba21f139525785e0bd6a131ff85335c0a","observation_id":"995c2a1b-5130-4a20-8406-02f87c9ddd44","resolution":{"observed_at":"2026-07-12T06:06:46.343211Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-12T06:06:46.343211Z","title":"Does restriction of pitch variation affect the perception of vocal emotions in Mandarin Chinese?","venue":null,"work_id":null,"year":2015},"citing_paper":{"arxiv_id":"2607.02920","last_updated":"2026-07-03T03:24:38Z","snapshot_observed_at":"2026-08-19T17:04:02.832931Z","submitted_at":"2026-07-03T03:24:38Z","title":"Layer-wise Cross-Lingual Depression Detection from Speech: Analysis with Contrastive Alignment","version":1},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-07-12T06:06:46.343211Z"},"links":{"citing_paper":"/paper/2607.02920"},"observation_digest":"sha256:47acbaf51282a6d020878ba2872e45e106fd2127261cb8ac445fb63a6ecc098f","observation_id":"f15d0693-dbc9-442c-bb3e-b2270426f0be","resolution":{"observed_at":"2026-07-12T06:06:46.343211Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-12T06:06:46.343211Z","title":"Cultural variations in the clinical presentation of depression and anxiety,","venue":null,"work_id":null,"year":2001},"citing_paper":{"arxiv_id":"2607.02920","last_updated":"2026-07-03T03:24:38Z","snapshot_observed_at":"2026-08-19T17:04:02.832931Z","submitted_at":"2026-07-03T03:24:38Z","title":"Layer-wise Cross-Lingual Depression Detection from Speech: Analysis with Contrastive Alignment","version":1},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-07-12T06:06:46.343211Z"},"links":{"citing_paper":"/paper/2607.02920"},"observation_digest":"sha256:0c9512dfbc9e698959374e12e1532e3ae207fdf6a183587aadb8c09a4bba44f3","observation_id":"5a3d1f9e-0594-487d-aa44-cfe0ba7c705a","resolution":{"observed_at":"2026-07-12T06:06:46.343211Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-12T06:06:46.343211Z","title":"Association between acoustic speech features and anxiety and depression symptoms across lifespan,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2607.02920","last_updated":"2026-07-03T03:24:38Z","snapshot_observed_at":"2026-08-19T17:04:02.832931Z","submitted_at":"2026-07-03T03:24:38Z","title":"Layer-wise Cross-Lingual Depression Detection from Speech: Analysis with Contrastive Alignment","version":1},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-07-12T06:06:46.343211Z"},"links":{"citing_paper":"/paper/2607.02920"},"observation_digest":"sha256:d0efa29de9b4faa9e858ebf0bc6db69b977b7dda1d47809d3cad1d7576972f19","observation_id":"809b2afc-193b-4a91-bb26-639eeafbf4ed","resolution":{"observed_at":"2026-07-12T06:06:46.343211Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-12T06:06:46.343211Z","title":"Detecting depression with audio/text sequence modeling of interviews,","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2607.02920","last_updated":"2026-07-03T03:24:38Z","snapshot_observed_at":"2026-08-19T17:04:02.832931Z","submitted_at":"2026-07-03T03:24:38Z","title":"Layer-wise Cross-Lingual Depression Detection from Speech: Analysis with Contrastive Alignment","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-07-12T06:06:46.343211Z"},"links":{"citing_paper":"/paper/2607.02920"},"observation_digest":"sha256:2b6c120ff2ef0cef86300b6dca66f656ef77d807a479db22421ff1e1e8b4f018","observation_id":"f1e87ae1-72f9-4c6f-a028-e9fede55573a","resolution":{"observed_at":"2026-07-12T06:06:46.343211Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-12T06:06:46.343211Z","title":"DEPA: Self-supervised audio embedding for depression detection,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2607.02920","last_updated":"2026-07-03T03:24:38Z","snapshot_observed_at":"2026-08-19T17:04:02.832931Z","submitted_at":"2026-07-03T03:24:38Z","title":"Layer-wise Cross-Lingual Depression Detection from Speech: Analysis with Contrastive Alignment","version":1},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-07-12T06:06:46.343211Z"},"links":{"citing_paper":"/paper/2607.02920"},"observation_digest":"sha256:55ec7fb7f135b10eed6b8d4dabe48e04fea8cc5ba537b2b6c24719ad514ab78a","observation_id":"9480a1de-2130-467c-98f0-2885f498b5c3","resolution":{"observed_at":"2026-07-12T06:06:46.343211Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2104.03502","last_updated":"2021-04-08T04:31:58Z","snapshot_observed_at":"2026-08-16T18:33:07.507126Z","submitted_at":"2021-04-08T04:31:58Z","title":"Emotion Recognition from Speech Using Wav2vec 2.0 Embeddings","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2104.03502","snapshot_observed_at":"2026-07-12T06:06:46.343211Z","title":"Emotion recognition from speech using wav2vec 2.0 embeddings,","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2607.02920","last_updated":"2026-07-03T03:24:38Z","snapshot_observed_at":"2026-08-19T17:04:02.832931Z","submitted_at":"2026-07-03T03:24:38Z","title":"Layer-wise Cross-Lingual Depression Detection from Speech: Analysis with Contrastive Alignment","version":1},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-07-12T06:06:46.343211Z"},"links":{"cited_paper":"/paper/2104.03502","citing_paper":"/paper/2607.02920"},"observation_digest":"sha256:890d70be885362bdcd0577777ccf4f8a7d36d59277144db61d57200565295651","observation_id":"3677867c-4038-4993-976c-4a15694b23b4","resolution":{"observed_at":"2026-07-12T06:06:46.343211Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-12T06:06:46.343211Z","title":"Layer-wise cross-lingual depression detection from speech: A HuBERT-based study on English and Mandarin,","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2607.02920","last_updated":"2026-07-03T03:24:38Z","snapshot_observed_at":"2026-08-19T17:04:02.832931Z","submitted_at":"2026-07-03T03:24:38Z","title":"Layer-wise Cross-Lingual Depression Detection from Speech: Analysis with Contrastive Alignment","version":1},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-07-12T06:06:46.343211Z"},"links":{"citing_paper":"/paper/2607.02920"},"observation_digest":"sha256:0404b5c50fec4978345fd156ae23ef2f0402783d20784dd88dd7fa1b46f755d9","observation_id":"31300b47-655a-486c-a333-dd416ca6bd9f","resolution":{"observed_at":"2026-07-12T06:06:46.343211Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-12T06:06:46.343211Z","title":"The distress analysis interview corpus of human and computer interviews,","venue":null,"work_id":null,"year":2014},"citing_paper":{"arxiv_id":"2607.02920","last_updated":"2026-07-03T03:24:38Z","snapshot_observed_at":"2026-08-19T17:04:02.832931Z","submitted_at":"2026-07-03T03:24:38Z","title":"Layer-wise Cross-Lingual Depression Detection from Speech: Analysis with Contrastive Alignment","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-07-12T06:06:46.343211Z"},"links":{"citing_paper":"/paper/2607.02920"},"observation_digest":"sha256:f479f4e8c010dd71b6a5b2182daf9c009f6c23ad828c8b39935c054f91ec3a84","observation_id":"385726b4-93e9-4106-aa42-e32500115b90","resolution":{"observed_at":"2026-07-12T06:06:46.343211Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-12T06:06:46.343211Z","title":"Semi-structural interview-based Chinese multimodal depression corpus,","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2607.02920","last_updated":"2026-07-03T03:24:38Z","snapshot_observed_at":"2026-08-19T17:04:02.832931Z","submitted_at":"2026-07-03T03:24:38Z","title":"Layer-wise Cross-Lingual Depression Detection from Speech: Analysis with Contrastive Alignment","version":1},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-07-12T06:06:46.343211Z"},"links":{"citing_paper":"/paper/2607.02920"},"observation_digest":"sha256:b9eaa0093973bdcdbaf237cee165d4d9ff9c82a49d92e110d6c7b902afdd8ab3","observation_id":"8b6e80b0-780d-46ab-9e31-8d88cb73e5c5","resolution":{"observed_at":"2026-07-12T06:06:46.343211Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-12T06:06:46.343211Z","title":"Supervised contrastive learning,","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2607.02920","last_updated":"2026-07-03T03:24:38Z","snapshot_observed_at":"2026-08-19T17:04:02.832931Z","submitted_at":"2026-07-03T03:24:38Z","title":"Layer-wise Cross-Lingual Depression Detection from Speech: Analysis with Contrastive Alignment","version":1},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-07-12T06:06:46.343211Z"},"links":{"citing_paper":"/paper/2607.02920"},"observation_digest":"sha256:8efb72cfcc71fec6044b6c20cb32862de2e73ee0416e9ba55ba70288bd494e32","observation_id":"b0b89fd0-ea8e-43a6-88f6-6064f197a061","resolution":{"observed_at":"2026-07-12T06:06:46.343211Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2209.05355","last_updated":"2023-09-20T20:20:45Z","snapshot_observed_at":"2026-08-20T01:11:09.278575Z","submitted_at":"2022-09-12T16:06:10Z","title":"Analysis and Comparison of Classification Metrics","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2209.05355","snapshot_observed_at":"2026-07-12T06:06:46.343211Z","title":"Analysis and comparison of classification metrics,","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2607.02920","last_updated":"2026-07-03T03:24:38Z","snapshot_observed_at":"2026-08-19T17:04:02.832931Z","submitted_at":"2026-07-03T03:24:38Z","title":"Layer-wise Cross-Lingual Depression Detection from Speech: Analysis with Contrastive Alignment","version":1},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-07-12T06:06:46.343211Z"},"links":{"cited_paper":"/paper/2209.05355","citing_paper":"/paper/2607.02920"},"observation_digest":"sha256:d2abf820fd7a2dc83cd9a4c140d88e763684834fe8451708ec2d4dc735700f5c","observation_id":"75d8fff7-cb43-47a9-8de8-8bb814179c2c","resolution":{"observed_at":"2026-07-12T06:06:46.343211Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-12T06:06:46.343211Z","title":"wav2vec 2.0: A framework for self-supervised learning of speech representations,","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2607.02920","last_updated":"2026-07-03T03:24:38Z","snapshot_observed_at":"2026-08-19T17:04:02.832931Z","submitted_at":"2026-07-03T03:24:38Z","title":"Layer-wise Cross-Lingual Depression Detection from Speech: Analysis with Contrastive Alignment","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-07-12T06:06:46.343211Z"},"links":{"citing_paper":"/paper/2607.02920"},"observation_digest":"sha256:9ec51bf0546aab5d88b153fe057840c39e8991aa14d30a7dc1cb3261d737abde","observation_id":"4c8f3b0f-a948-483c-83e9-b255c823df41","resolution":{"observed_at":"2026-07-12T06:06:46.343211Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-12T06:06:46.343211Z","title":"A VEC 2016: Depression, mood, and emotion recognition workshop and challenge,","venue":null,"work_id":null,"year":2016},"citing_paper":{"arxiv_id":"2607.02920","last_updated":"2026-07-03T03:24:38Z","snapshot_observed_at":"2026-08-19T17:04:02.832931Z","submitted_at":"2026-07-03T03:24:38Z","title":"Layer-wise Cross-Lingual Depression Detection from Speech: Analysis with Contrastive Alignment","version":1},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-07-12T06:06:46.343211Z"},"links":{"citing_paper":"/paper/2607.02920"},"observation_digest":"sha256:47f0540b1a2d3dd7d212295b01159256cf88de7cddc68c37ff4d51eb5df87675","observation_id":"796811d0-4d51-4176-9d0d-753194763b05","resolution":{"observed_at":"2026-07-12T06:06:46.343211Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-07-12T06:06:46.343211Z","title":"A VEC 2019 workshop and challenge: State-mindfulness and continuous emotion recognition,","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2607.02920","last_updated":"2026-07-03T03:24:38Z","snapshot_observed_at":"2026-08-19T17:04:02.832931Z","submitted_at":"2026-07-03T03:24:38Z","title":"Layer-wise Cross-Lingual Depression Detection from Speech: Analysis with Contrastive Alignment","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-07-12T06:06:46.343211Z"},"links":{"citing_paper":"/paper/2607.02920"},"observation_digest":"sha256:1bf58e2eee5b087432cebc9b28bd34ef12a31ae4185497ff6ce3bb0f5638e1ed","observation_id":"ae629191-8809-440a-8978-5c7c77ea0115","resolution":{"observed_at":"2026-07-12T06:06:46.343211Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"paper":{"arxiv_id":"2607.02920","last_updated":"2026-07-03T03:24:38Z","latest_version":1,"primary_category":"eess.AS","snapshot_observed_at":"2026-08-19T17:04:02.832931Z","submitted_at":"2026-07-03T03:24:38Z","title":"Layer-wise Cross-Lingual Depression Detection from Speech: Analysis with Contrastive Alignment"},"reference_resolution":{"displayed":25,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":25,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":25},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-20T06:33:59.587034+00:00","source":"crossref"},{"observed_at":"2026-08-20T06:33:54.927442+00:00","source":"retraction_watch"}],"thesis":"As of 20 August 2026, this Paper Citation Record lists 25 of 25 outbound references and 0 inbound Pith citation observations for arXiv:2607.02920."}