{"as_of":"2026-08-10T05:01:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:bcf0134512ee6ba2b7bc869303128737b2c3e7c69c3174850a3df3ea042b52e7","coverage":[{"denominator":138,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":100,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-06T21:29:05.314870Z","state":"measured"},{"denominator":103,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":103,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-09T06:31:02.800959+00:00","state":"measured"},{"denominator":3,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":3,"source":"paper_references, paper_reference_links","source_observed_at":"2026-06-30T01:29:58.474810Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":1,"source":"arxiv_reference","source_observed_at":"2026-08-05T02:28:24.338817Z","state":"measured"}],"external_citation_measurements":[{"count":0,"observed_at":"2026-08-05T02:28:24.338817Z","source":"arxiv_reference"}],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"cited_work":{"arxiv_id":"2506.24124","doi":"10.48550/arxiv.2506.24124","metadata_source":"arxiv_reference","pith_arxiv_id":"2506.24124","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"arXiv preprint arXiv:2506.24124 , year=","venue":"ArXiv.org","work_id":"b13601c4-5891-43bd-9914-51d6546973b7","year":2025},"citing_paper":{"arxiv_id":"2605.25943","last_updated":"2026-05-25T15:21:06Z","snapshot_observed_at":"2026-07-06T23:35:50.787324Z","submitted_at":"2026-05-25T15:21:06Z","title":"STaT: Resolving Shape Distortion in Non-Stationary Time Series via Tri-Modal Synergy","version":1},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-06-29T22:17:53.624667Z"},"links":{"cited_paper":"/paper/2506.24124","citing_paper":"/paper/2605.25943"},"observation_digest":"sha256:0915ed5057c2e7baa1f4d15cd3317893a05417ca966e0d6bdbea504b4edbc064","observation_id":"afd912c9-50f1-41cc-b294-7df8c4ad6a5e","resolution":{"observed_at":"2026-06-29T22:24:00.656941Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"cited_work":{"arxiv_id":"2506.24124","doi":"10.48550/arxiv.2506.24124","metadata_source":"arxiv_reference","pith_arxiv_id":"2506.24124","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"arXiv preprint arXiv:2506.24124 , year=","venue":"ArXiv.org","work_id":"b13601c4-5891-43bd-9914-51d6546973b7","year":2025},"citing_paper":{"arxiv_id":"2606.18986","last_updated":"2026-06-17T12:07:23Z","snapshot_observed_at":"2026-08-06T07:20:38.551006Z","submitted_at":"2026-06-17T12:07:23Z","title":"Beyond Tokenization: Direct Timestep Embedding and Contrastive Alignment for Time-Series Question Answering","version":1},"reference_index":24,"source":"arxiv_source","source_observed_at":"2026-06-26T20:42:50.385435Z"},"links":{"cited_paper":"/paper/2506.24124","citing_paper":"/paper/2606.18986"},"observation_digest":"sha256:5dc79923751fe74aa838df2b7ef0940352dc7f1465a553b9f8aeeaee535563c7","observation_id":"5c954d86-2670-4dce-bdaf-02a390bd3724","resolution":{"observed_at":"2026-07-04T01:09:18.666026Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"cited_work":{"arxiv_id":"2506.24124","doi":"10.48550/arxiv.2506.24124","metadata_source":"arxiv_reference","pith_arxiv_id":"2506.24124","snapshot_observed_at":"2026-08-05T02:28:24.338817Z","title":"arXiv preprint arXiv:2506.24124 , year=","venue":"ArXiv.org","work_id":"b13601c4-5891-43bd-9914-51d6546973b7","year":2025},"citing_paper":{"arxiv_id":"2606.28446","last_updated":"2026-06-26T08:35:55Z","snapshot_observed_at":"2026-08-07T10:51:25.338939Z","submitted_at":"2026-06-26T08:35:55Z","title":"Domain-Informed Multi-View Self-Distillation for Astronomical Light-Curve Representation Learning with JEPA","version":1},"reference_index":46,"source":"arxiv_source","source_observed_at":"2026-06-30T01:29:58.474810Z"},"links":{"cited_paper":"/paper/2506.24124","citing_paper":"/paper/2606.28446"},"observation_digest":"sha256:b68f3b7f6a68433e8821c2e161e95367b022b2ddb49f3f05d6c496726a22d896","observation_id":"7be0cf6d-ed19-4a25-aa4a-7bcbce33dec8","resolution":{"observed_at":"2026-06-30T01:34:09.349309Z","resolver_source":"arxiv_id","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}}],"links":{"evidence":"/evidence","html":"/paper/2506.24124/citation-record","integrity":"/paper/2506.24124/integrity","json":"/paper/2506.24124/citation-record.json","paper":"/paper/2506.24124"},"outbound":[{"citation":{"cited_paper":{"arxiv_id":"2403.07815","last_updated":"2024-11-04T17:42:45Z","snapshot_observed_at":"2026-07-06T17:43:27.034067Z","submitted_at":"2024-03-12T16:53:54Z","title":"Chronos: Learning the Language of Time Series","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.07815","snapshot_observed_at":"2026-08-06T21:28:52.573642Z","title":"Maddix, Hao Wang, Michael W","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":1,"source":"pdf_text","source_observed_at":"2026-08-06T21:28:52.573642Z"},"links":{"cited_paper":"/paper/2403.07815","citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:4937af3334c5bc901b904542f208bc5517b8be403399cbc8850059c136e753fc","observation_id":"12aad24f-4c1e-406c-b27d-1b7eb409f167","resolution":{"observed_at":"2026-08-06T21:28:52.573642Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1607.06450","last_updated":"2016-07-21T19:57:52Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2016-07-21T19:57:52Z","title":"Layer Normalization","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1607.06450","snapshot_observed_at":"2026-08-06T21:28:52.669344Z","title":"Layer normalization","venue":null,"work_id":null,"year":2016},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":2,"source":"pdf_text","source_observed_at":"2026-08-06T21:28:52.669344Z"},"links":{"cited_paper":"/paper/1607.06450","citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:1015e1a8e566dc516ed6cc5f0200873149bc681179d52cdeb7cad11a66705324","observation_id":"e4672016-a5af-41c7-88cd-de48d2772c31","resolution":{"observed_at":"2026-08-06T21:28:52.669344Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:28:52.820203Z","title":"Privacy preserving generative feature transformation","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":3,"source":"pdf_text","source_observed_at":"2026-08-06T21:28:52.820203Z"},"links":{"citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:ae1cd8f2753a7de4f745bd524d7202a7ae256da64532b76f4ef8e1bae6f0e83c","observation_id":"ff9ee70f-bbe6-4161-929c-375b1165a6cf","resolution":{"observed_at":"2026-08-06T21:28:52.820203Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:28:52.933313Z","title":"Gorec: a generative cold-start recommendation framework","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":4,"source":"pdf_text","source_observed_at":"2026-08-06T21:28:52.933313Z"},"links":{"citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:fb5068e43329987b77c4ca3c491fe04769aa1d28e7005c5a697e83b48fab0049","observation_id":"b2d385e1-2244-432c-b160-6c549fbb885a","resolution":{"observed_at":"2026-08-06T21:28:52.933313Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:28:53.085687Z","title":"Multimodality invariant learning for multimedia-based new item recommendation","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":5,"source":"pdf_text","source_observed_at":"2026-08-06T21:28:53.085687Z"},"links":{"citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:62e6e00a00887dd648862f9f006230c99883f7b4d438f815c6bc0d336618b22c","observation_id":"ad68d57d-d7fb-4549-8bea-0376225aa81f","resolution":{"observed_at":"2026-08-06T21:28:53.085687Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2505.18204","last_updated":"2025-05-21T14:11:46Z","snapshot_observed_at":"2026-08-09T08:26:36.162707Z","submitted_at":"2025-05-21T14:11:46Z","title":"Brownian Bridge Augmented Surrogate Simulation and Injection Planning for Geological CO$_2$ Storage","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2505.18204","snapshot_observed_at":"2026-08-06T21:28:53.188562Z","title":"Brownian bridge augmented surrogate simulation and injection planning for geological co _2 storage","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-08-06T21:28:53.188562Z"},"links":{"cited_paper":"/paper/2505.18204","citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:0ce17180ebaa24da0cff8ffec4eb787562381eba24c8de9f109cdc5687cad5a9","observation_id":"0889fec0-b78d-4e49-a1a1-f61225b3fb93","resolution":{"observed_at":"2026-08-06T21:28:53.188562Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:28:53.303227Z","title":"Deep learning and time series-to-image encoding for finan- cial forecasting","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":7,"source":"pdf_text","source_observed_at":"2026-08-06T21:28:53.303227Z"},"links":{"citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:b59ebb7f4ce576b6e8a7c77b74392bc3dd149ec2dc210c7d0b106476c7007e4f","observation_id":"3ea3dc0a-3a32-4928-83f7-f4879170aa3c","resolution":{"observed_at":"2026-08-06T21:28:53.303227Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:28:53.429090Z","title":"Fundamental limitations of foundational forecasting models: The need for multimodality and rigorous evaluation","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-06T21:28:53.429090Z"},"links":{"citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:ff4187e47bda3414815e53d111299a478e90dd83d5dc28789a009d75a91efeea","observation_id":"7620ff6f-c6f1-484d-b88d-a3495ed39952","resolution":{"observed_at":"2026-08-06T21:28:53.429090Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:28:53.591316Z","title":"Control charts in financial applications: An overview","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":9,"source":"pdf_text","source_observed_at":"2026-08-06T21:28:53.591316Z"},"links":{"citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:2b0b38ed32723e10a03a0f6eb64fcb74194826dd1e8a6f811a9b73139cb5e7a2","observation_id":"cb1d611f-4fea-4843-b258-3e9ee3ccafad","resolution":{"observed_at":"2026-08-06T21:28:53.591316Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:28:53.719457Z","title":"Language models are few-shot learners","venue":null,"work_id":null,"year":1901},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":10,"source":"pdf_text","source_observed_at":"2026-08-06T21:28:53.719457Z"},"links":{"citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:1646538ad4348dc543d521c435193d4d1ed1b91bbcd12509a08b868c6e30cdf3","observation_id":"a5f2d7b9-7b50-4880-bcd4-d23756c7c686","resolution":{"observed_at":"2026-08-06T21:28:53.719457Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:28:53.898252Z","title":"Time series forecasting for healthcare diagnosis and prognostics with the focus on cardiovascular diseases","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-08-06T21:28:53.898252Z"},"links":{"citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:d1d63ecd5d38e7f56841506a4c906c575b3db78147a1559ab77491699eac3221","observation_id":"48333e52-ec0f-4a8a-9809-899cd7f2d88f","resolution":{"observed_at":"2026-08-06T21:28:53.898252Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2405.16083","last_updated":"2024-05-25T06:26:02Z","snapshot_observed_at":"2026-07-06T18:19:42.946629Z","submitted_at":"2024-05-25T06:26:02Z","title":"From Orthogonality to Dependency: Learning Disentangled Representation for Multi-Modal Time-Series Sensing Signals","version":1},"cited_work":{"arxiv_id":"2405.16083","doi":null,"metadata_source":"pith","pith_arxiv_id":"2405.16083","snapshot_observed_at":"2026-08-06T21:29:13.800219Z","title":"From Orthogonality to Dependency: Learning Disentangled Representation for Multi-Modal Time-Series Sensing Signals","venue":"cs.LG","work_id":"f53019f9-0e5c-4fe3-86d9-31662c83fa20","year":2024},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":12,"source":"pdf_text","source_observed_at":"2026-08-06T21:28:54.021165Z"},"links":{"cited_paper":"/paper/2405.16083","citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:8b235bc15dc95175de25f3234d77560cb91dcbcb177aa56ce8818081111858b3","observation_id":"c9cbc451-dea2-4878-a69d-1f5034779d70","resolution":{"observed_at":"2026-08-06T21:29:13.831468Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:28:54.170478Z","title":"Lightts: Lightweight time series classification with adaptive ensemble distillation","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":13,"source":"pdf_text","source_observed_at":"2026-08-06T21:28:54.170478Z"},"links":{"citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:d2f5aa2eec763556b952e0dd3ede9557939a4d9aa2f00f2beb19a2115b5493b7","observation_id":"3aea84ee-749e-4558-82b7-d8de20d84d55","resolution":{"observed_at":"2026-08-06T21:28:54.170478Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2207.10362","last_updated":"2022-07-21T08:43:51Z","snapshot_observed_at":"2026-07-06T13:33:43.048614Z","submitted_at":"2022-07-21T08:43:51Z","title":"LocVTP: Video-Text Pre-training for Temporal Localization","version":1},"cited_work":{"arxiv_id":"2207.10362","doi":null,"metadata_source":"pith","pith_arxiv_id":"2207.10362","snapshot_observed_at":"2026-08-06T21:29:13.601719Z","title":"LocVTP: Video-Text Pre-training for Temporal Localization","venue":"cs.CV","work_id":"5ede4ca7-04f5-4199-98c3-3d60a5dcebd2","year":2022},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-08-06T21:28:54.312848Z"},"links":{"cited_paper":"/paper/2207.10362","citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:0811465ad0fc21cdf9a6caba10e31afe586d9727468a655c9aa8f5ce1c867ef4","observation_id":"2ce06a8f-22ac-4381-9d57-6b96aaa98ed3","resolution":{"observed_at":"2026-08-06T21:29:13.744649Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:28:54.438374Z","title":"Nhits: Neural hierarchical interpolation for time series forecasting","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-06T21:28:54.438374Z"},"links":{"citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:954c5bd8957d5cb6a3d9f19e64449906f744d17baa69d5e74701dd43efcf5897","observation_id":"fa67aee7-565e-495a-9fe2-7ae42c9ee75e","resolution":{"observed_at":"2026-08-06T21:28:54.438374Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:28:54.633314Z","title":"Multi- model approach for stock price prediction and trading recommendations","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":16,"source":"pdf_text","source_observed_at":"2026-08-06T21:28:54.633314Z"},"links":{"citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:32dc8741263cd1e93464eed9dc0c4fe066021a5dd9f2f3fe3d026adbb5859dc0","observation_id":"66be7787-6e16-469f-84fd-7187b163156b","resolution":{"observed_at":"2026-08-06T21:28:54.633314Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:28:54.924122Z","title":"Financial time series forecasting with multi-modality graph neural network","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":17,"source":"pdf_text","source_observed_at":"2026-08-06T21:28:54.924122Z"},"links":{"citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:27332d1060bdbc425e22ac5ec7d0c20f576377ba6e30133a35492032528d3a52","observation_id":"f4199341-2d8e-45ab-b958-1587667f075a","resolution":{"observed_at":"2026-08-06T21:28:54.924122Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2304.08424","last_updated":"2024-04-04T16:24:19Z","snapshot_observed_at":"2026-08-10T01:51:43.637808Z","submitted_at":"2023-04-17T16:46:48Z","title":"Long-term Forecasting with TiDE: Time-series Dense Encoder","version":5},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2304.08424","snapshot_observed_at":"2026-08-06T21:28:55.063158Z","title":"Long-term forecasting with tide: Time-series dense encoder","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":18,"source":"pdf_text","source_observed_at":"2026-08-06T21:28:55.063158Z"},"links":{"cited_paper":"/paper/2304.08424","citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:5b3f2f643244dfe2181c6131265c4872b3c0e3092d1dd59226fd431f9e72d1ce","observation_id":"965b7b11-eece-4c84-b22e-5c464ee84d34","resolution":{"observed_at":"2026-08-06T21:28:55.063158Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.02475","last_updated":"2024-06-07T07:46:26Z","snapshot_observed_at":"2026-08-01T20:24:08.204721Z","submitted_at":"2024-02-04T13:10:51Z","title":"TimeSiam: A Pre-Training Framework for Siamese Time-Series Modeling","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.02475","snapshot_observed_at":"2026-08-06T21:28:55.194907Z","title":"Timesiam: A pre-training framework for siamese time-series modeling","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-06T21:28:55.194907Z"},"links":{"cited_paper":"/paper/2402.02475","citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:37f9143ce720d03cb0c0d9aa0b8d6ca5def82635a80c7ee2acc0a3285b4ad8b5","observation_id":"1498ee0f-abdc-460f-b88e-4a487efcc00d","resolution":{"observed_at":"2026-08-06T21:28:55.194907Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:28:55.292212Z","title":"Weakly supervised video representation learning with unaligned text for sequential videos","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":20,"source":"pdf_text","source_observed_at":"2026-08-06T21:28:55.292212Z"},"links":{"citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:2af56906be12cf5fafa603cdc16f5b58d090186364fa50a678e9407d07ade3d4","observation_id":"6605f5ff-d14f-44bc-a0d4-b2498d82caf4","resolution":{"observed_at":"2026-08-06T21:28:55.292212Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2010.11929","last_updated":"2021-06-03T13:08:56Z","snapshot_observed_at":"2026-08-10T01:12:16.468283Z","submitted_at":"2020-10-22T17:55:59Z","title":"An Image is Worth 16x16 Words: Transformers for Image Recognition at Scale","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2010.11929","snapshot_observed_at":"2026-08-06T21:28:55.471004Z","title":"An image is worth 16x16 words: Transformers for image recognition at scale","venue":null,"work_id":null,"year":2010},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":21,"source":"pdf_text","source_observed_at":"2026-08-06T21:28:55.471004Z"},"links":{"cited_paper":"/paper/2010.11929","citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:69233904c78dc0d818bf4709026981926076c1a517f3a464d57eb6cee0a7b04b","observation_id":"8558c4ef-66fc-4f02-83bf-3fe25524dd33","resolution":{"observed_at":"2026-08-06T21:28:55.471004Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.16556","last_updated":"2024-04-24T17:37:52Z","snapshot_observed_at":"2026-07-06T15:33:42.279377Z","submitted_at":"2023-05-26T00:50:09Z","title":"LANISTR: Multimodal Learning from Structured and Unstructured Data","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.16556","snapshot_observed_at":"2026-08-06T21:28:55.619007Z","title":"Lanistr: Multimodal learning from structured and unstructured data","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":22,"source":"pdf_text","source_observed_at":"2026-08-06T21:28:55.619007Z"},"links":{"cited_paper":"/paper/2305.16556","citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:a4271d42a24a9a6e6b1c323518f35db4eb26ce7d42fcddb0b48385003b23cb58","observation_id":"a551c215-e2d7-4598-8764-5cd2377f00f4","resolution":{"observed_at":"2026-08-06T21:28:55.619007Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:28:55.745469Z","title":"Unsupervised scalable repre- sentation learning for multivariate time series","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-08-06T21:28:55.745469Z"},"links":{"citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:dace2ff4bff287d49b9ebe33372a73caa15aebc3c8f00fdae3c8e00e67e04711","observation_id":"f8e67e57-8267-40b8-941d-9cf3c8b00189","resolution":{"observed_at":"2026-08-06T21:28:55.745469Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2505.15076","last_updated":"2025-05-21T03:49:24Z","snapshot_observed_at":"2026-08-08T23:22:42.007781Z","submitted_at":"2025-05-21T03:49:24Z","title":"Agentic Feature Augmentation: Unifying Selection and Generation with Teaming, Planning, and Memories","version":1},"cited_work":{"arxiv_id":"2505.15076","doi":null,"metadata_source":"pith","pith_arxiv_id":"2505.15076","snapshot_observed_at":"2026-08-06T21:29:12.875192Z","title":"Agentic Feature Augmentation: Unifying Selection and Generation with Teaming, Planning, and Memories","venue":"cs.LG","work_id":"43b9cdaf-8af5-4948-ab38-7da32f569e4a","year":2025},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":24,"source":"pdf_text","source_observed_at":"2026-08-06T21:28:55.860154Z"},"links":{"cited_paper":"/paper/2505.15076","citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:cbfead6c5ab68ded3aefa3330a1b7ade22c5dbdc9507d4fd8d854278bdc3d470","observation_id":"d3b0f1d2-bdd1-4d87-9c78-c8024e18cb72","resolution":{"observed_at":"2026-08-06T21:29:13.008545Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2505.15152","last_updated":"2025-05-21T06:18:42Z","snapshot_observed_at":"2026-08-08T23:23:01.845391Z","submitted_at":"2025-05-21T06:18:42Z","title":"Sculpting Features from Noise: Reward-Guided Hierarchical Diffusion for Task-Optimal Feature Transformation","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2505.15152","snapshot_observed_at":"2026-08-06T21:28:55.930827Z","title":"Sculpting features from noise: Reward-guided hierarchical diffusion for task-optimal feature transformation","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-06T21:28:55.930827Z"},"links":{"cited_paper":"/paper/2505.15152","citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:cf02db309eb2dbaaebe63d0833658b2c0c1acb9453f470594361c364ed81d308","observation_id":"2402a515-8451-4756-a82f-890d49f7306e","resolution":{"observed_at":"2026-08-06T21:28:55.930827Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:28:56.052319Z","title":"Evolutionary large language model for automated feature transformation","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":26,"source":"pdf_text","source_observed_at":"2026-08-06T21:28:56.052319Z"},"links":{"citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:a58e10188271050131d76d528e746c029880b8c8496e084569d6fc01108e25e0","observation_id":"a3dec871-0c2a-4322-a69f-a47baacbfa58","resolution":{"observed_at":"2026-08-06T21:28:56.052319Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2504.21304","last_updated":"2025-04-30T04:26:03Z","snapshot_observed_at":"2026-08-07T15:57:55.449740Z","submitted_at":"2025-04-30T04:26:03Z","title":"Unsupervised Feature Transformation via In-context Generation, Generator-critic LLM Agents, and Duet-play Teaming","version":1},"cited_work":{"arxiv_id":"2504.21304","doi":null,"metadata_source":"pith","pith_arxiv_id":"2504.21304","snapshot_observed_at":"2026-08-06T21:29:12.645705Z","title":"Unsupervised Feature Transformation via In-context Generation, Generator-critic LLM Agents, and Duet-play Teaming","venue":"cs.LG","work_id":"4798feae-4239-44a4-9f8f-47e453952fe8","year":2025},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":27,"source":"pdf_text","source_observed_at":"2026-08-06T21:28:56.143876Z"},"links":{"cited_paper":"/paper/2504.21304","citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:a076c254fc49ce5346174cdbe5519e522863feb4d849e2dce0e88adaf72a30d6","observation_id":"dc277e5c-1e83-4425-a8bb-106ec7f42b40","resolution":{"observed_at":"2026-08-06T21:29:12.715332Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:28:56.260332Z","title":"Neuro-symbolic embedding for short and effective feature selection via autoregressive generation","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":28,"source":"pdf_text","source_observed_at":"2026-08-06T21:28:56.260332Z"},"links":{"citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:201f73bfb016d3d94c4ec7c17f9c333e0b60a59fafb6d66516647f02f02841e5","observation_id":"357861e3-64f2-4712-ad87-6460106ed578","resolution":{"observed_at":"2026-08-06T21:28:56.260332Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.03885","last_updated":"2024-10-10T15:37:45Z","snapshot_observed_at":"2026-08-07T22:30:00.741959Z","submitted_at":"2024-02-06T10:48:46Z","title":"MOMENT: A Family of Open Time-series Foundation Models","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.03885","snapshot_observed_at":"2026-08-06T21:28:56.368581Z","title":"Moment: a family of open time-series foundation models","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":29,"source":"pdf_text","source_observed_at":"2026-08-06T21:28:56.368581Z"},"links":{"cited_paper":"/paper/2402.03885","citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:ad1101b2eac1e326d5e33c42e2e9dae3c156675949aa3680cd1c14a72518cf9e","observation_id":"ff69d990-fea8-4712-a214-5936477b5c6f","resolution":{"observed_at":"2026-08-06T21:28:56.368581Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:28:56.457907Z","title":"Large language models are zero-shot time series forecasters","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-08-06T21:28:56.457907Z"},"links":{"citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:c99e500040534f7cae982f8e97296f8dc1eeca964c89a6ad560ee60ec5aea940","observation_id":"3a5dc60a-625c-468c-8e54-5c264a936fc2","resolution":{"observed_at":"2026-08-06T21:28:56.457907Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2111.00396","last_updated":"2022-08-05T17:54:38Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2021-10-31T03:32:18Z","title":"Efficiently Modeling Long Sequences with Structured State Spaces","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2111.00396","snapshot_observed_at":"2026-08-06T21:28:56.621717Z","title":"Efficiently modeling long sequences with structured state spaces","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":31,"source":"pdf_text","source_observed_at":"2026-08-06T21:28:56.621717Z"},"links":{"cited_paper":"/paper/2111.00396","citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:a42cae51b9e7649ed5509a5a0421f8ad222b86648507347718427c1aed486a64","observation_id":"dd212c19-eb43-465f-94af-a2fd243970d1","resolution":{"observed_at":"2026-08-06T21:28:56.621717Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:28:56.782411Z","title":"Audioclip: Extending clip to image, text and audio","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":32,"source":"pdf_text","source_observed_at":"2026-08-06T21:28:56.782411Z"},"links":{"citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:501a39cb34973250dd832e404253a700040f447198637c8faa9a64a96040a762","observation_id":"e0554412-4d71-469c-ac49-7eb65ca4e259","resolution":{"observed_at":"2026-08-06T21:28:56.782411Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:28:56.881025Z","title":"Temporal alignment networks for long- term video","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":33,"source":"pdf_text","source_observed_at":"2026-08-06T21:28:56.881025Z"},"links":{"citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:8ac47a466d8b586be4b293dc6e84e0be07576c8572c06ee2b3354725cdc754b3","observation_id":"93236bc6-6cba-4335-ab44-958f13b2417b","resolution":{"observed_at":"2026-08-06T21:28:56.881025Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2403.14608","last_updated":"2024-09-16T02:54:50Z","snapshot_observed_at":"2026-08-04T09:07:42.158421Z","submitted_at":"2024-03-21T17:55:50Z","title":"Parameter-Efficient Fine-Tuning for Large Models: A Comprehensive Survey","version":7},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2403.14608","snapshot_observed_at":"2026-08-06T21:28:57.009353Z","title":"Parameter-efficient fine-tuning for large models: A comprehensive survey","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":34,"source":"pdf_text","source_observed_at":"2026-08-06T21:28:57.009353Z"},"links":{"cited_paper":"/paper/2403.14608","citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:0e119484363fe19c28782b857beca9a3cfd7b66f67d2b3cb5a6f0eeddb5bd019","observation_id":"745fe211-ea90-42fa-b86a-b6f6ad5b3695","resolution":{"observed_at":"2026-08-06T21:28:57.009353Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:28:57.075192Z","title":"Deep residual learning for image recognition","venue":null,"work_id":null,"year":2016},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":35,"source":"pdf_text","source_observed_at":"2026-08-06T21:28:57.075192Z"},"links":{"citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:d69b3667e95615791ebc3a2f932b0598993861db2a4e30944914b9aa3a57ef64","observation_id":"81dba392-922a-4938-a472-5a701c5f2bcd","resolution":{"observed_at":"2026-08-06T21:28:57.075192Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:28:57.204592Z","title":"Double correction framework for denoising recommendation","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":36,"source":"pdf_text","source_observed_at":"2026-08-06T21:28:57.204592Z"},"links":{"citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:6b229d7bf60c63c9ce33c5c2e32a1265d11abbba5eade2227da0c0ff52fd20cc","observation_id":"3035ca14-c920-40dc-a34f-541a7b0f77b4","resolution":{"observed_at":"2026-08-06T21:28:57.204592Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:28:57.354424Z","title":"Long short-term memory","venue":null,"work_id":null,"year":1997},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-06T21:28:57.354424Z"},"links":{"citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:0609bbc3e2780daad99c93024f383b4313592c0ddcbf6f157b363b229a00b901","observation_id":"4a538c54-e9e8-492e-98d0-c9320fdadd15","resolution":{"observed_at":"2026-08-06T21:28:57.354424Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:28:57.425806Z","title":"Transrac: Encoding multi-scale temporal correlation with transformers for repetitive action counting","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":38,"source":"pdf_text","source_observed_at":"2026-08-06T21:28:57.425806Z"},"links":{"citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:2ef5a2622f0567629b69fd721eba209b49166b27e908407e4c26d56c3b07f265","observation_id":"23724760-8b4d-4a06-956c-f6d1a145cea9","resolution":{"observed_at":"2026-08-06T21:28:57.425806Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:28:57.540299Z","title":"Ct-patchtst: Channel-time patch time-series transformer for long-term renewable energy forecasting","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":39,"source":"pdf_text","source_observed_at":"2026-08-06T21:28:57.540299Z"},"links":{"citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:9bfec54bc00c179682b199a8dea55682214d618cc1c4c747ccc097e6a7b392cf","observation_id":"df2c1207-45b9-465a-8f84-55e5febb6d0e","resolution":{"observed_at":"2026-08-06T21:28:57.540299Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:28:57.622463Z","title":"Gpt4mts: Prompt-based large language model for multimodal time-series forecasting","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":40,"source":"pdf_text","source_observed_at":"2026-08-06T21:28:57.622463Z"},"links":{"citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:b5c388393af6673aa203a53a46db8a2df84a96b5645bd911c8efeecc3c89bed8","observation_id":"6c5fca9c-7645-485b-84fc-7d85708a9101","resolution":{"observed_at":"2026-08-06T21:28:57.622463Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.01728","last_updated":"2024-01-29T06:27:53Z","snapshot_observed_at":"2026-08-04T04:31:27.172482Z","submitted_at":"2023-10-03T01:31:25Z","title":"Time-LLM: Time Series Forecasting by Reprogramming Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.01728","snapshot_observed_at":"2026-08-06T21:28:57.826466Z","title":"Zhang, Xiaoming Shi, Pin-Yu Chen, Yuxuan Liang, Yuan-Fang Li, Shirui Pan, and Qingsong Wen","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":42,"source":"pdf_text","source_observed_at":"2026-08-06T21:28:57.826466Z"},"links":{"cited_paper":"/paper/2310.01728","citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:8052dce9c0bbd9f280dd113e06da0e7bccae3e067bc79e47d0a587c667b36576","observation_id":"09c60e33-9d1b-4142-b086-c2ad391f4f9a","resolution":{"observed_at":"2026-08-06T21:28:57.826466Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:28:57.925912Z","title":"Position: What can large language models tell us about time series analysis","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":43,"source":"pdf_text","source_observed_at":"2026-08-06T21:28:57.925912Z"},"links":{"citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:6cd25f060008fd53ec2b95b41a34f27b19d8e08ad1866a1869eb4ca5ac53344c","observation_id":"ebcc2524-0b47-4275-a195-439642c21a63","resolution":{"observed_at":"2026-08-06T21:28:57.925912Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:28:58.046606Z","title":"Ai in healthcare: time-series forecasting using statistical, neural, and ensemble architectures","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":44,"source":"pdf_text","source_observed_at":"2026-08-06T21:28:58.046606Z"},"links":{"citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:20850a1f897a359b1c0258fdbb1a7cb341bc987b1f60c9a8738a2d894d55c79c","observation_id":"cc426686-a8be-4e6d-ad4b-72fc1a8a1171","resolution":{"observed_at":"2026-08-06T21:28:58.046606Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:28:58.175064Z","title":"Bert: Pre-training of deep bidirectional transformers for language understanding","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":45,"source":"pdf_text","source_observed_at":"2026-08-06T21:28:58.175064Z"},"links":{"citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:01386039276f3a328230ba8f8237c8527ead6802f2bd46202c39b54a2449918e","observation_id":"cd0a5b4b-5790-4b3b-b2a2-47b0999c8f43","resolution":{"observed_at":"2026-08-06T21:28:58.175064Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2001.04451","last_updated":"2020-02-18T16:01:18Z","snapshot_observed_at":"2026-07-06T08:50:12.690900Z","submitted_at":"2020-01-13T18:38:28Z","title":"Reformer: The Efficient Transformer","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2001.04451","snapshot_observed_at":"2026-08-06T21:28:58.297086Z","title":"Reformer: The efficient transformer","venue":null,"work_id":null,"year":2001},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":46,"source":"pdf_text","source_observed_at":"2026-08-06T21:28:58.297086Z"},"links":{"cited_paper":"/paper/2001.04451","citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:fc324a5ed23d4d1e79a73543da4b6955ec887a8297080df8d06adc9fba6a295f","observation_id":"0549d233-f528-47fe-ab19-7c8a69a14404","resolution":{"observed_at":"2026-08-06T21:28:58.297086Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2404.01165","last_updated":"2024-08-10T12:03:44Z","snapshot_observed_at":"2026-07-06T17:54:00.048961Z","submitted_at":"2024-04-01T15:14:07Z","title":"LITE: Modeling Environmental Ecosystems with Multimodal Large Language Models","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2404.01165","snapshot_observed_at":"2026-08-06T21:28:58.424037Z","title":"Lite: Modeling environmental ecosystems with multimodal large language models","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":47,"source":"pdf_text","source_observed_at":"2026-08-06T21:28:58.424037Z"},"links":{"cited_paper":"/paper/2404.01165","citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:aa9cdf47daeb19ce8f987d32c8005a2ed405e39d458f6413ef51b1c25450b6a9","observation_id":"60be177e-e1d4-4a3b-853f-55cb20291c3e","resolution":{"observed_at":"2026-08-06T21:28:58.424037Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:28:58.573179Z","title":"Sehf: A summary- enhanced hierarchical framework for financial report sentiment analysis","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-06T21:28:58.573179Z"},"links":{"citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:6afc9b7dece9adaf64c3242295b1710001888ee91bdbea037082fb4d9eebe52f","observation_id":"85681163-5902-4ad8-9c2e-222e9aa2137e","resolution":{"observed_at":"2026-08-06T21:28:58.573179Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:28:58.681043Z","title":"Sade: A speaker- aware dual encoding model based on diagbert for medical triage and pre-diagnosis","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":49,"source":"pdf_text","source_observed_at":"2026-08-06T21:28:58.681043Z"},"links":{"citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:c837c200ede8a37f340a3279f9b3b0faa309ea804e3fdeab878f007f4e5789b3","observation_id":"c696f405-ae9b-4d0e-9541-322d02010e7f","resolution":{"observed_at":"2026-08-06T21:28:58.681043Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:28:58.804030Z","title":"Frozen language model helps ecg zero-shot learning","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":50,"source":"pdf_text","source_observed_at":"2026-08-06T21:28:58.804030Z"},"links":{"citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:75f7cb3f4efc306b37873f8934ec46a3774db5083be30b5e6fc95733c4322023","observation_id":"24431474-6f97-4885-a0ff-67268a44f9be","resolution":{"observed_at":"2026-08-06T21:28:58.804030Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:28:58.931484Z","title":"Blip-2: Bootstrapping language- image pre-training with frozen image encoders and large language models","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":51,"source":"pdf_text","source_observed_at":"2026-08-06T21:28:58.931484Z"},"links":{"citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:f998e374dfdd4154ee6da2ae896d876886087ec1c5d06869ee5f2d57f3521d88","observation_id":"b130c929-f07e-4c8a-830d-9c9e01f1e9e1","resolution":{"observed_at":"2026-08-06T21:28:58.931484Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.06355","last_updated":"2024-01-04T02:06:07Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-05-10T17:59:04Z","title":"VideoChat: Chat-Centric Video Understanding","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.06355","snapshot_observed_at":"2026-08-06T21:28:59.047091Z","title":"Videochat: Chat-centric video understanding","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":52,"source":"pdf_text","source_observed_at":"2026-08-06T21:28:59.047091Z"},"links":{"cited_paper":"/paper/2305.06355","citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:833b39fb74b166cde66203e3c9c7deab98cca72846eab9c997eb06ef25d99316","observation_id":"af2f22b6-21ba-46ab-afb4-c370dc8eafbd","resolution":{"observed_at":"2026-08-06T21:28:59.047091Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:28:59.149454Z","title":"Clip-event: Connecting text and images with event structures","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":53,"source":"pdf_text","source_observed_at":"2026-08-06T21:28:59.149454Z"},"links":{"citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:4cfbb93433f6465d4356314394189157fb252c743a4555b81c154b07ca629a12","observation_id":"aa420b5e-e238-4180-8fc4-61b8cf5e5151","resolution":{"observed_at":"2026-08-06T21:28:59.149454Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:28:59.266500Z","title":"Enhancing the locality and breaking the memory bottleneck of transformer on time series forecasting","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":54,"source":"pdf_text","source_observed_at":"2026-08-06T21:28:59.266500Z"},"links":{"citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:3d8144561accc8d22b7be9699d7c02a2846dba738efb5314c23f088b2860cc3d","observation_id":"6b59d7f7-01f5-4599-8fef-2b37ccd6bad2","resolution":{"observed_at":"2026-08-06T21:28:59.266500Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:28:59.368840Z","title":"Deep learning models for time series forecasting: a review","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":55,"source":"pdf_text","source_observed_at":"2026-08-06T21:28:59.368840Z"},"links":{"citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:88985ff3027af7c00a51e24db8dd6d9dc3712016246c12037ade990c409740ab","observation_id":"b6de7484-019b-4e50-b757-498817c3b545","resolution":{"observed_at":"2026-08-06T21:28:59.368840Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:28:59.457459Z","title":"Forecasting with time series imaging","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":56,"source":"pdf_text","source_observed_at":"2026-08-06T21:28:59.457459Z"},"links":{"citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:2f16fc3ea1f1af68702b1b9efb7aab9eacdc1204d58426ccbd6894d8301f3577","observation_id":"31a0e8ad-e176-495f-b3e0-71089a6034eb","resolution":{"observed_at":"2026-08-06T21:28:59.457459Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:28:59.558862Z","title":"Time series as images: Vision transformer for irregularly sampled time series","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":57,"source":"pdf_text","source_observed_at":"2026-08-06T21:28:59.558862Z"},"links":{"citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:af82b004fb5c70f39574d1570a9419a02ccd5efbbc6b59d4996b400b100c8058","observation_id":"da2b637c-03c0-41c8-ad19-7c9321028db5","resolution":{"observed_at":"2026-08-06T21:28:59.558862Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2305.10721","last_updated":"2026-05-17T02:25:14Z","snapshot_observed_at":"2026-08-03T04:41:00.104637Z","submitted_at":"2023-05-18T05:39:46Z","title":"Revisiting Long-term Time Series Forecasting: An Investigation on Linear Mapping","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2305.10721","snapshot_observed_at":"2026-08-06T21:28:59.666584Z","title":"Revisiting long-term time series forecasting: An investigation on linear mapping","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":58,"source":"pdf_text","source_observed_at":"2026-08-06T21:28:59.666584Z"},"links":{"cited_paper":"/paper/2305.10721","citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:c74b8b6b7bbe20f79811cc1ea4358ca9d13326d30f4fd64279a59b4d9e39a7b5","observation_id":"dd3481ad-3e57-44a6-9dc9-478e0c4fa6f2","resolution":{"observed_at":"2026-08-06T21:28:59.666584Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2308.11200","last_updated":"2023-08-22T05:23:04Z","snapshot_observed_at":"2026-08-09T23:13:06.346823Z","submitted_at":"2023-08-22T05:23:04Z","title":"SegRNN: Segment Recurrent Neural Network for Long-Term Time Series Forecasting","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2308.11200","snapshot_observed_at":"2026-08-06T21:28:59.787950Z","title":"Segrnn: Segment recurrent neural network for long-term time series forecasting.arXiv preprint arXiv:2308.11200, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":59,"source":"pdf_text","source_observed_at":"2026-08-06T21:28:59.787950Z"},"links":{"cited_paper":"/paper/2308.11200","citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:3b3df2466298f1e6bfacb17aab6fb4f606986033097492f001d47dd8b571a9e7","observation_id":"71d8c771-4547-453e-bd3c-b0c49dc6d522","resolution":{"observed_at":"2026-08-06T21:28:59.787950Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:28:59.941861Z","title":"Pth and the regulation of mesenchymal cells within the bone marrow niche","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":60,"source":"pdf_text","source_observed_at":"2026-08-06T21:28:59.941861Z"},"links":{"citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:5744f54428002038d06ecde93113364261f394110314f1d5710e408d6690fa2c","observation_id":"df36f1cf-1232-4907-9b6b-d5a3efe2a2e0","resolution":{"observed_at":"2026-08-06T21:28:59.941861Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:29:00.082909Z","title":"Visual instruction tuning","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":61,"source":"pdf_text","source_observed_at":"2026-08-06T21:29:00.082909Z"},"links":{"citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:1b3f6d8399102fb06bc730c611bff78fe5a6bfe70a4b68fca396f1b2ac74b69a","observation_id":"47fb0c6e-05a6-4874-82ff-1185c8628f69","resolution":{"observed_at":"2026-08-06T21:29:00.082909Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.16132","last_updated":"2024-02-25T16:14:26Z","snapshot_observed_at":"2026-08-07T22:38:57.643638Z","submitted_at":"2024-02-25T16:14:26Z","title":"LSTPrompt: Large Language Models as Zero-Shot Time Series Forecasters by Long-Short-Term Prompting","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.16132","snapshot_observed_at":"2026-08-06T21:29:00.230591Z","title":"Lstprompt: Large language models as zero-shot time series forecasters by long-short-term prompting","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":62,"source":"pdf_text","source_observed_at":"2026-08-06T21:29:00.230591Z"},"links":{"cited_paper":"/paper/2402.16132","citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:564e2ccbd4e2c963e3e9c9b2de6b2a624c4204d1b18c7bd6a93b1a660f7234a0","observation_id":"d66ff5d9-1eb6-493b-9437-89ff3428ee81","resolution":{"observed_at":"2026-08-06T21:29:00.230591Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:29:00.409086Z","title":"Edta enhances stromal cell–derived factor 1α–induced migration of dental pulp cells by up-regulating chemokine receptor 4 expression","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":63,"source":"pdf_text","source_observed_at":"2026-08-06T21:29:00.409086Z"},"links":{"citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:8b11cc6a553f6d1ceff3e6b995543ac8dc871a3d2b19d786317b4fbfa7c6c0ff","observation_id":"4a0ec73e-067e-414a-b936-e089c56b433a","resolution":{"observed_at":"2026-08-06T21:29:00.409086Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:29:00.503390Z","title":"Calorie restriction in mice impairs cortical but not trabecular peak bone mass by suppressing bone remodeling","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":64,"source":"pdf_text","source_observed_at":"2026-08-06T21:29:00.503390Z"},"links":{"citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:ab8a43dc8b141c7c41090863c3ad33406a782c50cdca44b462b91dcc562cefb4","observation_id":"f079f970-d6c2-4b8b-81fa-36e8c19e0858","resolution":{"observed_at":"2026-08-06T21:29:00.503390Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:29:00.626740Z","title":"Scinet: Time series modeling and forecasting with sample convolution and interaction","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":65,"source":"pdf_text","source_observed_at":"2026-08-06T21:29:00.626740Z"},"links":{"citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:fc8974186519cb9dee84cea8ea97c2741993d977ece79dc34557e377c1dc11c5","observation_id":"a5eabf00-0e70-4447-9965-89d2782fd179","resolution":{"observed_at":"2026-08-06T21:29:00.626740Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:29:00.728346Z","title":"Focal: Contrastive learning for multimodal time- series sensing signals in factorized orthogonal latent space","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":66,"source":"pdf_text","source_observed_at":"2026-08-06T21:29:00.728346Z"},"links":{"citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:342bddf3c19d3e0e91ee1b3e24c925b78ed57c9538b4bb81f86877c74a0541be","observation_id":"9a76fdb9-5d29-4675-9ce3-80e629141c7b","resolution":{"observed_at":"2026-08-06T21:29:00.728346Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:29:00.811287Z","title":"Pyraformer: Low-complexity pyramidal attention for long-range time series modeling and forecasting","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":67,"source":"pdf_text","source_observed_at":"2026-08-06T21:29:00.811287Z"},"links":{"citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:98684bee7968400a16f2ddb0f8166dcbf35d54ea3bdc123775505695910ec901","observation_id":"d2e4bba7-044d-42cc-acda-4ee820ac12ae","resolution":{"observed_at":"2026-08-06T21:29:00.811287Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:29:00.904259Z","title":"Unitime: A language-empowered unified model for cross-domain time series forecasting","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":68,"source":"pdf_text","source_observed_at":"2026-08-06T21:29:00.904259Z"},"links":{"citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:6ed016545c47727d9c4925035f249cb509770e503630b7f0db746c9aeff01b27","observation_id":"4b2c9914-8523-4ba1-8411-c6237ef1f62e","resolution":{"observed_at":"2026-08-06T21:29:00.904259Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:29:01.005278Z","title":"Non-stationary transformers: Ex- ploring the stationarity in time series forecasting","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":69,"source":"pdf_text","source_observed_at":"2026-08-06T21:29:01.005278Z"},"links":{"citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:932a6ad09414865b5b1c1d4a25843dddaf4f1d0ba3741907646313b1aaf33d56","observation_id":"4c9ab9be-20b4-46a3-8beb-cd07e9c6b4b0","resolution":{"observed_at":"2026-08-06T21:29:01.005278Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2310.06625","last_updated":"2024-03-14T11:45:57Z","snapshot_observed_at":"2026-07-06T16:30:29.783501Z","submitted_at":"2023-10-10T13:44:09Z","title":"iTransformer: Inverted Transformers Are Effective for Time Series Forecasting","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2310.06625","snapshot_observed_at":"2026-08-06T21:29:01.134359Z","title":"itransformer: Inverted transformers are effective for time series forecasting, 2024","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":70,"source":"pdf_text","source_observed_at":"2026-08-06T21:29:01.134359Z"},"links":{"cited_paper":"/paper/2310.06625","citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:be7af23d545ddd139b92946b39e317b5a096f67e7d9b60df014ff16156974dac","observation_id":"be74caf6-ba53-476d-8b93-c2718966e95a","resolution":{"observed_at":"2026-08-06T21:29:01.134359Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2402.02370","last_updated":"2024-10-31T11:37:41Z","snapshot_observed_at":"2026-08-09T22:53:58.003030Z","submitted_at":"2024-02-04T06:59:21Z","title":"AutoTimes: Autoregressive Time Series Forecasters via Large Language Models","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2402.02370","snapshot_observed_at":"2026-08-06T21:29:01.280307Z","title":"Auto- times: Autoregressive time series forecasters via large language models","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":71,"source":"pdf_text","source_observed_at":"2026-08-06T21:29:01.280307Z"},"links":{"cited_paper":"/paper/2402.02370","citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:bc2c7404b401f94781817926b933606c4c50de54b07e609297c9c52f0454c762","observation_id":"aacc89a5-f6e3-4e5e-adac-90cc3a3a59ec","resolution":{"observed_at":"2026-08-06T21:29:01.280307Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:29:01.404186Z","title":"Timer: Generative pre-trained transformers are large time series models","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":72,"source":"pdf_text","source_observed_at":"2026-08-06T21:29:01.404186Z"},"links":{"citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:9eaafc31bc4f191b2bb69a679b8fcea8a83c2eeebcc09ffaece79299f54196d5","observation_id":"78bc09ea-079c-4151-883d-5eca9fe8ee5a","resolution":{"observed_at":"2026-08-06T21:29:01.404186Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:29:01.514818Z","title":"Swin transformer: Hierarchical vision transformer using shifted windows","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":73,"source":"pdf_text","source_observed_at":"2026-08-06T21:29:01.514818Z"},"links":{"citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:410a0873df0969709ea660c8a6fabdce79c436e97b12fe7d3c58b25775524c66","observation_id":"6315744b-487e-448f-8e59-6245bfa6ae04","resolution":{"observed_at":"2026-08-06T21:29:01.514818Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:29:01.631306Z","title":"A cnn-bilstm-am method for stock price prediction","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":74,"source":"pdf_text","source_observed_at":"2026-08-06T21:29:01.631306Z"},"links":{"citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:d8eafdb0178affd4b5394f2be85abe25210300bfa1d7d274f3ec6ab2638e4ef7","observation_id":"b23abbb3-d2c6-4d74-ac0e-27b48ca99517","resolution":{"observed_at":"2026-08-06T21:29:01.631306Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:29:01.715495Z","title":"Howto100m: Learning a text-video embedding by watching hundred million narrated video clips","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":75,"source":"pdf_text","source_observed_at":"2026-08-06T21:29:01.715495Z"},"links":{"citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:3236552c58ddf94bb3920955493a26675393a2a9d9fb264d3b1389d2d8bc1581","observation_id":"f575aaf1-cb54-4661-a09f-f84009812730","resolution":{"observed_at":"2026-08-06T21:29:01.715495Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:29:01.824833Z","title":"Expanding language-image pretrained models for general video recognition","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":76,"source":"pdf_text","source_observed_at":"2026-08-06T21:29:01.824833Z"},"links":{"citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:2519d46ad1aec6ccde1914c18118e4f599d185e3c418b7c0929b09ff15ff97eb","observation_id":"5e4334fc-5e37-4b23-87fa-3f56cd0bbb56","resolution":{"observed_at":"2026-08-06T21:29:01.824833Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2211.14730","last_updated":"2023-03-05T22:11:56Z","snapshot_observed_at":"2026-07-31T22:45:35.561492Z","submitted_at":"2022-11-27T05:15:42Z","title":"A Time Series is Worth 64 Words: Long-term Forecasting with Transformers","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2211.14730","snapshot_observed_at":"2026-08-06T21:29:01.973202Z","title":"A time series is worth 64 words: Long-term forecasting with transformers","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":77,"source":"pdf_text","source_observed_at":"2026-08-06T21:29:01.973202Z"},"links":{"cited_paper":"/paper/2211.14730","citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:46eebb7a579993d34b155c2bac250dafd2e32ba96b102230467b6abe0ab4e230","observation_id":"a5c6dda5-0abd-44f3-a870-11a65f1f7a81","resolution":{"observed_at":"2026-08-06T21:29:01.973202Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1807.03748","last_updated":"2019-01-22T18:47:12Z","snapshot_observed_at":"2026-07-06T06:49:24.960992Z","submitted_at":"2018-07-10T16:52:11Z","title":"Representation Learning with Contrastive Predictive Coding","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1807.03748","snapshot_observed_at":"2026-08-06T21:29:02.114667Z","title":"Representation learning with contrastive predictive coding","venue":null,"work_id":null,"year":2018},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":78,"source":"pdf_text","source_observed_at":"2026-08-06T21:29:02.114667Z"},"links":{"cited_paper":"/paper/1807.03748","citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:029b179ab8f3affb29eb3f6a14d20060a9f6f510c23f4e9b15126498205a96cc","observation_id":"e5c409ef-0a7d-4e52-9790-1bf5715c70ae","resolution":{"observed_at":"2026-08-06T21:29:02.114667Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1905.10437","last_updated":"2020-02-20T21:08:57Z","snapshot_observed_at":"2026-08-06T09:33:52.442294Z","submitted_at":"2019-05-24T20:28:57Z","title":"N-BEATS: Neural basis expansion analysis for interpretable time series forecasting","version":4},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1905.10437","snapshot_observed_at":"2026-08-06T21:29:02.218628Z","title":"N-beats: Neural basis expansion analysis for interpretable time series forecasting","venue":null,"work_id":null,"year":1905},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":79,"source":"pdf_text","source_observed_at":"2026-08-06T21:29:02.218628Z"},"links":{"cited_paper":"/paper/1905.10437","citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:84555aebf25b015976203197bb6db032db887f7ec0c2a491a8762c5b2d576594","observation_id":"66d8dbab-91c0-4d5c-af8e-9b3c745ff5cd","resolution":{"observed_at":"2026-08-06T21:29:02.218628Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:29:02.308365Z","title":"Pytorch: An imperative style, high-performance deep learning library","venue":null,"work_id":null,"year":2019},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":80,"source":"pdf_text","source_observed_at":"2026-08-06T21:29:02.308365Z"},"links":{"citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:c10336230cfe6a22249d30c03e4cc258d3b87b86aec1bbee64403aba008f3ad6","observation_id":"2ee961b2-2bc9-4391-af25-348f23a172a0","resolution":{"observed_at":"2026-08-06T21:29:02.308365Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:29:02.411586Z","title":"Learning transferable visual models from natural language supervision","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":81,"source":"pdf_text","source_observed_at":"2026-08-06T21:29:02.411586Z"},"links":{"citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:d179ae87476c89cfbd399bab8f235b12f94274d11def554772b6c330f5a2685a","observation_id":"a9d6c929-8b36-4fb5-b390-bda00216289a","resolution":{"observed_at":"2026-08-06T21:29:02.411586Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:29:02.586917Z","title":"Exploring the limits of transfer learning with a unified text-to-text transformer","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":82,"source":"pdf_text","source_observed_at":"2026-08-06T21:29:02.586917Z"},"links":{"citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:4e013257fa45bae9be40e77de922b340e6594564e47ca14811cd8381df1b704e","observation_id":"f67a7772-c1b6-4031-ba0c-3ced2c7c61a6","resolution":{"observed_at":"2026-08-06T21:29:02.586917Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:29:02.729299Z","title":"Automatic diagnosis of the 12-lead ecg using a deep neural network","venue":null,"work_id":null,"year":2020},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":83,"source":"pdf_text","source_observed_at":"2026-08-06T21:29:02.729299Z"},"links":{"citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:6b4ad7b8f2cceb9add955072147906a3f9c90554a7c5feea46f13ec9706bf989","observation_id":"f8810c08-9755-4108-a7ad-296623248da6","resolution":{"observed_at":"2026-08-06T21:29:02.729299Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:29:02.918008Z","title":"High-resolution image synthesis with latent diffusion models","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":84,"source":"pdf_text","source_observed_at":"2026-08-06T21:29:02.918008Z"},"links":{"citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:c8297452c1714041dfefb433304824f1102d5b13539f4b14929a0cd25fc127d5","observation_id":"bc625d19-e4bd-4c4d-8895-a6d879bf533a","resolution":{"observed_at":"2026-08-06T21:29:02.918008Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:29:03.084368Z","title":"A review of deep learning techniques for forecasting energy use in buildings","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":85,"source":"pdf_text","source_observed_at":"2026-08-06T21:29:03.084368Z"},"links":{"citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:74444d025eef8ea98939ce0bb809c64718cbee92b4905277b456fa1ac658e3fa","observation_id":"1a3a896a-5155-47dd-9c0e-f3188e70785a","resolution":{"observed_at":"2026-08-06T21:29:03.084368Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:29:03.220999Z","title":"Image- based time series forecasting: A deep convolutional neural network approach.Neural Networks, 157:39–53, 2023","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":86,"source":"pdf_text","source_observed_at":"2026-08-06T21:29:03.220999Z"},"links":{"citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:bc6b37cc2228022b950e681fec6642f3d04c95acce62f5ede7da35816c9c5ff1","observation_id":"b31f76a3-7a37-41bb-85ed-3265589db840","resolution":{"observed_at":"2026-08-06T21:29:03.220999Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:29:03.332985Z","title":"Dust: Dual swin transformer for multi- modal video and time-series modeling","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":87,"source":"pdf_text","source_observed_at":"2026-08-06T21:29:03.332985Z"},"links":{"citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:2c13d8b77b2d286218b8ca14d0c33b907704dcbe9e871a40a0061a3e1993ebc7","observation_id":"37aee5d3-dbca-499e-a0be-57664450a828","resolution":{"observed_at":"2026-08-06T21:29:03.332985Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"1906.05743","last_updated":"2019-09-27T21:59:59Z","snapshot_observed_at":"2026-08-07T11:03:21.673722Z","submitted_at":"2019-06-13T15:03:52Z","title":"Learning Video Representations using Contrastive Bidirectional Transformer","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"1906.05743","snapshot_observed_at":"2026-08-06T21:29:03.461180Z","title":"Learning video representa- tions using contrastive bidirectional transformer","venue":null,"work_id":null,"year":1906},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":88,"source":"pdf_text","source_observed_at":"2026-08-06T21:29:03.461180Z"},"links":{"cited_paper":"/paper/1906.05743","citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:e3dc944f3de5d6903839fac8fdf2481a5894ebc5d18696080db1026d517add25","observation_id":"2e281257-1145-46ff-a2bd-75c1dd0387fb","resolution":{"observed_at":"2026-08-06T21:29:03.461180Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:29:19.545978Z","title":"Videobert: A joint model for video and language representation learning","venue":null,"work_id":"59497dca-f2e0-4536-a633-c882519c5345","year":2019},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":89,"source":"pdf_text","source_observed_at":"2026-08-06T21:29:03.627353Z"},"links":{"citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:d2ae4b365736bf89134ee107fb973b417746077a2209106f45648879033b4c78","observation_id":"2dc11e7f-c895-4848-b551-e8f1cec888c3","resolution":{"observed_at":"2026-08-06T21:29:19.700439Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2308.08241","last_updated":"2024-02-22T02:03:42Z","snapshot_observed_at":"2026-08-08T09:15:16.308286Z","submitted_at":"2023-08-16T09:16:02Z","title":"TEST: Text Prototype Aligned Embedding to Activate LLM's Ability for Time Series","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2308.08241","snapshot_observed_at":"2026-08-06T21:29:03.796086Z","title":"Test: Text prototype aligned embedding to activate llm’s ability for time series","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":90,"source":"pdf_text","source_observed_at":"2026-08-06T21:29:03.796086Z"},"links":{"cited_paper":"/paper/2308.08241","citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:ceb2ef4690aaf5e31e0eb6beede323fc807bf64aab0d7308fceb5ff7d72db209","observation_id":"abda9d6e-5505-4bd9-8d7f-2c4a62896087","resolution":{"observed_at":"2026-08-06T21:29:03.796086Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2210.06031","last_updated":"2023-03-02T09:05:43Z","snapshot_observed_at":"2026-08-01T22:41:21.854568Z","submitted_at":"2022-10-12T09:08:27Z","title":"Long-Form Video-Language Pre-Training with Multimodal Temporal Contrastive Learning","version":2},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2210.06031","snapshot_observed_at":"2026-08-06T21:29:03.936579Z","title":"Long-form video-language pre-training with multimodal temporal contrastive learning","venue":null,"work_id":null,"year":2022},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":91,"source":"pdf_text","source_observed_at":"2026-08-06T21:29:03.936579Z"},"links":{"cited_paper":"/paper/2210.06031","citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:2262dd116ee7e595ba773f83d58d1458e0c0b8e1c5f846c6170cd9ae3a9f8265","observation_id":"5c8101c3-3abd-4b70-8500-7ed013e5dc5e","resolution":{"observed_at":"2026-08-06T21:29:03.936579Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:29:19.309376Z","title":"Are language models actually useful for time series forecasting? In The Thirty-eighth Annual Conference on Neural Information Processing Systems, 2024","venue":null,"work_id":"6d17bb05-bccc-4bae-9be7-572302435a89","year":2024},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":92,"source":"pdf_text","source_observed_at":"2026-08-06T21:29:04.042255Z"},"links":{"citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:173d2b10655c3e399bad6efa39809add0c84502e9b9561506d634b394fb9fc8b","observation_id":"7b2ee3d7-90fb-4978-ac46-bbd58fe0f971","resolution":{"observed_at":"2026-08-06T21:29:19.417664Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:29:19.112767Z","title":"Are language models actually useful for time series forecasting? In The Thirty-eighth Annual Conference on Neural Information Processing Systems, 2024","venue":null,"work_id":"57479207-d8fd-4ec5-a42d-a9a4a446a3b5","year":2024},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":93,"source":"pdf_text","source_observed_at":"2026-08-06T21:29:04.148304Z"},"links":{"citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:496224e13da1b871e40b6d36b8145b757ce202bffc2fb86817c0b6d03341e842","observation_id":"a3147d8f-832b-427e-b3f3-8afbb9bde3d3","resolution":{"observed_at":"2026-08-06T21:29:19.196750Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2302.13971","last_updated":"2023-02-27T17:11:15Z","snapshot_observed_at":"2026-07-06T02:11:23.670680Z","submitted_at":"2023-02-27T17:11:15Z","title":"LLaMA: Open and Efficient Foundation Language Models","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2302.13971","snapshot_observed_at":"2026-08-06T21:29:04.295397Z","title":"Llama: Open and efficient foundation language models","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":94,"source":"pdf_text","source_observed_at":"2026-08-06T21:29:04.295397Z"},"links":{"cited_paper":"/paper/2302.13971","citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:d9afa309e5deba168562fa9fecff15ae5639d996f73b8e3876a9174f00fc0b78","observation_id":"e831c829-605f-4219-991a-3d502d35445b","resolution":{"observed_at":"2026-08-06T21:29:04.295397Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:29:04.435893Z","title":"Attention is all you need","venue":null,"work_id":null,"year":2017},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":95,"source":"pdf_text","source_observed_at":"2026-08-06T21:29:04.435893Z"},"links":{"citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:aefeb5a2c5190992a683e287b943d65106c66dbe2dc8e3fe8fab2a57f16e0f5b","observation_id":"167d606e-03eb-457b-952c-9774ae0356d7","resolution":{"observed_at":"2026-08-06T21:29:04.435893Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2501.10555","last_updated":"2025-01-17T21:05:09Z","snapshot_observed_at":"2026-07-06T20:22:41.032443Z","submitted_at":"2025-01-17T21:05:09Z","title":"Towards Data-Centric AI: A Comprehensive Survey of Traditional, Reinforcement, and Generative Approaches for Tabular Data Transformation","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2501.10555","snapshot_observed_at":"2026-08-06T21:29:04.583770Z","title":"Towards data-centric ai: A com- prehensive survey of traditional, reinforcement, and generative approaches for tabular data transformation","venue":null,"work_id":null,"year":2025},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":96,"source":"pdf_text","source_observed_at":"2026-08-06T21:29:04.583770Z"},"links":{"cited_paper":"/paper/2501.10555","citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:682cb1bd21c6612c2e60a29f39a6d87b166d533d297f19fcd1d2aaadb53d8214","observation_id":"a2726999-8d48-4c29-a8c9-2386cdd8ecf8","resolution":{"observed_at":"2026-08-06T21:29:04.583770Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:29:04.698553Z","title":"Micn: Multi-scale local and global context modeling for long-term series forecasting","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":97,"source":"pdf_text","source_observed_at":"2026-08-06T21:29:04.698553Z"},"links":{"citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:aaffbef1ec3c00029e41551d7bd377353440a241359cc3bfbe7d3063ea9e4fa1","observation_id":"66a86caf-e597-4e13-b649-582afd9a3f95","resolution":{"observed_at":"2026-08-06T21:29:04.698553Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:29:18.873157Z","title":"Long-short temporal contrastive learning of video transformers","venue":null,"work_id":"df325e12-d48f-42df-a826-7d653d527ae4","year":2022},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":98,"source":"pdf_text","source_observed_at":"2026-08-06T21:29:04.825726Z"},"links":{"citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:ae6e9e4b8e80abc08b2f6b43c9c18c06db8085a0c3a92b9f962871d7ecf04ec7","observation_id":"8e26c3b6-34a2-465d-8a82-efc4eac27530","resolution":{"observed_at":"2026-08-06T21:29:18.967420Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2109.08472","last_updated":"2021-09-17T11:21:34Z","snapshot_observed_at":"2026-07-06T11:48:42.000083Z","submitted_at":"2021-09-17T11:21:34Z","title":"ActionCLIP: A New Paradigm for Video Action Recognition","version":1},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2109.08472","snapshot_observed_at":"2026-08-06T21:29:05.004317Z","title":"Actionclip: A new paradigm for video action recognition","venue":null,"work_id":null,"year":2021},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":99,"source":"pdf_text","source_observed_at":"2026-08-06T21:29:05.004317Z"},"links":{"cited_paper":"/paper/2109.08472","citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:1c3bea1ca9a734fa3b1b5ae8d31c701ae74d7b8101e0af1fb7fb6b5ad58b0e8c","observation_id":"5e0caaa7-7186-43ea-a47c-e988f75f6f33","resolution":{"observed_at":"2026-08-06T21:29:05.004317Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":null,"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":"raw_reference","pith_arxiv_id":null,"snapshot_observed_at":"2026-08-06T21:29:18.598391Z","title":"A hierarchal bert structure for native speaker writing detection","venue":null,"work_id":"24bc2214-5057-4c49-92e4-2a07cab959ad","year":2022},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":100,"source":"pdf_text","source_observed_at":"2026-08-06T21:29:05.163304Z"},"links":{"citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:eb7ab4688d31fc703708def9d106448ccc8a7a3eae5d5c4aee5efc56a60777cc","observation_id":"2d9357f1-380d-45bb-b419-1bd210e8b1f3","resolution":{"observed_at":"2026-08-06T21:29:18.704941Z","resolver_source":"raw_fallback","status":"verified_fuzzy"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2410.03521","last_updated":"2025-02-25T02:17:05Z","snapshot_observed_at":"2026-07-06T19:27:47.617823Z","submitted_at":"2024-09-27T00:01:32Z","title":"Building a Chinese Medical Dialogue System: Integrating Large-scale Corpora and Novel Models","version":2},"cited_work":{"arxiv_id":"2410.03521","doi":null,"metadata_source":"pith","pith_arxiv_id":"2410.03521","snapshot_observed_at":"2026-08-06T21:29:12.003849Z","title":"Building a Chinese Medical Dialogue System: Integrating Large-scale Corpora and Novel Models","venue":"cs.CL","work_id":"45b13f6d-32d4-4d79-b334-5108c77aa71a","year":2024},"citing_paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives","version":2},"reference_index":101,"source":"pdf_text","source_observed_at":"2026-08-06T21:29:05.314870Z"},"links":{"cited_paper":"/paper/2410.03521","citing_paper":"/paper/2506.24124"},"observation_digest":"sha256:4b3ff3f98668c60e0b408d5ff4b0706db44908b36a7e736d38141e3d8fddec8e","observation_id":"2ec5269b-6731-4f8a-b4c0-58c71406f8cb","resolution":{"observed_at":"2026-08-06T21:29:12.093386Z","resolver_source":"local_arxiv","status":"verified_exact"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-09T06:31:02.800959+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"state":"measured"}}],"paper":{"arxiv_id":"2506.24124","last_updated":"2025-07-01T03:40:22Z","latest_version":2,"primary_category":"cs.LG","snapshot_observed_at":"2026-08-08T13:48:41.128139Z","submitted_at":"2025-06-30T17:59:14Z","title":"Teaching Time Series to See and Speak: Forecasting with Aligned Visual and Textual Perspectives"},"reference_resolution":{"displayed":100,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":90,"verified_exact":5,"verified_fuzzy":5},"total_outbound_references":138},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-09T06:31:02.800959+00:00","source":"crossref"},{"observed_at":"2026-08-09T06:30:57.326959+00:00","source":"retraction_watch"}],"thesis":"As of 10 August 2026, this Paper Citation Record lists 100 of 138 outbound references and 3 inbound Pith citation observations for arXiv:2506.24124."}