{"as_of":"2026-08-10T13:21:00Z","caps":{"database_statements":6,"inbound":100,"outbound":100},"context_digest":"sha256:1c7cb21d4b09a5c5c62910accbc51aa92fb2665cf05e66e629fd733918e0a944","coverage":[{"denominator":0,"lane":"reference_resolution","note":"Typed states for the displayed outbound observations.","records_observed":0,"source":"paper_references, paper_reference_links","source_observed_at":null,"state":"measured"},{"denominator":20,"lane":"standing_notices","note":"One-hop event checks from named stored sources.","records_observed":20,"source":"scholarly_work_events, retraction_status_cache","source_observed_at":"2026-08-10T06:31:04.303077+00:00","state":"measured"},{"denominator":20,"lane":"inbound_itemization","note":"Pith citing papers itemized under the disclosed page cap.","records_observed":20,"source":"paper_references, paper_reference_links","source_observed_at":"2026-08-10T00:36:46.453123Z","state":"measured"},{"denominator":1,"lane":"external_citation_measurements","note":"A source-named dated measurement, never combined with another source.","records_observed":0,"source":"arxiv_reference","source_observed_at":"2026-07-02T16:07:09.407610Z","state":"measured"}],"external_citation_measurements":[],"inbound":[{"citation":{"cited_paper":{"arxiv_id":"2312.15166","last_updated":"2024-04-04T01:53:38Z","snapshot_observed_at":"2026-08-09T10:33:25.832634Z","submitted_at":"2023-12-23T05:11:37Z","title":"SOLAR 10.7B: Scaling Large Language Models with Simple yet Effective Depth Up-Scaling","version":3},"cited_work":{"arxiv_id":"2312.15166","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2312.15166","snapshot_observed_at":"2026-07-02T16:07:09.407610Z","title":"Solar 10.7 b: Scaling large language models with simple yet effective depth up-scaling.arXiv preprint arXiv:2312.15166","venue":null,"work_id":"12c284dc-a58b-4942-8c65-c730a8a04685","year":2023},"citing_paper":{"arxiv_id":"2404.10981","last_updated":"2024-08-23T00:17:02Z","snapshot_observed_at":"2026-07-06T18:01:18.745347Z","submitted_at":"2024-04-17T01:27:42Z","title":"A Survey on Retrieval-Augmented Text Generation for Large Language Models","version":2},"reference_index":77,"source":"pdf_text","source_observed_at":"2026-05-24T02:15:05.379583Z"},"links":{"cited_paper":"/paper/2312.15166","citing_paper":"/paper/2404.10981"},"observation_digest":"sha256:c783cb426a8897684e308c4accbfe9488bce2f3e6c859c72e92ab2eeb37505b6","observation_id":"0eab44cc-62e6-49aa-903d-1710b18c0d08","resolution":{"observed_at":"2026-05-24T02:15:55.447265Z","resolver_source":"arxiv_id","status":"malformed_identifier"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2312.15166","last_updated":"2024-04-04T01:53:38Z","snapshot_observed_at":"2026-08-09T10:33:25.832634Z","submitted_at":"2023-12-23T05:11:37Z","title":"SOLAR 10.7B: Scaling Large Language Models with Simple yet Effective Depth Up-Scaling","version":3},"cited_work":{"arxiv_id":"2312.15166","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2312.15166","snapshot_observed_at":"2026-07-02T16:07:09.407610Z","title":"Solar 10.7 b: Scaling large language models with simple yet effective depth up-scaling.arXiv preprint arXiv:2312.15166","venue":null,"work_id":"12c284dc-a58b-4942-8c65-c730a8a04685","year":2023},"citing_paper":{"arxiv_id":"2406.00515","last_updated":"2024-11-10T22:02:27Z","snapshot_observed_at":"2026-07-29T20:40:25.374189Z","submitted_at":"2024-06-01T17:48:15Z","title":"A Survey on Large Language Models for Code Generation","version":2},"reference_index":133,"source":"pdf_text","source_observed_at":"2026-05-13T20:18:06.304134Z"},"links":{"cited_paper":"/paper/2312.15166","citing_paper":"/paper/2406.00515"},"observation_digest":"sha256:e2e0e968b7693f73d74a58ed284aac00a739262d5f6287edf2276547506658fd","observation_id":"59d9b44f-fe0b-4a66-b290-78c6dac7b04c","resolution":{"observed_at":"2026-05-13T20:18:06.706479Z","resolver_source":"arxiv_id","status":"malformed_identifier"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2312.15166","last_updated":"2024-04-04T01:53:38Z","snapshot_observed_at":"2026-08-09T10:33:25.832634Z","submitted_at":"2023-12-23T05:11:37Z","title":"SOLAR 10.7B: Scaling Large Language Models with Simple yet Effective Depth Up-Scaling","version":3},"cited_work":{"arxiv_id":"2312.15166","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2312.15166","snapshot_observed_at":"2026-07-02T16:07:09.407610Z","title":"Solar 10.7 b: Scaling large language models with simple yet effective depth up-scaling.arXiv preprint arXiv:2312.15166","venue":null,"work_id":"12c284dc-a58b-4942-8c65-c730a8a04685","year":2023},"citing_paper":{"arxiv_id":"2408.07666","last_updated":"2025-12-31T04:06:49Z","snapshot_observed_at":"2026-08-07T23:28:24.025478Z","submitted_at":"2024-08-14T16:58:48Z","title":"Model Merging in LLMs, MLLMs, and Beyond: Methods, Theories, Applications and Opportunities","version":5},"reference_index":109,"source":"pdf_text","source_observed_at":"2026-05-17T22:16:04.386706Z"},"links":{"cited_paper":"/paper/2312.15166","citing_paper":"/paper/2408.07666"},"observation_digest":"sha256:bb75817bf59917c929cc1fe009714e07f4bd8cf8f20d5b41d83134aafe8e6e11","observation_id":"e359f20a-47c1-48ac-b6cd-0e585d3e4ec6","resolution":{"observed_at":"2026-05-17T22:16:04.819850Z","resolver_source":"arxiv_id","status":"malformed_identifier"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2312.15166","last_updated":"2024-04-04T01:53:38Z","snapshot_observed_at":"2026-08-09T10:33:25.832634Z","submitted_at":"2023-12-23T05:11:37Z","title":"SOLAR 10.7B: Scaling Large Language Models with Simple yet Effective Depth Up-Scaling","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.15166","snapshot_observed_at":"2026-08-10T00:36:46.453123Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2501.18128","last_updated":"2025-01-30T04:20:16Z","snapshot_observed_at":"2026-08-10T00:31:43.505353Z","submitted_at":"2025-01-30T04:20:16Z","title":"Unraveling the Capabilities of Language Models in News Summarization","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-10T00:36:46.453123Z"},"links":{"cited_paper":"/paper/2312.15166","citing_paper":"/paper/2501.18128"},"observation_digest":"sha256:0d200557ca0b8ff66807793799e84ac6d2ebba44a7185de6299445161eba3ae8","observation_id":"13229869-7332-4292-b50c-16471c4779cf","resolution":{"observed_at":"2026-08-10T00:36:46.453123Z","resolver_source":null,"status":"unresolved"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2312.15166","last_updated":"2024-04-04T01:53:38Z","snapshot_observed_at":"2026-08-09T10:33:25.832634Z","submitted_at":"2023-12-23T05:11:37Z","title":"SOLAR 10.7B: Scaling Large Language Models with Simple yet Effective Depth Up-Scaling","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.15166","snapshot_observed_at":"2026-08-08T23:51:29.380623Z","title":"Solar 10.7 b: Scaling large language models with simple yet effective depth up-scaling","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2502.04030","last_updated":"2025-06-25T14:44:30Z","snapshot_observed_at":"2026-08-09T18:05:30.973187Z","submitted_at":"2025-02-06T12:47:25Z","title":"Fine, I'll Merge It Myself: A Multi-Fidelity Framework for Automated Model Merging","version":2},"reference_index":22,"source":"arxiv_source","source_observed_at":"2026-08-08T23:51:29.380623Z"},"links":{"cited_paper":"/paper/2312.15166","citing_paper":"/paper/2502.04030"},"observation_digest":"sha256:36eee14241201e3533d48cc61f15342464b52e4fa93b6edc2ad3cc8f9924d17d","observation_id":"d6b26ee5-8b09-48b2-b473-878efa0036a0","resolution":{"observed_at":"2026-08-08T23:51:29.380623Z","resolver_source":null,"status":"malformed_identifier"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2312.15166","last_updated":"2024-04-04T01:53:38Z","snapshot_observed_at":"2026-08-09T10:33:25.832634Z","submitted_at":"2023-12-23T05:11:37Z","title":"SOLAR 10.7B: Scaling Large Language Models with Simple yet Effective Depth Up-Scaling","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.15166","snapshot_observed_at":"2026-08-08T11:11:44.600880Z","title":"Solar 10.7 b: Scaling large language models with simple yet effective depth up-scaling","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2502.08020","last_updated":"2025-03-19T16:26:10Z","snapshot_observed_at":"2026-08-09T10:33:59.337803Z","submitted_at":"2025-02-11T23:40:53Z","title":"Speculate, then Collaborate: Fusing Knowledge of Language Models during Decoding","version":2},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-08T11:11:44.600880Z"},"links":{"cited_paper":"/paper/2312.15166","citing_paper":"/paper/2502.08020"},"observation_digest":"sha256:b14a789a4f5b2e42e1cfa40616f89284b7c3d107d9960f35632cc7166681c1b4","observation_id":"f5d946a5-ce6d-4e39-ad76-8bec83e82c11","resolution":{"observed_at":"2026-08-08T11:11:44.600880Z","resolver_source":null,"status":"malformed_identifier"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2312.15166","last_updated":"2024-04-04T01:53:38Z","snapshot_observed_at":"2026-08-09T10:33:25.832634Z","submitted_at":"2023-12-23T05:11:37Z","title":"SOLAR 10.7B: Scaling Large Language Models with Simple yet Effective Depth Up-Scaling","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.15166","snapshot_observed_at":"2026-08-07T15:42:05.936200Z","title":"Solar 10.7 b: Scaling large language models with simple yet effective depth up-scaling.arXiv preprint arXiv:2312.15166, 2023","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2505.14513","last_updated":"2025-05-20T15:41:05Z","snapshot_observed_at":"2026-08-10T04:40:39.370460Z","submitted_at":"2025-05-20T15:41:05Z","title":"Latent Flow Transformer","version":1},"reference_index":15,"source":"pdf_text","source_observed_at":"2026-08-07T15:42:05.936200Z"},"links":{"cited_paper":"/paper/2312.15166","citing_paper":"/paper/2505.14513"},"observation_digest":"sha256:a1c8a48e680e2ee5684ddba5f2b52120187903817bf087eb74d22dedd2ea3c2d","observation_id":"db79b3e1-1de4-4d1b-8839-711344341c75","resolution":{"observed_at":"2026-08-07T15:42:05.936200Z","resolver_source":null,"status":"malformed_identifier"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2312.15166","last_updated":"2024-04-04T01:53:38Z","snapshot_observed_at":"2026-08-09T10:33:25.832634Z","submitted_at":"2023-12-23T05:11:37Z","title":"SOLAR 10.7B: Scaling Large Language Models with Simple yet Effective Depth Up-Scaling","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.15166","snapshot_observed_at":"2026-08-07T14:04:48.602066Z","title":null,"venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2505.20237","last_updated":"2025-06-02T12:59:54Z","snapshot_observed_at":"2026-08-09T05:21:38.284414Z","submitted_at":"2025-05-26T17:17:08Z","title":"Efficient Speech Translation through Model Compression and Knowledge Distillation","version":2},"reference_index":12,"source":"arxiv_source","source_observed_at":"2026-08-07T14:04:48.602066Z"},"links":{"cited_paper":"/paper/2312.15166","citing_paper":"/paper/2505.20237"},"observation_digest":"sha256:cd74fbd94b3175d38bb8de2f7732f992f68364b1b6858dc667aae371dc5330b2","observation_id":"00b386a9-338e-4bfc-85fc-c037922ab477","resolution":{"observed_at":"2026-08-07T14:04:48.602066Z","resolver_source":null,"status":"malformed_identifier"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2312.15166","last_updated":"2024-04-04T01:53:38Z","snapshot_observed_at":"2026-08-09T10:33:25.832634Z","submitted_at":"2023-12-23T05:11:37Z","title":"SOLAR 10.7B: Scaling Large Language Models with Simple yet Effective Depth Up-Scaling","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.15166","snapshot_observed_at":"2026-08-07T15:09:37.628409Z","title":"SOLAR 10.7B: Scaling Large Language Mod- els with Simple yet Effective Depth Up-Scaling,","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2505.21518","last_updated":"2025-05-22T05:59:36Z","snapshot_observed_at":"2026-08-10T05:54:53.949160Z","submitted_at":"2025-05-22T05:59:36Z","title":"Resilient LLM-Empowered Semantic MAC Protocols via Zero-Shot Adaptation and Knowledge Distillation","version":1},"reference_index":37,"source":"pdf_text","source_observed_at":"2026-08-07T15:09:37.628409Z"},"links":{"cited_paper":"/paper/2312.15166","citing_paper":"/paper/2505.21518"},"observation_digest":"sha256:126b15339118b7e5809419d75c7521c8e3748888a64ae2fab3c333a0137d4711","observation_id":"a3c1da30-bbfd-4f14-8662-a686fbc5facc","resolution":{"observed_at":"2026-08-07T15:09:37.628409Z","resolver_source":null,"status":"malformed_identifier"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2312.15166","last_updated":"2024-04-04T01:53:38Z","snapshot_observed_at":"2026-08-09T10:33:25.832634Z","submitted_at":"2023-12-23T05:11:37Z","title":"SOLAR 10.7B: Scaling Large Language Models with Simple yet Effective Depth Up-Scaling","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.15166","snapshot_observed_at":"2026-08-07T12:39:00.710124Z","title":"Solar 10.7 b: Scaling large language models with simple yet effective depth up-scaling","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2505.24181","last_updated":"2025-05-30T03:43:24Z","snapshot_observed_at":"2026-08-09T10:33:44.915335Z","submitted_at":"2025-05-30T03:43:24Z","title":"SCOUT: Teaching Pre-trained Language Models to Enhance Reasoning via Flow Chain-of-Thought","version":1},"reference_index":25,"source":"pdf_text","source_observed_at":"2026-08-07T12:39:00.710124Z"},"links":{"cited_paper":"/paper/2312.15166","citing_paper":"/paper/2505.24181"},"observation_digest":"sha256:af08a2e8ee29825a6fdff26bcc4473ce88406e509694168fd026354a2aa0fd42","observation_id":"42830951-779f-4436-8dae-cfbc2b4694e2","resolution":{"observed_at":"2026-08-07T12:39:00.710124Z","resolver_source":null,"status":"malformed_identifier"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2312.15166","last_updated":"2024-04-04T01:53:38Z","snapshot_observed_at":"2026-08-09T10:33:25.832634Z","submitted_at":"2023-12-23T05:11:37Z","title":"SOLAR 10.7B: Scaling Large Language Models with Simple yet Effective Depth Up-Scaling","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.15166","snapshot_observed_at":"2026-08-06T23:28:39.250664Z","title":"So- lar 10.7b: Scaling large language models with simple yet effective depth up-scaling, 2024","venue":null,"work_id":null,"year":null},"citing_paper":{"arxiv_id":"2506.18945","last_updated":"2025-06-23T02:15:43Z","snapshot_observed_at":"2026-08-09T10:33:57.788178Z","submitted_at":"2025-06-23T02:15:43Z","title":"Chain-of-Experts: Unlocking the Communication Power of Mixture-of-Experts Models","version":1},"reference_index":48,"source":"pdf_text","source_observed_at":"2026-08-06T23:28:39.250664Z"},"links":{"cited_paper":"/paper/2312.15166","citing_paper":"/paper/2506.18945"},"observation_digest":"sha256:f452346f7b2a5126fd43cf6fdbe03e24aec6848753dfee227f5ad8773ea25e70","observation_id":"370e85a8-7d92-4788-91b8-a889fbe7214d","resolution":{"observed_at":"2026-08-06T23:28:39.250664Z","resolver_source":null,"status":"malformed_identifier"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2312.15166","last_updated":"2024-04-04T01:53:38Z","snapshot_observed_at":"2026-08-09T10:33:25.832634Z","submitted_at":"2023-12-23T05:11:37Z","title":"SOLAR 10.7B: Scaling Large Language Models with Simple yet Effective Depth Up-Scaling","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.15166","snapshot_observed_at":"2026-08-05T10:33:30.244944Z","title":"Solar 10.7 b: Scaling large language models with simple yet effective depth up-scaling","venue":null,"work_id":null,"year":2023},"citing_paper":{"arxiv_id":"2509.03972","last_updated":"2025-09-04T07:56:24Z","snapshot_observed_at":"2026-08-09T10:33:45.282694Z","submitted_at":"2025-09-04T07:56:24Z","title":"Expanding Foundational Language Capabilities in Open-Source LLMs through a Korean Case Study","version":1},"reference_index":8,"source":"pdf_text","source_observed_at":"2026-08-05T10:33:30.244944Z"},"links":{"cited_paper":"/paper/2312.15166","citing_paper":"/paper/2509.03972"},"observation_digest":"sha256:073551bd5e8a31bd31983aed82ae8e11196abfcf1d8fadb68e7da2667c083d9c","observation_id":"8e1aac1d-3f89-4cc4-8b9f-a4ba61f04de2","resolution":{"observed_at":"2026-08-05T10:33:30.244944Z","resolver_source":null,"status":"malformed_identifier"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}},{"citation":{"cited_paper":{"arxiv_id":"2312.15166","last_updated":"2024-04-04T01:53:38Z","snapshot_observed_at":"2026-08-09T10:33:25.832634Z","submitted_at":"2023-12-23T05:11:37Z","title":"SOLAR 10.7B: Scaling Large Language Models with Simple yet Effective Depth Up-Scaling","version":3},"cited_work":{"arxiv_id":"2312.15166","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2312.15166","snapshot_observed_at":"2026-07-02T16:07:09.407610Z","title":"Solar 10.7 b: Scaling large language models with simple yet effective depth up-scaling.arXiv preprint arXiv:2312.15166","venue":null,"work_id":"12c284dc-a58b-4942-8c65-c730a8a04685","year":2023},"citing_paper":{"arxiv_id":"2512.13751","last_updated":"2026-05-09T08:34:38Z","snapshot_observed_at":"2026-08-08T20:22:29.940148Z","submitted_at":"2025-12-15T05:50:45Z","title":"MIDUS: Memory-Infused Depth Up-Scaling","version":2},"reference_index":14,"source":"pdf_text","source_observed_at":"2026-05-16T22:02:42.297041Z"},"links":{"cited_paper":"/paper/2312.15166","citing_paper":"/paper/2512.13751"},"observation_digest":"sha256:2e5e17e8ad67b730aac5f66b78e636cc0004f8d1accc8fc8f9242f6b77c47567","observation_id":"424a487e-b4cb-4a64-ab4d-9eb99d4a7831","resolution":{"observed_at":"2026-05-16T22:03:36.128384Z","resolver_source":"arxiv_id","status":"malformed_identifier"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2312.15166","last_updated":"2024-04-04T01:53:38Z","snapshot_observed_at":"2026-08-09T10:33:25.832634Z","submitted_at":"2023-12-23T05:11:37Z","title":"SOLAR 10.7B: Scaling Large Language Models with Simple yet Effective Depth Up-Scaling","version":3},"cited_work":{"arxiv_id":"2312.15166","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2312.15166","snapshot_observed_at":"2026-07-02T16:07:09.407610Z","title":"Solar 10.7 b: Scaling large language models with simple yet effective depth up-scaling.arXiv preprint arXiv:2312.15166","venue":null,"work_id":"12c284dc-a58b-4942-8c65-c730a8a04685","year":2023},"citing_paper":{"arxiv_id":"2603.13224","last_updated":"2026-05-09T03:24:22Z","snapshot_observed_at":"2026-08-05T18:39:49.396872Z","submitted_at":"2026-03-13T17:58:14Z","title":"Visual-ERM: Reward Modeling for Visual Equivalence","version":2},"reference_index":11,"source":"pdf_text","source_observed_at":"2026-05-15T11:19:42.002790Z"},"links":{"cited_paper":"/paper/2312.15166","citing_paper":"/paper/2603.13224"},"observation_digest":"sha256:55a00112bb1a2a053cbe96bfcb05ab149c7c8948400af625e10fe65db55f1d3e","observation_id":"a4c31e22-759a-45d1-9bcf-cc0262ee1fb6","resolution":{"observed_at":"2026-05-15T11:19:57.913023Z","resolver_source":"arxiv_id","status":"malformed_identifier"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2312.15166","last_updated":"2024-04-04T01:53:38Z","snapshot_observed_at":"2026-08-09T10:33:25.832634Z","submitted_at":"2023-12-23T05:11:37Z","title":"SOLAR 10.7B: Scaling Large Language Models with Simple yet Effective Depth Up-Scaling","version":3},"cited_work":{"arxiv_id":"2312.15166","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2312.15166","snapshot_observed_at":"2026-07-02T16:07:09.407610Z","title":"Solar 10.7 b: Scaling large language models with simple yet effective depth up-scaling.arXiv preprint arXiv:2312.15166","venue":null,"work_id":"12c284dc-a58b-4942-8c65-c730a8a04685","year":2023},"citing_paper":{"arxiv_id":"2604.12391","last_updated":"2026-04-14T07:26:23Z","snapshot_observed_at":"2026-07-06T23:00:37.144068Z","submitted_at":"2026-04-14T07:26:23Z","title":"Chain-of-Models Pre-Training: Rethinking Training Acceleration of Vision Foundation Models","version":1},"reference_index":30,"source":"pdf_text","source_observed_at":"2026-05-10T16:16:51.109113Z"},"links":{"cited_paper":"/paper/2312.15166","citing_paper":"/paper/2604.12391"},"observation_digest":"sha256:1c724c4144609bbf32fc87f8d30617f56313dbf168088b65cf73a3c388685c90","observation_id":"8994c557-b41a-43e8-864d-602470089b39","resolution":{"observed_at":"2026-05-11T09:05:58.772360Z","resolver_source":"arxiv_id","status":"malformed_identifier"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2312.15166","last_updated":"2024-04-04T01:53:38Z","snapshot_observed_at":"2026-08-09T10:33:25.832634Z","submitted_at":"2023-12-23T05:11:37Z","title":"SOLAR 10.7B: Scaling Large Language Models with Simple yet Effective Depth Up-Scaling","version":3},"cited_work":{"arxiv_id":"2312.15166","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2312.15166","snapshot_observed_at":"2026-07-02T16:07:09.407610Z","title":"Solar 10.7 b: Scaling large language models with simple yet effective depth up-scaling.arXiv preprint arXiv:2312.15166","venue":null,"work_id":"12c284dc-a58b-4942-8c65-c730a8a04685","year":2023},"citing_paper":{"arxiv_id":"2605.05662","last_updated":"2026-05-07T04:35:03Z","snapshot_observed_at":"2026-07-06T23:18:17.333660Z","submitted_at":"2026-05-07T04:35:03Z","title":"XL-SafetyBench: A Country-Grounded Cross-Cultural Benchmark for LLM Safety and Cultural Sensitivity","version":1},"reference_index":23,"source":"pdf_text","source_observed_at":"2026-05-08T11:08:22.322879Z"},"links":{"cited_paper":"/paper/2312.15166","citing_paper":"/paper/2605.05662"},"observation_digest":"sha256:9e6ddf080b9b4b2c8ad04a590bb17909c166375a84579621103d3818b344a291","observation_id":"6ae07fc5-2413-441b-ac0c-e5e6ee13d588","resolution":{"observed_at":"2026-05-11T19:41:11.046774Z","resolver_source":"arxiv_id","status":"malformed_identifier"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2312.15166","last_updated":"2024-04-04T01:53:38Z","snapshot_observed_at":"2026-08-09T10:33:25.832634Z","submitted_at":"2023-12-23T05:11:37Z","title":"SOLAR 10.7B: Scaling Large Language Models with Simple yet Effective Depth Up-Scaling","version":3},"cited_work":{"arxiv_id":"2312.15166","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2312.15166","snapshot_observed_at":"2026-07-02T16:07:09.407610Z","title":"Solar 10.7 b: Scaling large language models with simple yet effective depth up-scaling.arXiv preprint arXiv:2312.15166","venue":null,"work_id":"12c284dc-a58b-4942-8c65-c730a8a04685","year":2023},"citing_paper":{"arxiv_id":"2606.04909","last_updated":"2026-06-03T14:06:27Z","snapshot_observed_at":"2026-08-03T00:26:15.308076Z","submitted_at":"2026-06-03T14:06:27Z","title":"BEATS: Bootstrapping E-commerce Attribute Taxonomies for Search through Iterative Human-AI Collaboration","version":1},"reference_index":6,"source":"pdf_text","source_observed_at":"2026-06-28T04:02:38.834181Z"},"links":{"cited_paper":"/paper/2312.15166","citing_paper":"/paper/2606.04909"},"observation_digest":"sha256:40d89edcd0cff0e3e01c950a3e16b649de1729c18b23bb74c96fc6c68170e7f6","observation_id":"d40349d5-5a06-4aaf-ac78-ed3e33a2493b","resolution":{"observed_at":"2026-07-02T11:26:54.195736Z","resolver_source":"arxiv_id","status":"malformed_identifier"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2312.15166","last_updated":"2024-04-04T01:53:38Z","snapshot_observed_at":"2026-08-09T10:33:25.832634Z","submitted_at":"2023-12-23T05:11:37Z","title":"SOLAR 10.7B: Scaling Large Language Models with Simple yet Effective Depth Up-Scaling","version":3},"cited_work":{"arxiv_id":"2312.15166","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2312.15166","snapshot_observed_at":"2026-07-02T16:07:09.407610Z","title":"Solar 10.7 b: Scaling large language models with simple yet effective depth up-scaling.arXiv preprint arXiv:2312.15166","venue":null,"work_id":"12c284dc-a58b-4942-8c65-c730a8a04685","year":2023},"citing_paper":{"arxiv_id":"2606.07404","last_updated":"2026-06-05T15:48:42Z","snapshot_observed_at":"2026-07-06T23:47:05.104295Z","submitted_at":"2026-06-05T15:48:42Z","title":"Reversible Foundations: Training a 120B Sparse MoE through State-Preserving Scaling","version":1},"reference_index":47,"source":"arxiv_source","source_observed_at":"2026-06-27T22:55:09.477413Z"},"links":{"cited_paper":"/paper/2312.15166","citing_paper":"/paper/2606.07404"},"observation_digest":"sha256:c9a0059295a69fb9855bca3756806db4ee19f5fa564fd0d7d0f2abf9dede5f5a","observation_id":"a529809d-f6b0-4389-a74c-2ab7c4d9a8f0","resolution":{"observed_at":"2026-07-02T16:07:09.409007Z","resolver_source":"arxiv_id","status":"malformed_identifier"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2312.15166","last_updated":"2024-04-04T01:53:38Z","snapshot_observed_at":"2026-08-09T10:33:25.832634Z","submitted_at":"2023-12-23T05:11:37Z","title":"SOLAR 10.7B: Scaling Large Language Models with Simple yet Effective Depth Up-Scaling","version":3},"cited_work":{"arxiv_id":"2312.15166","doi":null,"metadata_source":"arxiv_reference","pith_arxiv_id":"2312.15166","snapshot_observed_at":"2026-07-02T16:07:09.407610Z","title":"Solar 10.7 b: Scaling large language models with simple yet effective depth up-scaling.arXiv preprint arXiv:2312.15166","venue":null,"work_id":"12c284dc-a58b-4942-8c65-c730a8a04685","year":2023},"citing_paper":{"arxiv_id":"2606.31796","last_updated":"2026-07-23T16:29:18Z","snapshot_observed_at":"2026-08-09T07:09:32.790592Z","submitted_at":"2026-06-30T15:14:38Z","title":"CHERRY: Compressed Hierarchical Experts with Recurrent Representational Yield","version":1},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-07-01T05:46:40.510955Z"},"links":{"cited_paper":"/paper/2312.15166","citing_paper":"/paper/2606.31796"},"observation_digest":"sha256:5121355dc781ea91c66516f1cfa8e1e7f1009e40a861e8e3806b64625e672943","observation_id":"b5a252e4-1286-4eb0-99f9-bf436fd4440e","resolution":{"observed_at":"2026-07-01T10:15:44.059441Z","resolver_source":"arxiv_id","status":"malformed_identifier"},"standing_notice":{"events":[],"observation":"No event found in the named queried sources as of 2026-08-10T06:31:04.303077+00:00.","reason":null,"source_receipts":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"state":"measured"}},{"citation":{"cited_paper":{"arxiv_id":"2312.15166","last_updated":"2024-04-04T01:53:38Z","snapshot_observed_at":"2026-08-09T10:33:25.832634Z","submitted_at":"2023-12-23T05:11:37Z","title":"SOLAR 10.7B: Scaling Large Language Models with Simple yet Effective Depth Up-Scaling","version":3},"cited_work":{"arxiv_id":null,"doi":null,"metadata_source":null,"pith_arxiv_id":"2312.15166","snapshot_observed_at":"2026-08-02T09:24:13.924447Z","title":"NAACL-HLT Industry Track(2024)","venue":null,"work_id":null,"year":2024},"citing_paper":{"arxiv_id":"2606.31796","last_updated":"2026-07-23T16:29:18Z","snapshot_observed_at":"2026-08-09T07:09:32.790592Z","submitted_at":"2026-06-30T15:14:38Z","title":"CHERRY: Compressed Hierarchical Experts with Recurrent Representational Yield","version":2},"reference_index":19,"source":"pdf_text","source_observed_at":"2026-08-02T09:24:13.924447Z"},"links":{"cited_paper":"/paper/2312.15166","citing_paper":"/paper/2606.31796"},"observation_digest":"sha256:b34a2c15b698622676234ed9490205656b609821186028b109b420176d2223fc","observation_id":"e2795c3d-b840-4ccc-a228-dfc81ce0b248","resolution":{"observed_at":"2026-08-02T09:24:13.924447Z","resolver_source":null,"status":"malformed_identifier"},"standing_notice":{"events":[],"reason":"canonical_work_link_unavailable","source_receipts":[],"state":"unavailable"}}],"links":{"evidence":"/evidence","html":"/paper/2312.15166/citation-record","integrity":"/paper/2312.15166/integrity","json":"/paper/2312.15166/citation-record.json","paper":"/paper/2312.15166"},"outbound":[],"paper":{"arxiv_id":"2312.15166","last_updated":"2024-04-04T01:53:38Z","latest_version":3,"primary_category":"cs.CL","snapshot_observed_at":"2026-08-09T10:33:25.832634Z","submitted_at":"2023-12-23T05:11:37Z","title":"SOLAR 10.7B: Scaling Large Language Models with Simple yet Effective Depth Up-Scaling"},"reference_resolution":{"displayed":0,"state_counts":{"malformed_identifier":0,"metadata_mismatch":0,"parse_uncertain":0,"unresolved":0,"verified_exact":0,"verified_fuzzy":0},"total_outbound_references":0},"refusal":"A citation records a reference. It does not transfer a finding from one paper to another.","schema":"pith.paper-citation-record.v1","standing_sources":[{"observed_at":"2026-08-10T06:31:04.303077+00:00","source":"crossref"},{"observed_at":"2026-08-10T06:30:57.382061+00:00","source":"retraction_watch"}],"thesis":"As of 10 August 2026, this Paper Citation Record lists 0 of 0 outbound references and 20 inbound Pith citation observations for arXiv:2312.15166."}