{"slug": "running-a-35b-llm-at-128k-context-full-speed-on-eur870-of-used-hardware", "title": "Running a 35B LLM at 128K Context, Full Speed, on €870 of Used Hardware", "summary": "A developer demonstrated running a 35B-parameter large language model with 128K context at full speed using only €870 of used hardware, eliminating the need for cloud services. The setup leverages cost-effective second-hand components to achieve high-performance inference, highlighting the feasibility of local AI deployment.", "body_md": "Article URL: https://medium.com/ai-advances/running-a-35b-llm-at-128k-context-full-speed-on-870-of-used-hardware-no-cloud-required-c4f7629810b8\n\nComments URL: https://news.ycombinator.com/item?id=49142301\n\nPoints: 1\n\n# Comments: 0", "url": "https://wpnews.pro/news/running-a-35b-llm-at-128k-context-full-speed-on-eur870-of-used-hardware", "canonical_source": "https://medium.com/ai-advances/running-a-35b-llm-at-128k-context-full-speed-on-870-of-used-hardware-no-cloud-required-c4f7629810b8", "published_at": "2026-08-02 08:27:38+00:00", "updated_at": "2026-08-02 08:52:32.026351+00:00", "lang": "en", "topics": ["large-language-models", "ai-infrastructure", "ai-tools"], "entities": [], "alternates": {"html": "https://wpnews.pro/news/running-a-35b-llm-at-128k-context-full-speed-on-eur870-of-used-hardware", "markdown": "https://wpnews.pro/news/running-a-35b-llm-at-128k-context-full-speed-on-eur870-of-used-hardware.md", "text": "https://wpnews.pro/news/running-a-35b-llm-at-128k-context-full-speed-on-eur870-of-used-hardware.txt", "jsonld": "https://wpnews.pro/news/running-a-35b-llm-at-128k-context-full-speed-on-eur870-of-used-hardware.jsonld"}}