<?xml version="1.0" encoding="utf-8" standalone="yes"?>
<rss version="2.0" xmlns:atom="http://www.w3.org/2005/Atom">
	<channel>
		<title>Inference on libcom.de</title>
		<link>https://www.libcom.de/en/tags/inference/</link>
		<description>Recent content in Inference on libcom.de</description>
		<generator>Hugo</generator>
		<language>en-GB</language>
		
		
		
		
			<lastBuildDate>Tue, 18 Aug 2026 00:00:00 +0000</lastBuildDate>
		
			<atom:link href="https://www.libcom.de/en/tags/inference/index.xml" rel="self" type="application/rss+xml" />
			<item>
				<title>Local AI: From Desk to Inference Server</title>
				<link>https://www.libcom.de/en/blog/local-ai-from-desktop-to-server/</link>
				<pubDate>Tue, 18 Aug 2026 00:00:00 +0000</pubDate>
				<guid>https://www.libcom.de/en/blog/local-ai-from-desktop-to-server/</guid>
				<description>&lt;p&gt;The most revealing metric for local AI is not printed boldly on the box. It is called memory bandwidth, it is measured in gigabytes per second, and it decides whether a model runs smoothly on your machine or stutters slowly to death. Anyone serious about experimenting with AI who leans on the TOPS and NPU figures is looking in the wrong place. Compute has grown cheap. Bandwidth is the scarce resource.&lt;/p&gt;</description>
			</item>
	</channel>
</rss>
