{"model":{"id":"qwen3.5-9b-q8","name":"Qwen3.5 9B Instruct (Q8)","quantization":"Q8_0"},"hardware":{"slug":"rtx-4060","name":"NVIDIA GeForce RTX 4060","memoryGb":8},"verdict":"tight","localVerdict":"local_slow","fitLevel":"Heavy","estimatedLoadGb":10.7,"usableMemoryGb":7.2,"estimatedTokensPerSec":3.4,"kvCacheByContext":[{"contextTokens":8192,"kvGb":0.25,"totalGb":10.95,"fits":false},{"contextTokens":16384,"kvGb":0.5,"totalGb":11.2,"fits":false},{"contextTokens":32768,"kvGb":1,"totalGb":11.7,"fits":false},{"contextTokens":65536,"kvGb":2,"totalGb":12.7,"fits":false},{"contextTokens":131072,"kvGb":4,"totalGb":14.7,"fits":false}],"upgradeGpu":{"slug":"rx-7900-xt","name":"AMD Radeon RX 7900 XT"},"ollamaCommand":"ollama run qwen3.5:9b-q8_0","reportUrl":"https://modelfit.io/can-i-run/qwen3.5-9b-q8-on-rtx-4060/","source":"https://modelfit.io/","license":"CC BY 4.0 (https://creativecommons.org/licenses/by/4.0/)","caveat":"Estimates from the ModelFit engine, not measurements."}