diff --git a/resources/inference_model_cards/mlx-community--Qwen3-Coder-Next-4bit.toml b/resources/inference_model_cards/mlx-community--Qwen3-Coder-Next-4bit.toml new file mode 100644 index 00000000..9e6d3511 --- /dev/null +++ b/resources/inference_model_cards/mlx-community--Qwen3-Coder-Next-4bit.toml @@ -0,0 +1,8 @@ +model_id = "mlx-community/Qwen3-Coder-Next-4bit" +n_layers = 48 +hidden_size = 2048 +supports_tensor = true +tasks = ["TextGeneration"] + +[storage_size] +in_bytes = 45644286500 diff --git a/resources/inference_model_cards/mlx-community--Qwen3-Coder-Next-5bit.toml b/resources/inference_model_cards/mlx-community--Qwen3-Coder-Next-5bit.toml new file mode 100644 index 00000000..ebf487af --- /dev/null +++ b/resources/inference_model_cards/mlx-community--Qwen3-Coder-Next-5bit.toml @@ -0,0 +1,8 @@ +model_id = "mlx-community/Qwen3-Coder-Next-5bit" +n_layers = 48 +hidden_size = 2048 +supports_tensor = true +tasks = ["TextGeneration"] + +[storage_size] +in_bytes = 57657697020 diff --git a/resources/inference_model_cards/mlx-community--Qwen3-Coder-Next-6bit.toml b/resources/inference_model_cards/mlx-community--Qwen3-Coder-Next-6bit.toml new file mode 100644 index 00000000..ef801631 --- /dev/null +++ b/resources/inference_model_cards/mlx-community--Qwen3-Coder-Next-6bit.toml @@ -0,0 +1,8 @@ +model_id = "mlx-community/Qwen3-Coder-Next-6bit" +n_layers = 48 +hidden_size = 2048 +supports_tensor = true +tasks = ["TextGeneration"] + +[storage_size] +in_bytes = 68899327465 diff --git a/resources/inference_model_cards/mlx-community--Qwen3-Coder-Next-8bit.toml b/resources/inference_model_cards/mlx-community--Qwen3-Coder-Next-8bit.toml new file mode 100644 index 00000000..05b92116 --- /dev/null +++ b/resources/inference_model_cards/mlx-community--Qwen3-Coder-Next-8bit.toml @@ -0,0 +1,8 @@ +model_id = "mlx-community/Qwen3-Coder-Next-8bit" +n_layers = 48 +hidden_size = 2048 +supports_tensor = true +tasks = ["TextGeneration"] + +[storage_size] +in_bytes = 89357758772 diff --git a/resources/inference_model_cards/mlx-community--Qwen3-Coder-Next-bf16.toml b/resources/inference_model_cards/mlx-community--Qwen3-Coder-Next-bf16.toml new file mode 100644 index 00000000..93798231 --- /dev/null +++ b/resources/inference_model_cards/mlx-community--Qwen3-Coder-Next-bf16.toml @@ -0,0 +1,8 @@ +model_id = "mlx-community/Qwen3-Coder-Next-bf16" +n_layers = 48 +hidden_size = 2048 +supports_tensor = true +tasks = ["TextGeneration"] + +[storage_size] +in_bytes = 157548627945