Code
Hub
Workspaces
Following
Trending
Connect
MCP
copy
Create free account
hub
/
github.com/ml-explore/mlx-lm
/ functions
Functions
2,622 in github.com/ml-explore/mlx-lm
⨍
Functions
2,622
◇
Types & classes
816
↳
Endpoints
9
Method
__post_init__
(self)
mlx_lm/models/mamba.py:30
Method
__post_init__
(self)
mlx_lm/models/ministral3.py:34
Method
__post_init__
(self)
mlx_lm/models/hunyuan_v1_dense.py:31
Method
__post_init__
(self)
mlx_lm/models/mixtral.py:30
Method
__post_init__
(self)
mlx_lm/models/gpt_bigcode.py:29
Method
__post_init__
(self)
mlx_lm/models/phi.py:26
Method
__post_init__
(self)
mlx_lm/models/mistral3.py:19
Method
__post_init__
(self)
mlx_lm/models/hunyuan.py:37
Method
__post_init__
(self)
mlx_lm/models/gpt_neox.py:29
Method
__post_init__
(self)
mlx_lm/models/nemotron-nas.py:27
Method
__post_init__
(self)
mlx_lm/models/nemotron-nas.py:53
Method
__post_init__
(self)
mlx_lm/models/nemotron-nas.py:118
Method
__post_init__
(self)
mlx_lm/models/jamba.py:47
Method
__post_init__
(self)
mlx_lm/models/mamba2.py:37
Method
__post_init__
(self)
mlx_lm/models/internlm3.py:31
Method
__post_init__
(self)
mlx_lm/models/internlm2.py:30
Method
__post_init__
(self)
mlx_lm/models/olmo.py:34
Method
__post_init__
(self)
mlx_lm/models/qwen.py:26
Method
__post_init__
(self)
mlx_lm/models/olmoe.py:36
Method
__post_init__
(self)
mlx_lm/models/pixtral.py:19
Method
__post_init__
(self)
mlx_lm/models/nemotron.py:34
Method
__post_init__
(self)
mlx_lm/models/phi3.py:32
Method
__post_init__
(self)
mlx_lm/models/nemotron_h.py:63
Method
__repr__
(self)
mlx_lm/gguf.py:94
Method
__setitem__
(self, idx, value)
mlx_lm/models/cache.py:618
Function
_capture
(m)
mlx_lm/tool_parsers/gemma4.py:31
Method
_extra_repr
(self)
mlx_lm/models/minimax.py:68
Method
_generate
(self)
mlx_lm/server.py:688
Function
_get_classes
Retrieve the model and model args classes based on the configuration. Args: config (dict): The model configuration. Returns:
mlx_lm/utils.py:175
Method
_inner
()
mlx_lm/server.py:1035
Function
_inner_sharded_rms_norm
(x, w, eps)
mlx_lm/models/minimax.py:52
Function
_left_pad_prompts
(prompts, max_length=None)
mlx_lm/generate.py:802
Function
_make_cache
Convert a list of regular caches into their corresponding batch-aware caches.
mlx_lm/generate.py:838
Function
_maybe_qq
(m)
mlx_lm/utils.py:394
Function
_parse_size
(x)
mlx_lm/utils.py:60
Method
_r
(text, state, match=None)
tests/test_server.py:97
Method
_repeat
(p)
mlx_lm/models/qwen3_5.py:415
Method
add_token
(self, token)
mlx_lm/tokenizer_utils.py:81
Method
add_token
(self, token)
mlx_lm/tokenizer_utils.py:144
Method
add_token
(self, token)
mlx_lm/tokenizer_utils.py:206
Method
apply_block_linear
(h, w, b)
mlx_lm/models/recurrent_gemma.py:136
Function
apply_chat_template
(self, chat_history, add_generation_prompt=True)
mlx_lm/evaluate.py:59
Function
apply_chat_template
( messages, continue_final_message=False, add_generation_prompt=False, **kwargs )
mlx_lm/chat_templates/deepseek_v32.py:333
Function
apply_clip
(path, module)
mlx_lm/quant/awq.py:384
Method
as_linear
(self, x)
mlx_lm/tuner/lora.py:282
Function
batch_bench
()
mlx_lm/benchmark.py:128
Method
batch_decode
(self, *args, **kwargs)
mlx_lm/tokenizer_utils.py:500
Method
batch_size
(self)
mlx_lm/models/cache.py:607
Function
callback
(processed, total_tokens)
mlx_lm/cache_prompt.py:118
Function
capture
(module)
mlx_lm/quant/awq.py:426
Method
cast_predicate
(self)
mlx_lm/models/kimi_k25.py:79
Method
cast_predicate
(self)
mlx_lm/models/minimax.py:381
Method
cast_predicate
(self)
mlx_lm/models/bailing_moe_linear.py:578
Method
cast_predicate
(self)
mlx_lm/models/lfm2_moe.py:383
Method
cast_predicate
(self)
mlx_lm/models/glm4_moe.py:399
Method
cast_predicate
(self)
mlx_lm/models/mimo_v2_flash.py:371
Method
cast_predicate
(self)
mlx_lm/models/kimi_linear.py:594
Method
cast_predicate
(self)
mlx_lm/models/step3p5.py:443
Method
cast_predicate
(self)
mlx_lm/models/exaone_moe.py:368
Method
cast_predicate
(self)
mlx_lm/models/glm4_moe_lite.py:527
Method
cast_predicate
(self)
mlx_lm/models/qwen3_5.py:346
Method
cast_predicate
(self)
mlx_lm/models/qwen3_5.py:530
Method
cast_predicate
(self)
mlx_lm/models/deepseek_v3.py:549
Method
cast_predicate
(self)
mlx_lm/models/kimi_vl.py:112
Method
cast_predicate
(self)
mlx_lm/models/deepseek_v32.py:647
Method
cast_predicate
(self)
mlx_lm/models/Klear.py:259
Method
cast_predicate
(self)
mlx_lm/models/bailing_moe.py:393
Method
cast_predicate
(self)
mlx_lm/models/longcat_flash_ngram.py:197
Method
cast_predicate
(self)
mlx_lm/models/longcat_flash.py:385
Method
cast_predicate
(self)
mlx_lm/models/afmoe.py:392
Method
cast_predicate
(self)
mlx_lm/models/nemotron_h.py:558
Method
cat
(a, b)
mlx_lm/models/cache.py:650
Method
check
(tokens)
tests/test_tokenizers.py:19
Method
check_config
(params, expected_trainable_parameters=None)
tests/test_finetune.py:46
Method
check_config
(params, expected_params=None)
tests/test_finetune.py:189
Function
checkpointed_fn
(model, *args, **kwargs)
mlx_lm/tuner/trainer.py:31
Function
class_predicate
(p, m)
mlx_lm/utils.py:349
Method
cli_args
(self)
mlx_lm/server.py:1055
Function
common_prefix_len
Calculates the length of the common prefix of two lists. Args: list1: The first list of strings. list2: The second list of s
mlx_lm/utils.py:953
Function
compute_sensitivity
(gradient, low_q_weight, original_weight)
mlx_lm/quant/dynamic_quant.py:88
Method
conv_sharding
(key_dim)
mlx_lm/models/qwen3_5.py:406
Method
custom_get_classes
(config)
tests/test_utils.py:116
Function
decode_dsml_to_arguments
( tool_name: str, tool_args: Dict[str, Tuple[str, str]] )
mlx_lm/chat_templates/deepseek_v32.py:113
Function
default_loss
(model, batch, lengths)
mlx_lm/tuner/trainer.py:86
Method
default_scale
(factor)
mlx_lm/models/rope_utils.py:52
Method
dequant
(weight, scale_inv)
mlx_lm/models/minimax.py:279
Method
dequant
(weight, scale_inv)
mlx_lm/models/mimo_v2_flash.py:322
Method
dequant
(weight, scale_inv)
mlx_lm/models/deepseek_v3.py:381
Method
dequant
(weight, scale_inv)
mlx_lm/models/deepseek_v32.py:505
Method
dequantize_weight
(quantized_linear)
tests/test_finetune.py:291
Method
detokenizer
Get a stateful streaming detokenizer.
mlx_lm/tokenizer_utils.py:451
Method
do_GET
Respond to a GET request from a client.
mlx_lm/server.py:1620
Method
do_OPTIONS
(self)
mlx_lm/server.py:1097
Method
do_POST
Respond to a POST request from a client.
mlx_lm/server.py:1101
Method
empty
( cls, model: nn.Module, fallback_sampler: Callable[[mx.array], mx.array], pre
mlx_lm/generate.py:1209
Method
empty
Return if the cache is empty or not.
mlx_lm/models/cache.py:163
Method
empty
(self)
mlx_lm/models/cache.py:222
Method
empty
(self)
mlx_lm/models/cache.py:317
Method
empty
(self)
mlx_lm/models/cache.py:584
Method
empty
(self)
mlx_lm/models/cache.py:723
← previous
next →
1,801–1,900 of 2,622, ranked by callers