File size: 465 Bytes

daf9abb

quant_stage:
  quant_modifiers:
    QuantizationModifier:
      targets: [Linear]
      ignore: []
      kv_cache_scheme:
        num_bits: 8
        type: float
        symmetric: true
        group_size: null
        strategy: tensor
        block_structure: null
        dynamic: false
        actorder: null
        scale_dtype: null
        zp_dtype: null
        observer: memoryless_minmax
        observer_kwargs: {}
      bypass_divisibility_checks: false