Upload pruned model

Browse files

Files changed (9) hide show

1_Pooling/config.json +7 -0
README.md +43 -0
config.json +50 -0
model.safetensors +3 -0
modules.json +20 -0
sentence_bert_config.json +4 -0
special_tokens_map.json +51 -0
tokenizer.json +0 -0
tokenizer_config.json +54 -0

1_Pooling/config.json ADDED Viewed

	@@ -0,0 +1,7 @@

+{
+ "word_embedding_dimension": 768,
+ "pooling_mode_cls_token": true,
+ "pooling_mode_mean_tokens": false,
+ "pooling_mode_max_tokens": false,
+ "pooling_mode_mean_sqrt_len_tokens": false
+}

README.md ADDED Viewed

	@@ -0,0 +1,43 @@

+---
+pipeline_tag: sentence-similarity
+language: fr
+license: apache-2.0
+tags:
+- passage-retrieval
+- sentence-similarity
+- pruned
+library_name: sentence-transformers
+base_model: Alibaba-NLP/gte-multilingual-base
+base_model_relation: quantized
+---
+# 🇫🇷 french-gte-multilingual-base
+This model is a 53.5% smaller version of [Alibaba-NLP/gte-multilingual-base](https://huggingface.co/Alibaba-NLP/gte-multilingual-base)
+for the French language, created using the [mtem-pruner](https://huggingface.co/spaces/antoinelouis/mtem-pruner) space.
+This pruned model should perform similarly to the original model for French language tasks with a much smaller
+memory footprint. However, it may not perform well for other languages present in the original multilingual model as tokens not
+commonly used in French were removed from the original multilingual model's vocabulary.
+## Usage
+You can use this model with the Transformers library:
+```python
+from transformers import AutoModel, AutoTokenizer
+model_name = "antoinelouis/french-gte-multilingual-base"
+model = AutoModel.from_pretrained(model_name, trust_remote_code=True)
+tokenizer = AutoTokenizer.from_pretrained(model_name, trust_remote_code=True, use_fast=True)
+```
+Or with the sentence-transformers library:
+```python
+from sentence_transformers import SentenceTransformer
+model = SentenceTransformer("antoinelouis/french-gte-multilingual-base")
+```
+**Credits**: cc [@antoinelouis](https://huggingface.co/antoinelouis)

config.json ADDED Viewed

	@@ -0,0 +1,50 @@

+{
+ "_name_or_path": "Alibaba-NLP/gte-multilingual-base",
+ "architectures": [
+ "NewModel"
+ ],
+ "attention_probs_dropout_prob": 0.0,
+ "auto_map": {
+ "AutoConfig": "Alibaba-NLP/new-impl--configuration.NewConfig",
+ "AutoModel": "Alibaba-NLP/new-impl--modeling.NewModel",
+ "AutoModelForMaskedLM": "Alibaba-NLP/new-impl--modeling.NewForMaskedLM",
+ "AutoModelForMultipleChoice": "Alibaba-NLP/new-impl--modeling.NewForMultipleChoice",
+ "AutoModelForQuestionAnswering": "Alibaba-NLP/new-impl--modeling.NewForQuestionAnswering",
+ "AutoModelForSequenceClassification": "Alibaba-NLP/new-impl--modeling.NewForSequenceClassification",
+ "AutoModelForTokenClassification": "Alibaba-NLP/new-impl--modeling.NewForTokenClassification"
+ },
+ "classifier_dropout": 0.0,
+ "hidden_act": "gelu",
+ "hidden_dropout_prob": 0.1,
+ "hidden_size": 768,
+ "id2label": {
+ "0": "LABEL_0"
+ },
+ "initializer_range": 0.02,
+ "intermediate_size": 3072,
+ "label2id": {
+ "LABEL_0": 0
+ },
+ "layer_norm_eps": 1e-12,
+ "layer_norm_type": "layer_norm",
+ "logn_attention_clip1": false,
+ "logn_attention_scale": false,
+ "max_position_embeddings": 8192,
+ "model_type": "new",
+ "num_attention_heads": 12,
+ "num_hidden_layers": 12,
+ "pack_qkv": true,
+ "pad_token_id": 1,
+ "position_embedding_type": "rope",
+ "rope_scaling": {
+ "factor": 8.0,
+ "type": "ntk"
+ },
+ "rope_theta": 20000,
+ "torch_dtype": "float32",
+ "transformers_version": "4.45.1",
+ "type_vocab_size": 1,
+ "unpad_inputs": false,
+ "use_memory_efficient_attention": false,
+ "vocab_size": 37200
+}

model.safetensors ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:5b34f6f8e4686505d2a44fd944bbd38ace6b922cdde73a43e2088ad6a9985c81
+size 567618688

modules.json ADDED Viewed

	@@ -0,0 +1,20 @@

+[
+ {
+ "idx": 0,
+ "name": "0",
+ "path": "",
+ "type": "sentence_transformers.models.Transformer"
+ },
+ {
+ "idx": 1,
+ "name": "1",
+ "path": "1_Pooling",
+ "type": "sentence_transformers.models.Pooling"
+ },
+ {
+ "idx": 2,
+ "name": "2",
+ "path": "2_Normalize",
+ "type": "sentence_transformers.models.Normalize"
+ }
+]

sentence_bert_config.json ADDED Viewed

	@@ -0,0 +1,4 @@

+{
+ "max_seq_length": 8192,
+ "do_lower_case": false
+}

special_tokens_map.json ADDED Viewed

	@@ -0,0 +1,51 @@

+{
+ "bos_token": {
+ "content": "<s>",
+ "lstrip": false,
+ "normalized": false,
+ "rstrip": false,
+ "single_word": false
+ },
+ "cls_token": {
+ "content": "<s>",
+ "lstrip": false,
+ "normalized": false,
+ "rstrip": false,
+ "single_word": false
+ },
+ "eos_token": {
+ "content": "</s>",
+ "lstrip": false,
+ "normalized": false,
+ "rstrip": false,
+ "single_word": false
+ },
+ "mask_token": {
+ "content": "<mask>",
+ "lstrip": true,
+ "normalized": false,
+ "rstrip": false,
+ "single_word": false
+ },
+ "pad_token": {
+ "content": "<pad>",
+ "lstrip": false,
+ "normalized": false,
+ "rstrip": false,
+ "single_word": false
+ },
+ "sep_token": {
+ "content": "</s>",
+ "lstrip": false,
+ "normalized": false,
+ "rstrip": false,
+ "single_word": false
+ },
+ "unk_token": {
+ "content": "<unk>",
+ "lstrip": false,
+ "normalized": false,
+ "rstrip": false,
+ "single_word": false
+ }
+}

tokenizer.json ADDED Viewed

The diff for this file is too large to render. See raw diff

tokenizer_config.json ADDED Viewed

	@@ -0,0 +1,54 @@

+{
+ "added_tokens_decoder": {
+ "0": {
+ "content": "<s>",
+ "lstrip": false,
+ "normalized": false,
+ "rstrip": false,
+ "single_word": false,
+ "special": true
+ },
+ "1": {
+ "content": "<pad>",
+ "lstrip": false,
+ "normalized": false,
+ "rstrip": false,
+ "single_word": false,
+ "special": true
+ },
+ "2": {
+ "content": "</s>",
+ "lstrip": false,
+ "normalized": false,
+ "rstrip": false,
+ "single_word": false,
+ "special": true
+ },
+ "3": {
+ "content": "<unk>",
+ "lstrip": false,
+ "normalized": false,
+ "rstrip": false,
+ "single_word": false,
+ "special": true
+ },
+ "37199": {
+ "content": "<mask>",
+ "lstrip": true,
+ "normalized": false,
+ "rstrip": false,
+ "single_word": false,
+ "special": true
+ }
+ },
+ "bos_token": "<s>",
+ "clean_up_tokenization_spaces": true,
+ "cls_token": "<s>",
+ "eos_token": "</s>",
+ "mask_token": "<mask>",
+ "model_max_length": 32768,
+ "pad_token": "<pad>",
+ "sep_token": "</s>",
+ "tokenizer_class": "PreTrainedTokenizerFast",
+ "unk_token": "<unk>"
+}