Upload 6 files

Browse files

Files changed (6) hide show

added_tokens.json +3 -0
config.json +50 -58
merges.txt +0 -0
special_tokens_map.json +24 -0
tokenizer_config.json +30 -0
vocab.json +0 -0

added_tokens.json ADDED Viewed

	@@ -0,0 +1,3 @@

+{
+  "<pad>": 50257
+}

config.json CHANGED Viewed

@@ -1,23 +1,23 @@
 {
-  "_name_or_path": "ayjays132/CustomGPT2Conversational",
-  "activation_function": "gelu_new",
-  "architectures": [
-    "GPT2LMHeadModel"
-  ],
-  "attn_pdrop": 0.1,
-  "bos_token_id": 50256,
   "config": {
     "activation_function": "gelu_new",
-    "attn_pdrop": 0.1,
-    "embd_pdrop": 0.1,
-    "initializer_range": 0.02,
-    "layer_norm_epsilon": 1e-05,
     "n_ctx": 2048,
     "n_embd": 2048,
     "n_head": 16,
     "n_layer": 24,
     "n_positions": 2048,
     "n_special": 0,
     "resid_pdrop": 0.1,
     "summary_activation": null,
     "summary_first_dropout": 0.1,
@@ -25,61 +25,53 @@
     "summary_type": "cls_index",
     "summary_use_proj": true
   },
-  "datasets": [
-    "vicgalle/alpaca-gpt4"
-  ],
-  "embd_pdrop": 0.1,
-  "eos_token_id": 50256,
-  "initializer_range": 0.02,
-  "language": [
-    "en"
-  ],
-  "layer_norm_epsilon": 1e-05,
-  "library_name": "transformers",
-  "license": "apache-2.0",
-  "metrics": [
-    "perplexity",
-    "accuracy"
-  ],
-  "model_type": "gpt2",
-  "n_embd": 768,
-  "n_head": 12,
-  "n_inner": null,
-  "n_layer": 12,
-  "n_positions": 1024,
-  "pipeline_tag": "conversational",
-  "reorder_and_upcast_attn": false,
-  "resid_pdrop": 0.1,
-  "scale_attn_by_inverse_layer_idx": false,
-  "scale_attn_weights": true,
-  "summary_activation": null,
-  "summary_first_dropout": 0.1,
-  "summary_proj_to_labels": true,
-  "summary_type": "cls_index",
-  "summary_use_proj": true,
-  "tags": [
-    "conversational"
-  ],
   "task_specific_params": {
     "conversational": {
-      "early_stopping": true,
-      "length_penalty": 1.5,
       "max_length": 1024,
       "min_length": 20,
-      "no_repeat_ngram_size": 3,
       "num_beams": 5,
       "temperature": 0.7,
       "top_k": 50,
       "top_p": 0.9
     }
   },
-  "tokenizer_config": {
-    "bos_token_id": 50256,
-    "eos_token_id": 50256,
-    "n_positions": 2048
-  },
-  "torch_dtype": "float32",
-  "transformers_version": "4.37.2",
-  "use_cache": true,
-  "vocab_size": 50257
 }

 {
+  "model_type": "gpt2",
+  "architectures": ["GPT2LMHeadModel"],
+  "tokenizer_config": {
+    "bos_token_id": 50256,
+    "eos_token_id": 50256,
+    "n_positions": 2048
+  },
   "config": {
     "activation_function": "gelu_new",
     "n_ctx": 2048,
     "n_embd": 2048,
     "n_head": 16,
     "n_layer": 24,
     "n_positions": 2048,
     "n_special": 0,
+    "attn_pdrop": 0.1,
+    "embd_pdrop": 0.1,
+    "initializer_range": 0.02,
+    "layer_norm_epsilon": 1e-05,
     "resid_pdrop": 0.1,
     "summary_activation": null,
     "summary_first_dropout": 0.1,
     "summary_type": "cls_index",
     "summary_use_proj": true
   },
   "task_specific_params": {
     "conversational": {
       "max_length": 1024,
       "min_length": 20,
+      "length_penalty": 1.5,
       "num_beams": 5,
+      "early_stopping": true,
+      "no_repeat_ngram_size": 3,
       "temperature": 0.7,
       "top_k": 50,
       "top_p": 0.9
     }
   },
+  "transformers_version": "4.34.0",
+  "language": ["en"],
+  "tags": ["conversational"],
+  "metrics": ["perplexity", "accuracy"],
+  "pipeline_tag": "conversational",
+  "library_name": "transformers",
+  "datasets": ["vicgalle/alpaca-gpt4"],
+  "license": "apache-2.0",
+  "custom_params": {
+    "adaptation_rate": 0.05,
+    "desired_improvement_rate": 0.02,
+    "ecosystem_dynamics": {
+      "environmental_volatility": 0.1,
+      "resource_pool": 1
+    },
+    "growth_improvement_threshold": 0.01,
+    "hidden_dim": 2048,
+    "initial_neuron_count": 5000,
+    "innovative_growth_net": {
+      "adaptation_rate": 0.05,
+      "initial_capacity": 250000,
+      "input_size": 2048
+    },
+    "input_dimension": 768,
+    "low_stability_threshold": 0.01,
+    "max_complexity": 10000,
+    "max_neurons": 250000,
+    "max_sequence_length": 1024,
+    "min_epochs_before_growth": 5,
+    "model_filename": "pytorch_model.bin",
+    "num_embeddings": 25000,
+    "pruning_improvement_threshold": 0.005,
+    "some_adaptation_rate": 0.05,
+    "stability_threshold": 0.02,
+    "start_token_index": 2
+  }
 }

merges.txt ADDED Viewed

The diff for this file is too large to render. See raw diff

special_tokens_map.json ADDED Viewed

	@@ -0,0 +1,24 @@

+{
+  "bos_token": {
+    "content": "<|endoftext|>",
+    "lstrip": false,
+    "normalized": true,
+    "rstrip": false,
+    "single_word": false
+  },
+  "eos_token": {
+    "content": "<|endoftext|>",
+    "lstrip": false,
+    "normalized": true,
+    "rstrip": false,
+    "single_word": false
+  },
+  "pad_token": "<pad>",
+  "unk_token": {
+    "content": "<|endoftext|>",
+    "lstrip": false,
+    "normalized": true,
+    "rstrip": false,
+    "single_word": false
+  }
+}

tokenizer_config.json ADDED Viewed

	@@ -0,0 +1,30 @@

+{
+  "add_bos_token": false,
+  "add_prefix_space": false,
+  "added_tokens_decoder": {
+    "50256": {
+      "content": "<|endoftext|>",
+      "lstrip": false,
+      "normalized": true,
+      "rstrip": false,
+      "single_word": false,
+      "special": true
+    },
+    "50257": {
+      "content": "<pad>",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": true
+    }
+  },
+  "bos_token": "<|endoftext|>",
+  "clean_up_tokenization_spaces": true,
+  "eos_token": "<|endoftext|>",
+  "errors": "replace",
+  "model_max_length": 2048,
+  "pad_token": "<pad>",
+  "tokenizer_class": "GPT2Tokenizer",
+  "unk_token": "<|endoftext|>"
+}

vocab.json ADDED Viewed

The diff for this file is too large to render. See raw diff