Buckets:
| import{s as po,o as ho,n as Ke}from"../chunks/scheduler.9991993c.js";import{S as go,i as uo,g as s,s as o,r as g,A as _o,h as i,f as c,c as n,j as C,u,x as p,k as $,y as t,a as w,v as _,d as b,t as y,w as v}from"../chunks/index.7fc9a5e7.js";import{T as co}from"../chunks/Tip.9de92fc6.js";import{D as T}from"../chunks/Docstring.0d7e3ebb.js";import{C as fo}from"../chunks/CodeBlock.e11cba92.js";import{E as lo}from"../chunks/ExampleCodeBlock.46b9776a.js";import{H as mo,E as bo}from"../chunks/EditOnGithub.84ab7f0e.js";function yo(j){let d,k=`A configuration file can be loaded and saved to disk. Loading the configuration file and using this file to | |
| initialize a model does <strong>not</strong> load the model weights. It only affects the model’s configuration.`;return{c(){d=s("p"),d.innerHTML=k},l(f){d=i(f,"P",{"data-svelte-h":!0}),p(d)!=="svelte-s3sff7"&&(d.innerHTML=k)},m(f,h){w(f,d,h)},p:Ke,d(f){f&&c(d)}}}function vo(j){let d,k="Examples:",f,h,x;return h=new fo({props:{code:"ZnJvbSUyMHRyYW5zZm9ybWVycyUyMGltcG9ydCUyMEF1dG9Db25maWclMEElMEFjb25maWclMjAlM0QlMjBBdXRvQ29uZmlnLmZyb21fcHJldHJhaW5lZCglMjJnb29nbGUtYmVydCUyRmJlcnQtYmFzZS1jYXNlZCUyMiklMEElMEElMjMlMjBQdXNoJTIwdGhlJTIwY29uZmlnJTIwdG8lMjB5b3VyJTIwbmFtZXNwYWNlJTIwd2l0aCUyMHRoZSUyMG5hbWUlMjAlMjJteS1maW5ldHVuZWQtYmVydCUyMi4lMEFjb25maWcucHVzaF90b19odWIoJTIybXktZmluZXR1bmVkLWJlcnQlMjIpJTBBJTBBJTIzJTIwUHVzaCUyMHRoZSUyMGNvbmZpZyUyMHRvJTIwYW4lMjBvcmdhbml6YXRpb24lMjB3aXRoJTIwdGhlJTIwbmFtZSUyMCUyMm15LWZpbmV0dW5lZC1iZXJ0JTIyLiUwQWNvbmZpZy5wdXNoX3RvX2h1YiglMjJodWdnaW5nZmFjZSUyRm15LWZpbmV0dW5lZC1iZXJ0JTIyKQ==",highlighted:`<span class="hljs-keyword">from</span> transformers <span class="hljs-keyword">import</span> AutoConfig | |
| config = AutoConfig.from_pretrained(<span class="hljs-string">"google-bert/bert-base-cased"</span>) | |
| <span class="hljs-comment"># Push the config to your namespace with the name "my-finetuned-bert".</span> | |
| config.push_to_hub(<span class="hljs-string">"my-finetuned-bert"</span>) | |
| <span class="hljs-comment"># Push the config to an organization with the name "my-finetuned-bert".</span> | |
| config.push_to_hub(<span class="hljs-string">"huggingface/my-finetuned-bert"</span>)`,wrap:!1}}),{c(){d=s("p"),d.textContent=k,f=o(),g(h.$$.fragment)},l(l){d=i(l,"P",{"data-svelte-h":!0}),p(d)!=="svelte-kvfsh7"&&(d.textContent=k),f=n(l),u(h.$$.fragment,l)},m(l,P){w(l,d,P),w(l,f,P),_(h,l,P),x=!0},p:Ke,i(l){x||(b(h.$$.fragment,l),x=!0)},o(l){y(h.$$.fragment,l),x=!1},d(l){l&&(c(d),c(f)),v(h,l)}}}function wo(j){let d,k="Examples:",f,h,x;return h=new fo({props:{code:"JTIzJTIwV2UlMjBjYW4ndCUyMGluc3RhbnRpYXRlJTIwZGlyZWN0bHklMjB0aGUlMjBiYXNlJTIwY2xhc3MlMjAqUHJldHJhaW5lZENvbmZpZyolMjBzbyUyMGxldCdzJTIwc2hvdyUyMHRoZSUyMGV4YW1wbGVzJTIwb24lMjBhJTBBJTIzJTIwZGVyaXZlZCUyMGNsYXNzJTNBJTIwQmVydENvbmZpZyUwQWNvbmZpZyUyMCUzRCUyMEJlcnRDb25maWcuZnJvbV9wcmV0cmFpbmVkKCUwQSUyMCUyMCUyMCUyMCUyMmdvb2dsZS1iZXJ0JTJGYmVydC1iYXNlLXVuY2FzZWQlMjIlMEEpJTIwJTIwJTIzJTIwRG93bmxvYWQlMjBjb25maWd1cmF0aW9uJTIwZnJvbSUyMGh1Z2dpbmdmYWNlLmNvJTIwYW5kJTIwY2FjaGUuJTBBY29uZmlnJTIwJTNEJTIwQmVydENvbmZpZy5mcm9tX3ByZXRyYWluZWQoJTBBJTIwJTIwJTIwJTIwJTIyLiUyRnRlc3QlMkZzYXZlZF9tb2RlbCUyRiUyMiUwQSklMjAlMjAlMjMlMjBFLmcuJTIwY29uZmlnJTIwKG9yJTIwbW9kZWwpJTIwd2FzJTIwc2F2ZWQlMjB1c2luZyUyMCpzYXZlX3ByZXRyYWluZWQoJy4lMkZ0ZXN0JTJGc2F2ZWRfbW9kZWwlMkYnKSolMEFjb25maWclMjAlM0QlMjBCZXJ0Q29uZmlnLmZyb21fcHJldHJhaW5lZCglMjIuJTJGdGVzdCUyRnNhdmVkX21vZGVsJTJGbXlfY29uZmlndXJhdGlvbi5qc29uJTIyKSUwQWNvbmZpZyUyMCUzRCUyMEJlcnRDb25maWcuZnJvbV9wcmV0cmFpbmVkKCUyMmdvb2dsZS1iZXJ0JTJGYmVydC1iYXNlLXVuY2FzZWQlMjIlMkMlMjBvdXRwdXRfYXR0ZW50aW9ucyUzRFRydWUlMkMlMjBmb28lM0RGYWxzZSklMEFhc3NlcnQlMjBjb25maWcub3V0cHV0X2F0dGVudGlvbnMlMjAlM0QlM0QlMjBUcnVlJTBBY29uZmlnJTJDJTIwdW51c2VkX2t3YXJncyUyMCUzRCUyMEJlcnRDb25maWcuZnJvbV9wcmV0cmFpbmVkKCUwQSUyMCUyMCUyMCUyMCUyMmdvb2dsZS1iZXJ0JTJGYmVydC1iYXNlLXVuY2FzZWQlMjIlMkMlMjBvdXRwdXRfYXR0ZW50aW9ucyUzRFRydWUlMkMlMjBmb28lM0RGYWxzZSUyQyUyMHJldHVybl91bnVzZWRfa3dhcmdzJTNEVHJ1ZSUwQSklMEFhc3NlcnQlMjBjb25maWcub3V0cHV0X2F0dGVudGlvbnMlMjAlM0QlM0QlMjBUcnVlJTBBYXNzZXJ0JTIwdW51c2VkX2t3YXJncyUyMCUzRCUzRCUyMCU3QiUyMmZvbyUyMiUzQSUyMEZhbHNlJTdE",highlighted:`<span class="hljs-comment"># We can't instantiate directly the base class *PretrainedConfig* so let's show the examples on a</span> | |
| <span class="hljs-comment"># derived class: BertConfig</span> | |
| config = BertConfig.from_pretrained( | |
| <span class="hljs-string">"google-bert/bert-base-uncased"</span> | |
| ) <span class="hljs-comment"># Download configuration from huggingface.co and cache.</span> | |
| config = BertConfig.from_pretrained( | |
| <span class="hljs-string">"./test/saved_model/"</span> | |
| ) <span class="hljs-comment"># E.g. config (or model) was saved using *save_pretrained('./test/saved_model/')*</span> | |
| config = BertConfig.from_pretrained(<span class="hljs-string">"./test/saved_model/my_configuration.json"</span>) | |
| config = BertConfig.from_pretrained(<span class="hljs-string">"google-bert/bert-base-uncased"</span>, output_attentions=<span class="hljs-literal">True</span>, foo=<span class="hljs-literal">False</span>) | |
| <span class="hljs-keyword">assert</span> config.output_attentions == <span class="hljs-literal">True</span> | |
| config, unused_kwargs = BertConfig.from_pretrained( | |
| <span class="hljs-string">"google-bert/bert-base-uncased"</span>, output_attentions=<span class="hljs-literal">True</span>, foo=<span class="hljs-literal">False</span>, return_unused_kwargs=<span class="hljs-literal">True</span> | |
| ) | |
| <span class="hljs-keyword">assert</span> config.output_attentions == <span class="hljs-literal">True</span> | |
| <span class="hljs-keyword">assert</span> unused_kwargs == {<span class="hljs-string">"foo"</span>: <span class="hljs-literal">False</span>}`,wrap:!1}}),{c(){d=s("p"),d.textContent=k,f=o(),g(h.$$.fragment)},l(l){d=i(l,"P",{"data-svelte-h":!0}),p(d)!=="svelte-kvfsh7"&&(d.textContent=k),f=n(l),u(h.$$.fragment,l)},m(l,P){w(l,d,P),w(l,f,P),_(h,l,P),x=!0},p:Ke,i(l){x||(b(h.$$.fragment,l),x=!0)},o(l){y(h.$$.fragment,l),x=!1},d(l){l&&(c(d),c(f)),v(h,l)}}}function xo(j){let d,k="This API is experimental and may have some slight breaking changes in the next releases.";return{c(){d=s("p"),d.textContent=k},l(f){d=i(f,"P",{"data-svelte-h":!0}),p(d)!=="svelte-15rpg4"&&(d.textContent=k)},m(f,h){w(f,d,h)},p:Ke,d(f){f&&c(d)}}}function Co(j){let d,k,f,h,x,l,P,qt='基类<a href="/docs/transformers/v4.44.2/zh/main_classes/configuration#transformers.PretrainedConfig">PretrainedConfig</a>实现了从本地文件或目录加载/保存配置的常见方法,或下载库提供的预训练模型配置(从HuggingFace的AWS S3库中下载)。',qe,A,Vt="每个派生的配置类都实现了特定于模型的属性。所有配置类中共同存在的属性有:<code>hidden_size</code>、<code>num_attention_heads</code> 和 <code>num_hidden_layers</code>。文本模型进一步添加了 <code>vocab_size</code>。",Ve,G,He,r,Q,et,ge,Ht=`Base class for all configuration classes. Handles a few parameters common to all models’ configurations as well as | |
| methods for loading/downloading/saving configurations.`,tt,Z,ot,ue,Nt="Class attributes (overridden by derived classes):",nt,_e,Bt=`<li><strong>model_type</strong> (<code>str</code>) — An identifier for the model type, serialized into the JSON file, and used to recreate | |
| the correct object in <code>AutoConfig</code>.</li> <li><strong>is_composition</strong> (<code>bool</code>) — Whether the config class is composed of multiple sub-configs. In this case the | |
| config has to be initialized from two or more configs of type <a href="/docs/transformers/v4.44.2/zh/main_classes/configuration#transformers.PretrainedConfig">PretrainedConfig</a> like: | |
| <code>EncoderDecoderConfig</code> or <code>~RagConfig</code>.</li> <li><strong>keys_to_ignore_at_inference</strong> (<code>List[str]</code>) — A list of keys to ignore by default when looking at dictionary | |
| outputs of the model during inference.</li> <li><strong>attribute_map</strong> (<code>Dict[str, str]</code>) — A dict that maps model specific attribute names to the standardized | |
| naming of attributes.</li>`,rt,be,Et="Common attributes (present in all subclasses):",at,ye,St=`<li><strong>vocab_size</strong> (<code>int</code>) — The number of tokens in the vocabulary, which is also the first dimension of the | |
| embeddings matrix (this attribute may be missing for models that don’t have a text modality like ViT).</li> <li><strong>hidden_size</strong> (<code>int</code>) — The hidden size of the model.</li> <li><strong>num_attention_heads</strong> (<code>int</code>) — The number of attention heads used in the multi-head attention layers of the | |
| model.</li> <li><strong>num_hidden_layers</strong> (<code>int</code>) — The number of blocks in the model.</li>`,st,z,O,it,ve,Xt="Upload the configuration file to the 🤗 Model Hub.",dt,I,ct,L,K,lt,we,Rt=`Checks whether the passed dictionary and its nested dicts have a <em>torch_dtype</em> key and if it’s not None, | |
| converts torch.dtype to a string of just the type. For example, <code>torch.float32</code> get converted into <em>“float32”</em> | |
| string, which can then be stored in the json format.`,mt,F,ee,ft,xe,Yt='Instantiates a <a href="/docs/transformers/v4.44.2/zh/main_classes/configuration#transformers.PretrainedConfig">PretrainedConfig</a> from a Python dictionary of parameters.',pt,D,te,ht,Ce,At='Instantiates a <a href="/docs/transformers/v4.44.2/zh/main_classes/configuration#transformers.PretrainedConfig">PretrainedConfig</a> from the path to a JSON file of parameters.',gt,J,oe,ut,$e,Gt='Instantiate a <a href="/docs/transformers/v4.44.2/zh/main_classes/configuration#transformers.PretrainedConfig">PretrainedConfig</a> (or a derived class) from a pretrained model configuration.',_t,q,bt,V,ne,yt,ke,Qt=`From a <code>pretrained_model_name_or_path</code>, resolve to a dictionary of parameters, to be used for instantiating a | |
| <a href="/docs/transformers/v4.44.2/zh/main_classes/configuration#transformers.PretrainedConfig">PretrainedConfig</a> using <code>from_dict</code>.`,vt,U,re,wt,Pe,Ot=`Register this class with a given auto class. This should only be used for custom configurations as the ones in | |
| the library are already mapped with <code>AutoConfig</code>.`,xt,H,Ct,N,ae,$t,Te,Kt=`Save a configuration object to the directory <code>save_directory</code>, so that it can be re-loaded using the | |
| <a href="/docs/transformers/v4.44.2/zh/main_classes/configuration#transformers.PretrainedConfig.from_pretrained">from_pretrained()</a> class method.`,kt,B,se,Pt,Me,eo="Serializes this instance to a Python dictionary.",Tt,E,ie,Mt,je,to=`Removes all attributes from config which correspond to the default config attributes for better readability and | |
| serializes to a Python dictionary.`,jt,S,de,zt,ze,oo="Save this instance to a JSON file.",Jt,X,ce,Ut,Je,no="Serializes this instance to a JSON string.",Wt,R,le,Zt,Ue,ro="Updates attributes of this class with attributes from <code>config_dict</code>.",It,M,me,Lt,We,ao="Updates attributes of this class with attributes from <code>update_str</code>.",Ft,Ze,so=`The expected format is ints, floats and strings as is, and for booleans use <code>true</code> or <code>false</code>. For example: | |
| “n_embd=10,resid_pdrop=0.2,scale_attn_weights=false,summary_type=cls_index”`,Dt,Ie,io="The keys to change have to already exist in the config object.",Ne,fe,Be,De,Ee;return x=new mo({props:{title:"Configuration",local:"configuration",headingTag:"h1"}}),G=new mo({props:{title:"PretrainedConfig",local:"transformers.PretrainedConfig",headingTag:"h2"}}),Q=new T({props:{name:"class transformers.PretrainedConfig",anchor:"transformers.PretrainedConfig",parameters:[{name:"**kwargs",val:""}],parametersDescription:[{anchor:"transformers.PretrainedConfig.name_or_path",description:`<strong>name_or_path</strong> (<code>str</code>, <em>optional</em>, defaults to <code>""</code>) — | |
| Store the string that was passed to <a href="/docs/transformers/v4.44.2/zh/main_classes/model#transformers.PreTrainedModel.from_pretrained">PreTrainedModel.from_pretrained()</a> or | |
| <a href="/docs/transformers/v4.44.2/zh/main_classes/model#transformers.TFPreTrainedModel.from_pretrained">TFPreTrainedModel.from_pretrained()</a> as <code>pretrained_model_name_or_path</code> if the configuration was created | |
| with such a method.`,name:"name_or_path"},{anchor:"transformers.PretrainedConfig.output_hidden_states",description:`<strong>output_hidden_states</strong> (<code>bool</code>, <em>optional</em>, defaults to <code>False</code>) — | |
| Whether or not the model should return all hidden-states.`,name:"output_hidden_states"},{anchor:"transformers.PretrainedConfig.output_attentions",description:`<strong>output_attentions</strong> (<code>bool</code>, <em>optional</em>, defaults to <code>False</code>) — | |
| Whether or not the model should returns all attentions.`,name:"output_attentions"},{anchor:"transformers.PretrainedConfig.return_dict",description:`<strong>return_dict</strong> (<code>bool</code>, <em>optional</em>, defaults to <code>True</code>) — | |
| Whether or not the model should return a <a href="/docs/transformers/v4.44.2/zh/main_classes/output#transformers.utils.ModelOutput">ModelOutput</a> instead of a plain tuple.`,name:"return_dict"},{anchor:"transformers.PretrainedConfig.is_encoder_decoder",description:`<strong>is_encoder_decoder</strong> (<code>bool</code>, <em>optional</em>, defaults to <code>False</code>) — | |
| Whether the model is used as an encoder/decoder or not.`,name:"is_encoder_decoder"},{anchor:"transformers.PretrainedConfig.is_decoder",description:`<strong>is_decoder</strong> (<code>bool</code>, <em>optional</em>, defaults to <code>False</code>) — | |
| Whether the model is used as decoder or not (in which case it’s used as an encoder).`,name:"is_decoder"},{anchor:"transformers.PretrainedConfig.cross_attention_hidden_size**",description:`<strong>cross_attention_hidden_size**</strong> (<code>bool</code>, <em>optional</em>) — | |
| The hidden size of the cross-attention layer in case the model is used as a decoder in an encoder-decoder | |
| setting and the cross-attention hidden dimension differs from <code>self.config.hidden_size</code>.`,name:"cross_attention_hidden_size**"},{anchor:"transformers.PretrainedConfig.add_cross_attention",description:`<strong>add_cross_attention</strong> (<code>bool</code>, <em>optional</em>, defaults to <code>False</code>) — | |
| Whether cross-attention layers should be added to the model. Note, this option is only relevant for models | |
| that can be used as decoder models within the <code>EncoderDecoderModel</code> class, which consists of all models | |
| in <code>AUTO_MODELS_FOR_CAUSAL_LM</code>.`,name:"add_cross_attention"},{anchor:"transformers.PretrainedConfig.tie_encoder_decoder",description:`<strong>tie_encoder_decoder</strong> (<code>bool</code>, <em>optional</em>, defaults to <code>False</code>) — | |
| Whether all encoder weights should be tied to their equivalent decoder weights. This requires the encoder | |
| and decoder model to have the exact same parameter names.`,name:"tie_encoder_decoder"},{anchor:"transformers.PretrainedConfig.prune_heads",description:`<strong>prune_heads</strong> (<code>Dict[int, List[int]]</code>, <em>optional</em>, defaults to <code>{}</code>) — | |
| Pruned heads of the model. The keys are the selected layer indices and the associated values, the list of | |
| heads to prune in said layer.</p> | |
| <p>For instance <code>{1: [0, 2], 2: [2, 3]}</code> will prune heads 0 and 2 on layer 1 and heads 2 and 3 on layer 2.`,name:"prune_heads"},{anchor:"transformers.PretrainedConfig.chunk_size_feed_forward",description:`<strong>chunk_size_feed_forward</strong> (<code>int</code>, <em>optional</em>, defaults to <code>0</code>) — | |
| The chunk size of all feed forward layers in the residual attention blocks. A chunk size of <code>0</code> means that | |
| the feed forward layer is not chunked. A chunk size of n means that the feed forward layer processes <code>n</code> < | |
| sequence_length embeddings at a time. For more information on feed forward chunking, see <a href="../glossary.html#feed-forward-chunking">How does Feed | |
| Forward Chunking work?</a>.`,name:"chunk_size_feed_forward"}],source:"https://github.com/huggingface/transformers/blob/v4.44.2/src/transformers/configuration_utils.py#L50",parameterGroups:[{title:"Parameters for sequence generation",parametersDescription:[{anchor:"transformers.PretrainedConfig.max_length",description:`<strong>max_length</strong> (<code>int</code>, <em>optional</em>, defaults to 20) — | |
| Maximum length that will be used by default in the <code>generate</code> method of the model.`,name:"max_length"},{anchor:"transformers.PretrainedConfig.min_length",description:`<strong>min_length</strong> (<code>int</code>, <em>optional</em>, defaults to 0) — | |
| Minimum length that will be used by default in the <code>generate</code> method of the model.`,name:"min_length"},{anchor:"transformers.PretrainedConfig.do_sample",description:`<strong>do_sample</strong> (<code>bool</code>, <em>optional</em>, defaults to <code>False</code>) — | |
| Flag that will be used by default in the <code>generate</code> method of the model. Whether or not to use sampling ; | |
| use greedy decoding otherwise.`,name:"do_sample"},{anchor:"transformers.PretrainedConfig.early_stopping",description:`<strong>early_stopping</strong> (<code>bool</code>, <em>optional</em>, defaults to <code>False</code>) — | |
| Flag that will be used by default in the <code>generate</code> method of the model. Whether to stop the beam search | |
| when at least <code>num_beams</code> sentences are finished per batch or not.`,name:"early_stopping"},{anchor:"transformers.PretrainedConfig.num_beams",description:`<strong>num_beams</strong> (<code>int</code>, <em>optional</em>, defaults to 1) — | |
| Number of beams for beam search that will be used by default in the <code>generate</code> method of the model. 1 means | |
| no beam search.`,name:"num_beams"},{anchor:"transformers.PretrainedConfig.num_beam_groups",description:`<strong>num_beam_groups</strong> (<code>int</code>, <em>optional</em>, defaults to 1) — | |
| Number of groups to divide <code>num_beams</code> into in order to ensure diversity among different groups of beams | |
| that will be used by default in the <code>generate</code> method of the model. 1 means no group beam search.`,name:"num_beam_groups"},{anchor:"transformers.PretrainedConfig.diversity_penalty",description:`<strong>diversity_penalty</strong> (<code>float</code>, <em>optional</em>, defaults to 0.0) — | |
| Value to control diversity for group beam search. that will be used by default in the <code>generate</code> method of | |
| the model. 0 means no diversity penalty. The higher the penalty, the more diverse are the outputs.`,name:"diversity_penalty"},{anchor:"transformers.PretrainedConfig.temperature",description:`<strong>temperature</strong> (<code>float</code>, <em>optional</em>, defaults to 1.0) — | |
| The value used to module the next token probabilities that will be used by default in the <code>generate</code> method | |
| of the model. Must be strictly positive.`,name:"temperature"},{anchor:"transformers.PretrainedConfig.top_k",description:`<strong>top_k</strong> (<code>int</code>, <em>optional</em>, defaults to 50) — | |
| Number of highest probability vocabulary tokens to keep for top-k-filtering that will be used by default in | |
| the <code>generate</code> method of the model.`,name:"top_k"},{anchor:"transformers.PretrainedConfig.top_p",description:`<strong>top_p</strong> (<code>float</code>, <em>optional</em>, defaults to 1) — | |
| Value that will be used by default in the <code>generate</code> method of the model for <code>top_p</code>. If set to float < 1, | |
| only the most probable tokens with probabilities that add up to <code>top_p</code> or higher are kept for generation.`,name:"top_p"},{anchor:"transformers.PretrainedConfig.typical_p",description:`<strong>typical_p</strong> (<code>float</code>, <em>optional</em>, defaults to 1) — | |
| Local typicality measures how similar the conditional probability of predicting a target token next is to | |
| the expected conditional probability of predicting a random token next, given the partial text already | |
| generated. If set to float < 1, the smallest set of the most locally typical tokens with probabilities that | |
| add up to <code>typical_p</code> or higher are kept for generation. See <a href="https://arxiv.org/pdf/2202.00666.pdf" rel="nofollow">this | |
| paper</a> for more details.`,name:"typical_p"},{anchor:"transformers.PretrainedConfig.repetition_penalty",description:`<strong>repetition_penalty</strong> (<code>float</code>, <em>optional</em>, defaults to 1) — | |
| Parameter for repetition penalty that will be used by default in the <code>generate</code> method of the model. 1.0 | |
| means no penalty.`,name:"repetition_penalty"},{anchor:"transformers.PretrainedConfig.length_penalty",description:`<strong>length_penalty</strong> (<code>float</code>, <em>optional</em>, defaults to 1) — | |
| Exponential penalty to the length that is used with beam-based generation. It is applied as an exponent to | |
| the sequence length, which in turn is used to divide the score of the sequence. Since the score is the log | |
| likelihood of the sequence (i.e. negative), <code>length_penalty</code> > 0.0 promotes longer sequences, while | |
| <code>length_penalty</code> < 0.0 encourages shorter sequences.`,name:"length_penalty"},{anchor:"transformers.PretrainedConfig.no_repeat_ngram_size",description:`<strong>no_repeat_ngram_size</strong> (<code>int</code>, <em>optional</em>, defaults to 0) — Value that will be used by default in the — | |
| <code>generate</code> method of the model for <code>no_repeat_ngram_size</code>. If set to int > 0, all ngrams of that size can | |
| only occur once.`,name:"no_repeat_ngram_size"},{anchor:"transformers.PretrainedConfig.encoder_no_repeat_ngram_size",description:`<strong>encoder_no_repeat_ngram_size</strong> (<code>int</code>, <em>optional</em>, defaults to 0) — Value that will be used by — | |
| default in the <code>generate</code> method of the model for <code>encoder_no_repeat_ngram_size</code>. If set to int > 0, all | |
| ngrams of that size that occur in the <code>encoder_input_ids</code> cannot occur in the <code>decoder_input_ids</code>.`,name:"encoder_no_repeat_ngram_size"},{anchor:"transformers.PretrainedConfig.bad_words_ids",description:`<strong>bad_words_ids</strong> (<code>List[int]</code>, <em>optional</em>) — | |
| List of token ids that are not allowed to be generated that will be used by default in the <code>generate</code> | |
| method of the model. In order to get the tokens of the words that should not appear in the generated text, | |
| use <code>tokenizer.encode(bad_word, add_prefix_space=True)</code>.`,name:"bad_words_ids"},{anchor:"transformers.PretrainedConfig.num_return_sequences",description:`<strong>num_return_sequences</strong> (<code>int</code>, <em>optional</em>, defaults to 1) — | |
| Number of independently computed returned sequences for each element in the batch that will be used by | |
| default in the <code>generate</code> method of the model.`,name:"num_return_sequences"},{anchor:"transformers.PretrainedConfig.output_scores",description:`<strong>output_scores</strong> (<code>bool</code>, <em>optional</em>, defaults to <code>False</code>) — | |
| Whether the model should return the logits when used for generation.`,name:"output_scores"},{anchor:"transformers.PretrainedConfig.return_dict_in_generate",description:`<strong>return_dict_in_generate</strong> (<code>bool</code>, <em>optional</em>, defaults to <code>False</code>) — | |
| Whether the model should return a <a href="/docs/transformers/v4.44.2/zh/main_classes/output#transformers.utils.ModelOutput">ModelOutput</a> instead of a <code>torch.LongTensor</code>.`,name:"return_dict_in_generate"},{anchor:"transformers.PretrainedConfig.forced_bos_token_id",description:`<strong>forced_bos_token_id</strong> (<code>int</code>, <em>optional</em>) — | |
| The id of the token to force as the first generated token after the <code>decoder_start_token_id</code>. Useful for | |
| multilingual models like <a href="../model_doc/mbart">mBART</a> where the first generated token needs to be the target | |
| language token.`,name:"forced_bos_token_id"},{anchor:"transformers.PretrainedConfig.forced_eos_token_id",description:`<strong>forced_eos_token_id</strong> (<code>int</code>, <em>optional</em>) — | |
| The id of the token to force as the last generated token when <code>max_length</code> is reached.`,name:"forced_eos_token_id"},{anchor:"transformers.PretrainedConfig.remove_invalid_values",description:`<strong>remove_invalid_values</strong> (<code>bool</code>, <em>optional</em>) — | |
| Whether to remove possible <em>nan</em> and <em>inf</em> outputs of the model to prevent the generation method to crash. | |
| Note that using <code>remove_invalid_values</code> can slow down generation.`,name:"remove_invalid_values"}]},{title:"Parameters for fine-tuning tasks",parametersDescription:[{anchor:"transformers.PretrainedConfig.architectures",description:`<strong>architectures</strong> (<code>List[str]</code>, <em>optional</em>) — | |
| Model architectures that can be used with the model pretrained weights.`,name:"architectures"},{anchor:"transformers.PretrainedConfig.finetuning_task",description:`<strong>finetuning_task</strong> (<code>str</code>, <em>optional</em>) — | |
| Name of the task used to fine-tune the model. This can be used when converting from an original (TensorFlow | |
| or PyTorch) checkpoint.`,name:"finetuning_task"},{anchor:"transformers.PretrainedConfig.id2label",description:`<strong>id2label</strong> (<code>Dict[int, str]</code>, <em>optional</em>) — | |
| A map from index (for instance prediction index, or target index) to label.`,name:"id2label"},{anchor:"transformers.PretrainedConfig.label2id",description:"<strong>label2id</strong> (<code>Dict[str, int]</code>, <em>optional</em>) — A map from label to index for the model.",name:"label2id"},{anchor:"transformers.PretrainedConfig.num_labels",description:`<strong>num_labels</strong> (<code>int</code>, <em>optional</em>) — | |
| Number of labels to use in the last layer added to the model, typically for a classification task.`,name:"num_labels"},{anchor:"transformers.PretrainedConfig.task_specific_params",description:`<strong>task_specific_params</strong> (<code>Dict[str, Any]</code>, <em>optional</em>) — | |
| Additional keyword arguments to store for the current task.`,name:"task_specific_params"},{anchor:"transformers.PretrainedConfig.problem_type",description:`<strong>problem_type</strong> (<code>str</code>, <em>optional</em>) — | |
| Problem type for <code>XxxForSequenceClassification</code> models. Can be one of <code>"regression"</code>, | |
| <code>"single_label_classification"</code> or <code>"multi_label_classification"</code>.`,name:"problem_type"}]},{title:"Parameters linked to the tokenizer",parametersDescription:[{anchor:"transformers.PretrainedConfig.tokenizer_class",description:`<strong>tokenizer_class</strong> (<code>str</code>, <em>optional</em>) — | |
| The name of the associated tokenizer class to use (if none is set, will use the tokenizer associated to the | |
| model by default).`,name:"tokenizer_class"},{anchor:"transformers.PretrainedConfig.prefix",description:`<strong>prefix</strong> (<code>str</code>, <em>optional</em>) — | |
| A specific prompt that should be added at the beginning of each text before calling the model.`,name:"prefix"},{anchor:"transformers.PretrainedConfig.bos_token_id",description:"<strong>bos_token_id</strong> (<code>int</code>, <em>optional</em>) — The id of the <em>beginning-of-stream</em> token.",name:"bos_token_id"},{anchor:"transformers.PretrainedConfig.pad_token_id",description:"<strong>pad_token_id</strong> (<code>int</code>, <em>optional</em>) — The id of the <em>padding</em> token.",name:"pad_token_id"},{anchor:"transformers.PretrainedConfig.eos_token_id",description:"<strong>eos_token_id</strong> (<code>int</code>, <em>optional</em>) — The id of the <em>end-of-stream</em> token.",name:"eos_token_id"},{anchor:"transformers.PretrainedConfig.decoder_start_token_id",description:`<strong>decoder_start_token_id</strong> (<code>int</code>, <em>optional</em>) — | |
| If an encoder-decoder model starts decoding with a different token than <em>bos</em>, the id of that token.`,name:"decoder_start_token_id"},{anchor:"transformers.PretrainedConfig.sep_token_id",description:"<strong>sep_token_id</strong> (<code>int</code>, <em>optional</em>) — The id of the <em>separation</em> token.",name:"sep_token_id"}]},{title:"PyTorch specific parameters",parametersDescription:[{anchor:"transformers.PretrainedConfig.torchscript",description:`<strong>torchscript</strong> (<code>bool</code>, <em>optional</em>, defaults to <code>False</code>) — | |
| Whether or not the model should be used with Torchscript.`,name:"torchscript"},{anchor:"transformers.PretrainedConfig.tie_word_embeddings",description:`<strong>tie_word_embeddings</strong> (<code>bool</code>, <em>optional</em>, defaults to <code>True</code>) — | |
| Whether the model’s input and output word embeddings should be tied. Note that this is only relevant if the | |
| model has a output word embedding layer.`,name:"tie_word_embeddings"},{anchor:"transformers.PretrainedConfig.torch_dtype",description:`<strong>torch_dtype</strong> (<code>str</code>, <em>optional</em>) — | |
| The <code>dtype</code> of the weights. This attribute can be used to initialize the model to a non-default <code>dtype</code> | |
| (which is normally <code>float32</code>) and thus allow for optimal storage allocation. For example, if the saved | |
| model is <code>float16</code>, ideally we want to load it back using the minimal amount of memory needed to load | |
| <code>float16</code> weights. Since the config object is stored in plain text, this attribute contains just the | |
| floating type string without the <code>torch.</code> prefix. For example, for <code>torch.float16</code> \`<code>torch_dtype</code> is the | |
| <code>"float16"</code> string.</p> | |
| <p>This attribute is currently not being used during model loading time, but this may change in the future | |
| versions. But we can already start preparing for the future by saving the dtype with save_pretrained.`,name:"torch_dtype"}]},{title:"TensorFlow specific parameters",parametersDescription:[{anchor:"transformers.PretrainedConfig.use_bfloat16",description:`<strong>use_bfloat16</strong> (<code>bool</code>, <em>optional</em>, defaults to <code>False</code>) — | |
| Whether or not the model should use BFloat16 scalars (only used by some TensorFlow models).`,name:"use_bfloat16"},{anchor:"transformers.PretrainedConfig.tf_legacy_loss",description:`<strong>tf_legacy_loss</strong> (<code>bool</code>, <em>optional</em>, defaults to <code>False</code>) — | |
| Whether the model should use legacy TensorFlow losses. Legacy losses have variable output shapes and may | |
| not be XLA-compatible. This option is here for backward compatibility and will be removed in Transformers | |
| v5.`,name:"tf_legacy_loss"}]}]}}),Z=new co({props:{$$slots:{default:[yo]},$$scope:{ctx:j}}}),O=new T({props:{name:"push_to_hub",anchor:"transformers.PretrainedConfig.push_to_hub",parameters:[{name:"repo_id",val:": str"},{name:"use_temp_dir",val:": Optional = None"},{name:"commit_message",val:": Optional = None"},{name:"private",val:": Optional = None"},{name:"token",val:": Union = None"},{name:"max_shard_size",val:": Union = '5GB'"},{name:"create_pr",val:": bool = False"},{name:"safe_serialization",val:": bool = True"},{name:"revision",val:": str = None"},{name:"commit_description",val:": str = None"},{name:"tags",val:": Optional = None"},{name:"**deprecated_kwargs",val:""}],parametersDescription:[{anchor:"transformers.PretrainedConfig.push_to_hub.repo_id",description:`<strong>repo_id</strong> (<code>str</code>) — | |
| The name of the repository you want to push your config to. It should contain your organization name | |
| when pushing to a given organization.`,name:"repo_id"},{anchor:"transformers.PretrainedConfig.push_to_hub.use_temp_dir",description:`<strong>use_temp_dir</strong> (<code>bool</code>, <em>optional</em>) — | |
| Whether or not to use a temporary directory to store the files saved before they are pushed to the Hub. | |
| Will default to <code>True</code> if there is no directory named like <code>repo_id</code>, <code>False</code> otherwise.`,name:"use_temp_dir"},{anchor:"transformers.PretrainedConfig.push_to_hub.commit_message",description:`<strong>commit_message</strong> (<code>str</code>, <em>optional</em>) — | |
| Message to commit while pushing. Will default to <code>"Upload config"</code>.`,name:"commit_message"},{anchor:"transformers.PretrainedConfig.push_to_hub.private",description:`<strong>private</strong> (<code>bool</code>, <em>optional</em>) — | |
| Whether or not the repository created should be private.`,name:"private"},{anchor:"transformers.PretrainedConfig.push_to_hub.token",description:`<strong>token</strong> (<code>bool</code> or <code>str</code>, <em>optional</em>) — | |
| The token to use as HTTP bearer authorization for remote files. If <code>True</code>, will use the token generated | |
| when running <code>huggingface-cli login</code> (stored in <code>~/.huggingface</code>). Will default to <code>True</code> if <code>repo_url</code> | |
| is not specified.`,name:"token"},{anchor:"transformers.PretrainedConfig.push_to_hub.max_shard_size",description:`<strong>max_shard_size</strong> (<code>int</code> or <code>str</code>, <em>optional</em>, defaults to <code>"5GB"</code>) — | |
| Only applicable for models. The maximum size for a checkpoint before being sharded. Checkpoints shard | |
| will then be each of size lower than this size. If expressed as a string, needs to be digits followed | |
| by a unit (like <code>"5MB"</code>). We default it to <code>"5GB"</code> so that users can easily load models on free-tier | |
| Google Colab instances without any CPU OOM issues.`,name:"max_shard_size"},{anchor:"transformers.PretrainedConfig.push_to_hub.create_pr",description:`<strong>create_pr</strong> (<code>bool</code>, <em>optional</em>, defaults to <code>False</code>) — | |
| Whether or not to create a PR with the uploaded files or directly commit.`,name:"create_pr"},{anchor:"transformers.PretrainedConfig.push_to_hub.safe_serialization",description:`<strong>safe_serialization</strong> (<code>bool</code>, <em>optional</em>, defaults to <code>True</code>) — | |
| Whether or not to convert the model weights in safetensors format for safer serialization.`,name:"safe_serialization"},{anchor:"transformers.PretrainedConfig.push_to_hub.revision",description:`<strong>revision</strong> (<code>str</code>, <em>optional</em>) — | |
| Branch to push the uploaded files to.`,name:"revision"},{anchor:"transformers.PretrainedConfig.push_to_hub.commit_description",description:`<strong>commit_description</strong> (<code>str</code>, <em>optional</em>) — | |
| The description of the commit that will be created`,name:"commit_description"},{anchor:"transformers.PretrainedConfig.push_to_hub.tags",description:`<strong>tags</strong> (<code>List[str]</code>, <em>optional</em>) — | |
| List of tags to push on the Hub.`,name:"tags"}],source:"https://github.com/huggingface/transformers/blob/v4.44.2/src/transformers/utils/hub.py#L809"}}),I=new lo({props:{anchor:"transformers.PretrainedConfig.push_to_hub.example",$$slots:{default:[vo]},$$scope:{ctx:j}}}),K=new T({props:{name:"dict_torch_dtype_to_str",anchor:"transformers.PretrainedConfig.dict_torch_dtype_to_str",parameters:[{name:"d",val:": Dict"}],source:"https://github.com/huggingface/transformers/blob/v4.44.2/src/transformers/configuration_utils.py#L1013"}}),ee=new T({props:{name:"from_dict",anchor:"transformers.PretrainedConfig.from_dict",parameters:[{name:"config_dict",val:": Dict"},{name:"**kwargs",val:""}],parametersDescription:[{anchor:"transformers.PretrainedConfig.from_dict.config_dict",description:`<strong>config_dict</strong> (<code>Dict[str, Any]</code>) — | |
| Dictionary that will be used to instantiate the configuration object. Such a dictionary can be | |
| retrieved from a pretrained checkpoint by leveraging the <a href="/docs/transformers/v4.44.2/zh/main_classes/configuration#transformers.PretrainedConfig.get_config_dict">get_config_dict()</a> method.`,name:"config_dict"},{anchor:"transformers.PretrainedConfig.from_dict.kwargs",description:`<strong>kwargs</strong> (<code>Dict[str, Any]</code>) — | |
| Additional parameters from which to initialize the configuration object.`,name:"kwargs"}],source:"https://github.com/huggingface/transformers/blob/v4.44.2/src/transformers/configuration_utils.py#L745",returnDescription:`<script context="module">export const metadata = 'undefined';<\/script> | |
| <p>The configuration object instantiated from those parameters.</p> | |
| `,returnType:`<script context="module">export const metadata = 'undefined';<\/script> | |
| <p><a | |
| href="/docs/transformers/v4.44.2/zh/main_classes/configuration#transformers.PretrainedConfig" | |
| >PretrainedConfig</a></p> | |
| `}}),te=new T({props:{name:"from_json_file",anchor:"transformers.PretrainedConfig.from_json_file",parameters:[{name:"json_file",val:": Union"}],parametersDescription:[{anchor:"transformers.PretrainedConfig.from_json_file.json_file",description:`<strong>json_file</strong> (<code>str</code> or <code>os.PathLike</code>) — | |
| Path to the JSON file containing the parameters.`,name:"json_file"}],source:"https://github.com/huggingface/transformers/blob/v4.44.2/src/transformers/configuration_utils.py#L806",returnDescription:`<script context="module">export const metadata = 'undefined';<\/script> | |
| <p>The configuration object instantiated from that JSON file.</p> | |
| `,returnType:`<script context="module">export const metadata = 'undefined';<\/script> | |
| <p><a | |
| href="/docs/transformers/v4.44.2/zh/main_classes/configuration#transformers.PretrainedConfig" | |
| >PretrainedConfig</a></p> | |
| `}}),oe=new T({props:{name:"from_pretrained",anchor:"transformers.PretrainedConfig.from_pretrained",parameters:[{name:"pretrained_model_name_or_path",val:": Union"},{name:"cache_dir",val:": Union = None"},{name:"force_download",val:": bool = False"},{name:"local_files_only",val:": bool = False"},{name:"token",val:": Union = None"},{name:"revision",val:": str = 'main'"},{name:"**kwargs",val:""}],parametersDescription:[{anchor:"transformers.PretrainedConfig.from_pretrained.pretrained_model_name_or_path",description:`<strong>pretrained_model_name_or_path</strong> (<code>str</code> or <code>os.PathLike</code>) — | |
| This can be either:</p> | |
| <ul> | |
| <li>a string, the <em>model id</em> of a pretrained model configuration hosted inside a model repo on | |
| huggingface.co.</li> | |
| <li>a path to a <em>directory</em> containing a configuration file saved using the | |
| <a href="/docs/transformers/v4.44.2/zh/main_classes/configuration#transformers.PretrainedConfig.save_pretrained">save_pretrained()</a> method, e.g., <code>./my_model_directory/</code>.</li> | |
| <li>a path or url to a saved configuration JSON <em>file</em>, e.g., <code>./my_model_directory/configuration.json</code>.</li> | |
| </ul>`,name:"pretrained_model_name_or_path"},{anchor:"transformers.PretrainedConfig.from_pretrained.cache_dir",description:`<strong>cache_dir</strong> (<code>str</code> or <code>os.PathLike</code>, <em>optional</em>) — | |
| Path to a directory in which a downloaded pretrained model configuration should be cached if the | |
| standard cache should not be used.`,name:"cache_dir"},{anchor:"transformers.PretrainedConfig.from_pretrained.force_download",description:`<strong>force_download</strong> (<code>bool</code>, <em>optional</em>, defaults to <code>False</code>) — | |
| Whether or not to force to (re-)download the configuration files and override the cached versions if | |
| they exist. | |
| resume_download — | |
| Deprecated and ignored. All downloads are now resumed by default when possible. | |
| Will be removed in v5 of Transformers.`,name:"force_download"},{anchor:"transformers.PretrainedConfig.from_pretrained.proxies",description:`<strong>proxies</strong> (<code>Dict[str, str]</code>, <em>optional</em>) — | |
| A dictionary of proxy servers to use by protocol or endpoint, e.g., <code>{'http': 'foo.bar:3128', 'http://hostname': 'foo.bar:4012'}.</code> The proxies are used on each request.`,name:"proxies"},{anchor:"transformers.PretrainedConfig.from_pretrained.token",description:`<strong>token</strong> (<code>str</code> or <code>bool</code>, <em>optional</em>) — | |
| The token to use as HTTP bearer authorization for remote files. If <code>True</code>, or not specified, will use | |
| the token generated when running <code>huggingface-cli login</code> (stored in <code>~/.huggingface</code>).`,name:"token"},{anchor:"transformers.PretrainedConfig.from_pretrained.revision",description:`<strong>revision</strong> (<code>str</code>, <em>optional</em>, defaults to <code>"main"</code>) — | |
| The specific model version to use. It can be a branch name, a tag name, or a commit id, since we use a | |
| git-based system for storing models and other artifacts on huggingface.co, so <code>revision</code> can be any | |
| identifier allowed by git.</p> | |
| <div class="course-tip bg-gradient-to-br dark:bg-gradient-to-r before:border-green-500 dark:before:border-green-800 from-green-50 dark:from-gray-900 to-white dark:to-gray-950 border border-green-50 text-green-700 dark:text-gray-400"> | |
| <p>To test a pull request you made on the Hub, you can pass \`revision=“refs/pr/<pr_number>“.</pr_number></p> | |
| </div>`,name:"revision"},{anchor:"transformers.PretrainedConfig.from_pretrained.return_unused_kwargs",description:`<strong>return_unused_kwargs</strong> (<code>bool</code>, <em>optional</em>, defaults to <code>False</code>) — | |
| If <code>False</code>, then this function returns just the final configuration object.</p> | |
| <p>If <code>True</code>, then this functions returns a <code>Tuple(config, unused_kwargs)</code> where <em>unused_kwargs</em> is a | |
| dictionary consisting of the key/value pairs whose keys are not configuration attributes: i.e., the | |
| part of <code>kwargs</code> which has not been used to update <code>config</code> and is otherwise ignored.`,name:"return_unused_kwargs"},{anchor:"transformers.PretrainedConfig.from_pretrained.subfolder",description:`<strong>subfolder</strong> (<code>str</code>, <em>optional</em>, defaults to <code>""</code>) — | |
| In case the relevant files are located inside a subfolder of the model repo on huggingface.co, you can | |
| specify the folder name here.`,name:"subfolder"},{anchor:"transformers.PretrainedConfig.from_pretrained.kwargs",description:`<strong>kwargs</strong> (<code>Dict[str, Any]</code>, <em>optional</em>) — | |
| The values in kwargs of any keys which are configuration attributes will be used to override the loaded | |
| values. Behavior concerning key/value pairs whose keys are <em>not</em> configuration attributes is controlled | |
| by the <code>return_unused_kwargs</code> keyword parameter.`,name:"kwargs"}],source:"https://github.com/huggingface/transformers/blob/v4.44.2/src/transformers/configuration_utils.py#L510",returnDescription:`<script context="module">export const metadata = 'undefined';<\/script> | |
| <p>The configuration object instantiated from this pretrained model.</p> | |
| `,returnType:`<script context="module">export const metadata = 'undefined';<\/script> | |
| <p><a | |
| href="/docs/transformers/v4.44.2/zh/main_classes/configuration#transformers.PretrainedConfig" | |
| >PretrainedConfig</a></p> | |
| `}}),q=new lo({props:{anchor:"transformers.PretrainedConfig.from_pretrained.example",$$slots:{default:[wo]},$$scope:{ctx:j}}}),ne=new T({props:{name:"get_config_dict",anchor:"transformers.PretrainedConfig.get_config_dict",parameters:[{name:"pretrained_model_name_or_path",val:": Union"},{name:"**kwargs",val:""}],parametersDescription:[{anchor:"transformers.PretrainedConfig.get_config_dict.pretrained_model_name_or_path",description:`<strong>pretrained_model_name_or_path</strong> (<code>str</code> or <code>os.PathLike</code>) — | |
| The identifier of the pre-trained checkpoint from which we want the dictionary of parameters.`,name:"pretrained_model_name_or_path"}],source:"https://github.com/huggingface/transformers/blob/v4.44.2/src/transformers/configuration_utils.py#L612",returnDescription:`<script context="module">export const metadata = 'undefined';<\/script> | |
| <p>The dictionary(ies) that will be used to instantiate the configuration object.</p> | |
| `,returnType:`<script context="module">export const metadata = 'undefined';<\/script> | |
| <p><code>Tuple[Dict, Dict]</code></p> | |
| `}}),re=new T({props:{name:"register_for_auto_class",anchor:"transformers.PretrainedConfig.register_for_auto_class",parameters:[{name:"auto_class",val:" = 'AutoConfig'"}],parametersDescription:[{anchor:"transformers.PretrainedConfig.register_for_auto_class.auto_class",description:`<strong>auto_class</strong> (<code>str</code> or <code>type</code>, <em>optional</em>, defaults to <code>"AutoConfig"</code>) — | |
| The auto class to register this new configuration with.`,name:"auto_class"}],source:"https://github.com/huggingface/transformers/blob/v4.44.2/src/transformers/configuration_utils.py#L1025"}}),H=new co({props:{warning:!0,$$slots:{default:[xo]},$$scope:{ctx:j}}}),ae=new T({props:{name:"save_pretrained",anchor:"transformers.PretrainedConfig.save_pretrained",parameters:[{name:"save_directory",val:": Union"},{name:"push_to_hub",val:": bool = False"},{name:"**kwargs",val:""}],parametersDescription:[{anchor:"transformers.PretrainedConfig.save_pretrained.save_directory",description:`<strong>save_directory</strong> (<code>str</code> or <code>os.PathLike</code>) — | |
| Directory where the configuration JSON file will be saved (will be created if it does not exist).`,name:"save_directory"},{anchor:"transformers.PretrainedConfig.save_pretrained.push_to_hub",description:`<strong>push_to_hub</strong> (<code>bool</code>, <em>optional</em>, defaults to <code>False</code>) — | |
| Whether or not to push your model to the Hugging Face model hub after saving it. You can specify the | |
| repository you want to push to with <code>repo_id</code> (will default to the name of <code>save_directory</code> in your | |
| namespace).`,name:"push_to_hub"},{anchor:"transformers.PretrainedConfig.save_pretrained.kwargs",description:`<strong>kwargs</strong> (<code>Dict[str, Any]</code>, <em>optional</em>) — | |
| Additional key word arguments passed along to the <a href="/docs/transformers/v4.44.2/zh/main_classes/model#transformers.utils.PushToHubMixin.push_to_hub">push_to_hub()</a> method.`,name:"kwargs"}],source:"https://github.com/huggingface/transformers/blob/v4.44.2/src/transformers/configuration_utils.py#L423"}}),se=new T({props:{name:"to_dict",anchor:"transformers.PretrainedConfig.to_dict",parameters:[],source:"https://github.com/huggingface/transformers/blob/v4.44.2/src/transformers/configuration_utils.py#L891",returnDescription:`<script context="module">export const metadata = 'undefined';<\/script> | |
| <p>Dictionary of all the attributes that make up this configuration instance.</p> | |
| `,returnType:`<script context="module">export const metadata = 'undefined';<\/script> | |
| <p><code>Dict[str, Any]</code></p> | |
| `}}),ie=new T({props:{name:"to_diff_dict",anchor:"transformers.PretrainedConfig.to_diff_dict",parameters:[],source:"https://github.com/huggingface/transformers/blob/v4.44.2/src/transformers/configuration_utils.py#L834",returnDescription:`<script context="module">export const metadata = 'undefined';<\/script> | |
| <p>Dictionary of all the attributes that make up this configuration instance,</p> | |
| `,returnType:`<script context="module">export const metadata = 'undefined';<\/script> | |
| <p><code>Dict[str, Any]</code></p> | |
| `}}),de=new T({props:{name:"to_json_file",anchor:"transformers.PretrainedConfig.to_json_file",parameters:[{name:"json_file_path",val:": Union"},{name:"use_diff",val:": bool = True"}],parametersDescription:[{anchor:"transformers.PretrainedConfig.to_json_file.json_file_path",description:`<strong>json_file_path</strong> (<code>str</code> or <code>os.PathLike</code>) — | |
| Path to the JSON file in which this configuration instance’s parameters will be saved.`,name:"json_file_path"},{anchor:"transformers.PretrainedConfig.to_json_file.use_diff",description:`<strong>use_diff</strong> (<code>bool</code>, <em>optional</em>, defaults to <code>True</code>) — | |
| If set to <code>True</code>, only the difference between the config instance and the default <code>PretrainedConfig()</code> | |
| is serialized to JSON file.`,name:"use_diff"}],source:"https://github.com/huggingface/transformers/blob/v4.44.2/src/transformers/configuration_utils.py#L951"}}),ce=new T({props:{name:"to_json_string",anchor:"transformers.PretrainedConfig.to_json_string",parameters:[{name:"use_diff",val:": bool = True"}],parametersDescription:[{anchor:"transformers.PretrainedConfig.to_json_string.use_diff",description:`<strong>use_diff</strong> (<code>bool</code>, <em>optional</em>, defaults to <code>True</code>) — | |
| If set to <code>True</code>, only the difference between the config instance and the default <code>PretrainedConfig()</code> | |
| is serialized to JSON string.`,name:"use_diff"}],source:"https://github.com/huggingface/transformers/blob/v4.44.2/src/transformers/configuration_utils.py#L933",returnDescription:`<script context="module">export const metadata = 'undefined';<\/script> | |
| <p>String containing all the attributes that make up this configuration instance in JSON format.</p> | |
| `,returnType:`<script context="module">export const metadata = 'undefined';<\/script> | |
| <p><code>str</code></p> | |
| `}}),le=new T({props:{name:"update",anchor:"transformers.PretrainedConfig.update",parameters:[{name:"config_dict",val:": Dict"}],parametersDescription:[{anchor:"transformers.PretrainedConfig.update.config_dict",description:"<strong>config_dict</strong> (<code>Dict[str, Any]</code>) — Dictionary of attributes that should be updated for this class.",name:"config_dict"}],source:"https://github.com/huggingface/transformers/blob/v4.44.2/src/transformers/configuration_utils.py#L965"}}),me=new T({props:{name:"update_from_string",anchor:"transformers.PretrainedConfig.update_from_string",parameters:[{name:"update_str",val:": str"}],parametersDescription:[{anchor:"transformers.PretrainedConfig.update_from_string.update_str",description:"<strong>update_str</strong> (<code>str</code>) — String with attributes that should be updated for this class.",name:"update_str"}],source:"https://github.com/huggingface/transformers/blob/v4.44.2/src/transformers/configuration_utils.py#L975"}}),fe=new bo({props:{source:"https://github.com/huggingface/transformers/blob/main/docs/source/zh/main_classes/configuration.md"}}),{c(){d=s("meta"),k=o(),f=s("p"),h=o(),g(x.$$.fragment),l=o(),P=s("p"),P.innerHTML=qt,qe=o(),A=s("p"),A.innerHTML=Vt,Ve=o(),g(G.$$.fragment),He=o(),r=s("div"),g(Q.$$.fragment),et=o(),ge=s("p"),ge.textContent=Ht,tt=o(),g(Z.$$.fragment),ot=o(),ue=s("p"),ue.textContent=Nt,nt=o(),_e=s("ul"),_e.innerHTML=Bt,rt=o(),be=s("p"),be.textContent=Et,at=o(),ye=s("ul"),ye.innerHTML=St,st=o(),z=s("div"),g(O.$$.fragment),it=o(),ve=s("p"),ve.textContent=Xt,dt=o(),g(I.$$.fragment),ct=o(),L=s("div"),g(K.$$.fragment),lt=o(),we=s("p"),we.innerHTML=Rt,mt=o(),F=s("div"),g(ee.$$.fragment),ft=o(),xe=s("p"),xe.innerHTML=Yt,pt=o(),D=s("div"),g(te.$$.fragment),ht=o(),Ce=s("p"),Ce.innerHTML=At,gt=o(),J=s("div"),g(oe.$$.fragment),ut=o(),$e=s("p"),$e.innerHTML=Gt,_t=o(),g(q.$$.fragment),bt=o(),V=s("div"),g(ne.$$.fragment),yt=o(),ke=s("p"),ke.innerHTML=Qt,vt=o(),U=s("div"),g(re.$$.fragment),wt=o(),Pe=s("p"),Pe.innerHTML=Ot,xt=o(),g(H.$$.fragment),Ct=o(),N=s("div"),g(ae.$$.fragment),$t=o(),Te=s("p"),Te.innerHTML=Kt,kt=o(),B=s("div"),g(se.$$.fragment),Pt=o(),Me=s("p"),Me.textContent=eo,Tt=o(),E=s("div"),g(ie.$$.fragment),Mt=o(),je=s("p"),je.textContent=to,jt=o(),S=s("div"),g(de.$$.fragment),zt=o(),ze=s("p"),ze.textContent=oo,Jt=o(),X=s("div"),g(ce.$$.fragment),Ut=o(),Je=s("p"),Je.textContent=no,Wt=o(),R=s("div"),g(le.$$.fragment),Zt=o(),Ue=s("p"),Ue.innerHTML=ro,It=o(),M=s("div"),g(me.$$.fragment),Lt=o(),We=s("p"),We.innerHTML=ao,Ft=o(),Ze=s("p"),Ze.innerHTML=so,Dt=o(),Ie=s("p"),Ie.textContent=io,Ne=o(),g(fe.$$.fragment),Be=o(),De=s("p"),this.h()},l(e){const m=_o("svelte-u9bgzb",document.head);d=i(m,"META",{name:!0,content:!0}),m.forEach(c),k=n(e),f=i(e,"P",{}),C(f).forEach(c),h=n(e),u(x.$$.fragment,e),l=n(e),P=i(e,"P",{"data-svelte-h":!0}),p(P)!=="svelte-s6rnoi"&&(P.innerHTML=qt),qe=n(e),A=i(e,"P",{"data-svelte-h":!0}),p(A)!=="svelte-1u9ssan"&&(A.innerHTML=Vt),Ve=n(e),u(G.$$.fragment,e),He=n(e),r=i(e,"DIV",{class:!0});var a=C(r);u(Q.$$.fragment,a),et=n(a),ge=i(a,"P",{"data-svelte-h":!0}),p(ge)!=="svelte-1p9qi0i"&&(ge.textContent=Ht),tt=n(a),u(Z.$$.fragment,a),ot=n(a),ue=i(a,"P",{"data-svelte-h":!0}),p(ue)!=="svelte-1qxxkvo"&&(ue.textContent=Nt),nt=n(a),_e=i(a,"UL",{"data-svelte-h":!0}),p(_e)!=="svelte-16ll04v"&&(_e.innerHTML=Bt),rt=n(a),be=i(a,"P",{"data-svelte-h":!0}),p(be)!=="svelte-1yi3c0w"&&(be.textContent=Et),at=n(a),ye=i(a,"UL",{"data-svelte-h":!0}),p(ye)!=="svelte-1cjcwii"&&(ye.innerHTML=St),st=n(a),z=i(a,"DIV",{class:!0});var W=C(z);u(O.$$.fragment,W),it=n(W),ve=i(W,"P",{"data-svelte-h":!0}),p(ve)!=="svelte-j50kqf"&&(ve.textContent=Xt),dt=n(W),u(I.$$.fragment,W),W.forEach(c),ct=n(a),L=i(a,"DIV",{class:!0});var pe=C(L);u(K.$$.fragment,pe),lt=n(pe),we=i(pe,"P",{"data-svelte-h":!0}),p(we)!=="svelte-m61tyl"&&(we.innerHTML=Rt),pe.forEach(c),mt=n(a),F=i(a,"DIV",{class:!0});var he=C(F);u(ee.$$.fragment,he),ft=n(he),xe=i(he,"P",{"data-svelte-h":!0}),p(xe)!=="svelte-36v0yf"&&(xe.innerHTML=Yt),he.forEach(c),pt=n(a),D=i(a,"DIV",{class:!0});var Se=C(D);u(te.$$.fragment,Se),ht=n(Se),Ce=i(Se,"P",{"data-svelte-h":!0}),p(Ce)!=="svelte-1ne5ckw"&&(Ce.innerHTML=At),Se.forEach(c),gt=n(a),J=i(a,"DIV",{class:!0});var Le=C(J);u(oe.$$.fragment,Le),ut=n(Le),$e=i(Le,"P",{"data-svelte-h":!0}),p($e)!=="svelte-1mdoy7e"&&($e.innerHTML=Gt),_t=n(Le),u(q.$$.fragment,Le),Le.forEach(c),bt=n(a),V=i(a,"DIV",{class:!0});var Xe=C(V);u(ne.$$.fragment,Xe),yt=n(Xe),ke=i(Xe,"P",{"data-svelte-h":!0}),p(ke)!=="svelte-2rltva"&&(ke.innerHTML=Qt),Xe.forEach(c),vt=n(a),U=i(a,"DIV",{class:!0});var Fe=C(U);u(re.$$.fragment,Fe),wt=n(Fe),Pe=i(Fe,"P",{"data-svelte-h":!0}),p(Pe)!=="svelte-30y31e"&&(Pe.innerHTML=Ot),xt=n(Fe),u(H.$$.fragment,Fe),Fe.forEach(c),Ct=n(a),N=i(a,"DIV",{class:!0});var Re=C(N);u(ae.$$.fragment,Re),$t=n(Re),Te=i(Re,"P",{"data-svelte-h":!0}),p(Te)!=="svelte-ygts64"&&(Te.innerHTML=Kt),Re.forEach(c),kt=n(a),B=i(a,"DIV",{class:!0});var Ye=C(B);u(se.$$.fragment,Ye),Pt=n(Ye),Me=i(Ye,"P",{"data-svelte-h":!0}),p(Me)!=="svelte-1ww3wqq"&&(Me.textContent=eo),Ye.forEach(c),Tt=n(a),E=i(a,"DIV",{class:!0});var Ae=C(E);u(ie.$$.fragment,Ae),Mt=n(Ae),je=i(Ae,"P",{"data-svelte-h":!0}),p(je)!=="svelte-1p6bdas"&&(je.textContent=to),Ae.forEach(c),jt=n(a),S=i(a,"DIV",{class:!0});var Ge=C(S);u(de.$$.fragment,Ge),zt=n(Ge),ze=i(Ge,"P",{"data-svelte-h":!0}),p(ze)!=="svelte-1g70y32"&&(ze.textContent=oo),Ge.forEach(c),Jt=n(a),X=i(a,"DIV",{class:!0});var Qe=C(X);u(ce.$$.fragment,Qe),Ut=n(Qe),Je=i(Qe,"P",{"data-svelte-h":!0}),p(Je)!=="svelte-5ayq1f"&&(Je.textContent=no),Qe.forEach(c),Wt=n(a),R=i(a,"DIV",{class:!0});var Oe=C(R);u(le.$$.fragment,Oe),Zt=n(Oe),Ue=i(Oe,"P",{"data-svelte-h":!0}),p(Ue)!=="svelte-1hh5fg7"&&(Ue.innerHTML=ro),Oe.forEach(c),It=n(a),M=i(a,"DIV",{class:!0});var Y=C(M);u(me.$$.fragment,Y),Lt=n(Y),We=i(Y,"P",{"data-svelte-h":!0}),p(We)!=="svelte-a2aodj"&&(We.innerHTML=ao),Ft=n(Y),Ze=i(Y,"P",{"data-svelte-h":!0}),p(Ze)!=="svelte-179z5e8"&&(Ze.innerHTML=so),Dt=n(Y),Ie=i(Y,"P",{"data-svelte-h":!0}),p(Ie)!=="svelte-5bouux"&&(Ie.textContent=io),Y.forEach(c),a.forEach(c),Ne=n(e),u(fe.$$.fragment,e),Be=n(e),De=i(e,"P",{}),C(De).forEach(c),this.h()},h(){$(d,"name","hf:doc:metadata"),$(d,"content",$o),$(z,"class","docstring border-l-2 border-t-2 pl-4 pt-3.5 border-gray-100 rounded-tl-xl mb-6 mt-8"),$(L,"class","docstring border-l-2 border-t-2 pl-4 pt-3.5 border-gray-100 rounded-tl-xl mb-6 mt-8"),$(F,"class","docstring border-l-2 border-t-2 pl-4 pt-3.5 border-gray-100 rounded-tl-xl mb-6 mt-8"),$(D,"class","docstring border-l-2 border-t-2 pl-4 pt-3.5 border-gray-100 rounded-tl-xl mb-6 mt-8"),$(J,"class","docstring border-l-2 border-t-2 pl-4 pt-3.5 border-gray-100 rounded-tl-xl mb-6 mt-8"),$(V,"class","docstring border-l-2 border-t-2 pl-4 pt-3.5 border-gray-100 rounded-tl-xl mb-6 mt-8"),$(U,"class","docstring border-l-2 border-t-2 pl-4 pt-3.5 border-gray-100 rounded-tl-xl mb-6 mt-8"),$(N,"class","docstring border-l-2 border-t-2 pl-4 pt-3.5 border-gray-100 rounded-tl-xl mb-6 mt-8"),$(B,"class","docstring border-l-2 border-t-2 pl-4 pt-3.5 border-gray-100 rounded-tl-xl mb-6 mt-8"),$(E,"class","docstring border-l-2 border-t-2 pl-4 pt-3.5 border-gray-100 rounded-tl-xl mb-6 mt-8"),$(S,"class","docstring border-l-2 border-t-2 pl-4 pt-3.5 border-gray-100 rounded-tl-xl mb-6 mt-8"),$(X,"class","docstring border-l-2 border-t-2 pl-4 pt-3.5 border-gray-100 rounded-tl-xl mb-6 mt-8"),$(R,"class","docstring border-l-2 border-t-2 pl-4 pt-3.5 border-gray-100 rounded-tl-xl mb-6 mt-8"),$(M,"class","docstring border-l-2 border-t-2 pl-4 pt-3.5 border-gray-100 rounded-tl-xl mb-6 mt-8"),$(r,"class","docstring border-l-2 border-t-2 pl-4 pt-3.5 border-gray-100 rounded-tl-xl mb-6 mt-8")},m(e,m){t(document.head,d),w(e,k,m),w(e,f,m),w(e,h,m),_(x,e,m),w(e,l,m),w(e,P,m),w(e,qe,m),w(e,A,m),w(e,Ve,m),_(G,e,m),w(e,He,m),w(e,r,m),_(Q,r,null),t(r,et),t(r,ge),t(r,tt),_(Z,r,null),t(r,ot),t(r,ue),t(r,nt),t(r,_e),t(r,rt),t(r,be),t(r,at),t(r,ye),t(r,st),t(r,z),_(O,z,null),t(z,it),t(z,ve),t(z,dt),_(I,z,null),t(r,ct),t(r,L),_(K,L,null),t(L,lt),t(L,we),t(r,mt),t(r,F),_(ee,F,null),t(F,ft),t(F,xe),t(r,pt),t(r,D),_(te,D,null),t(D,ht),t(D,Ce),t(r,gt),t(r,J),_(oe,J,null),t(J,ut),t(J,$e),t(J,_t),_(q,J,null),t(r,bt),t(r,V),_(ne,V,null),t(V,yt),t(V,ke),t(r,vt),t(r,U),_(re,U,null),t(U,wt),t(U,Pe),t(U,xt),_(H,U,null),t(r,Ct),t(r,N),_(ae,N,null),t(N,$t),t(N,Te),t(r,kt),t(r,B),_(se,B,null),t(B,Pt),t(B,Me),t(r,Tt),t(r,E),_(ie,E,null),t(E,Mt),t(E,je),t(r,jt),t(r,S),_(de,S,null),t(S,zt),t(S,ze),t(r,Jt),t(r,X),_(ce,X,null),t(X,Ut),t(X,Je),t(r,Wt),t(r,R),_(le,R,null),t(R,Zt),t(R,Ue),t(r,It),t(r,M),_(me,M,null),t(M,Lt),t(M,We),t(M,Ft),t(M,Ze),t(M,Dt),t(M,Ie),w(e,Ne,m),_(fe,e,m),w(e,Be,m),w(e,De,m),Ee=!0},p(e,[m]){const a={};m&2&&(a.$$scope={dirty:m,ctx:e}),Z.$set(a);const W={};m&2&&(W.$$scope={dirty:m,ctx:e}),I.$set(W);const pe={};m&2&&(pe.$$scope={dirty:m,ctx:e}),q.$set(pe);const he={};m&2&&(he.$$scope={dirty:m,ctx:e}),H.$set(he)},i(e){Ee||(b(x.$$.fragment,e),b(G.$$.fragment,e),b(Q.$$.fragment,e),b(Z.$$.fragment,e),b(O.$$.fragment,e),b(I.$$.fragment,e),b(K.$$.fragment,e),b(ee.$$.fragment,e),b(te.$$.fragment,e),b(oe.$$.fragment,e),b(q.$$.fragment,e),b(ne.$$.fragment,e),b(re.$$.fragment,e),b(H.$$.fragment,e),b(ae.$$.fragment,e),b(se.$$.fragment,e),b(ie.$$.fragment,e),b(de.$$.fragment,e),b(ce.$$.fragment,e),b(le.$$.fragment,e),b(me.$$.fragment,e),b(fe.$$.fragment,e),Ee=!0)},o(e){y(x.$$.fragment,e),y(G.$$.fragment,e),y(Q.$$.fragment,e),y(Z.$$.fragment,e),y(O.$$.fragment,e),y(I.$$.fragment,e),y(K.$$.fragment,e),y(ee.$$.fragment,e),y(te.$$.fragment,e),y(oe.$$.fragment,e),y(q.$$.fragment,e),y(ne.$$.fragment,e),y(re.$$.fragment,e),y(H.$$.fragment,e),y(ae.$$.fragment,e),y(se.$$.fragment,e),y(ie.$$.fragment,e),y(de.$$.fragment,e),y(ce.$$.fragment,e),y(le.$$.fragment,e),y(me.$$.fragment,e),y(fe.$$.fragment,e),Ee=!1},d(e){e&&(c(k),c(f),c(h),c(l),c(P),c(qe),c(A),c(Ve),c(He),c(r),c(Ne),c(Be),c(De)),c(d),v(x,e),v(G,e),v(Q),v(Z),v(O),v(I),v(K),v(ee),v(te),v(oe),v(q),v(ne),v(re),v(H),v(ae),v(se),v(ie),v(de),v(ce),v(le),v(me),v(fe,e)}}}const $o='{"title":"Configuration","local":"configuration","sections":[{"title":"PretrainedConfig","local":"transformers.PretrainedConfig","sections":[],"depth":2}],"depth":1}';function ko(j){return ho(()=>{new URLSearchParams(window.location.search).get("fw")}),[]}class Wo extends go{constructor(d){super(),uo(this,d,ko,Co,po,{})}}export{Wo as component}; | |
Xet Storage Details
- Size:
- 60.4 kB
- Xet hash:
- ca41f744ed2b77a94b9a3a6bb040067aea861b12975f2670f6c71480737fc4d6
·
Xet efficiently stores files, intelligently splitting them into unique chunks and accelerating uploads and downloads. More info.