§
    ‚Štjä  ã                   óF  — d dl mZ ddlmZ ddlmZmZ  ej        e¦  «        Z	 ed¬¦  «        e G d„ de¦  «        ¦   «         ¦   «         Z
 ed	¬¦  «        e G d
„ de¦  «        ¦   «         ¦   «         Z ed	¬¦  «        e G d„ de¦  «        ¦   «         ¦   «         Zg d¢ZdS )é    )Ústricté   )ÚPreTrainedConfig)Úauto_docstringÚloggingzgoogle/videoprism-base-f16r288)Ú
checkpointc                   ón  — e Zd ZU dZdZdZeee         z  eeef         z  e	d<   dZ
ee	d<   dZee         eedf         z  e	d	<   d
Zee	d<   dZee	d<   dZee	d<   dZee	d<   dZee	d<   dZeez  e	d<   dZeez  e	d<   dZee	d<   dZee	d<   dZee	d<   dZdZee	d<   dZee	d <   d!Zee	d"<   d#Zee	d$<   dZee	d%<   d&S )'ÚVideoPrismVisionConfiga™  
    num_frames (`int`, *optional*, defaults to 16):
        The number of frames in the input video.
    tubelet_size (`List[int]`, *optional*, defaults to `[1, 18, 18]`):
        The size of the tubelet patch.
    num_spatial_layers (`int`, *optional*, defaults to 12):
        Number of spatial transformer blocks.
    num_temporal_layers (`int`, *optional*, defaults to 4):
        Number of temporal transformer blocks.
    attn_logit_softcapping (`float`, *optional*, defaults to 50.0):
        Softcapping constant for attention logits.
    num_auxiliary_layers (`int`, *optional*, defaults to 2):
        Number of auxiliary layers. This is used in the VideoPrismVideoModel that is a part of VideoPrismClipModel.
    apply_l2norm (`bool`, *optional*, defaults to `True`):
        Whether to apply L2 normalization to the output. This is used in the VideoPrismVideoModel that is a part of VideoPrismClipModel.
    Úvideoprism_vision_modeli   Ú
image_sizeé   Ú
num_frames)é   é   r   .Útubelet_sizer   Únum_channelsé   Úhidden_sizeé   Únum_attention_headsé   Úintermediate_sizeÚgelu_pythonÚ
hidden_actç        Úhidden_dropout_probÚattention_probs_dropout_probç{®Gáz”?Úinitializer_rangeç�íµ ÷Æ°>Úlayer_norm_epsTÚqkv_biasÚvision_configÚnum_spatial_layersé   Únum_temporal_layersç      I@Úattn_logit_softcappingé   Únum_auxiliary_layersÚapply_l2normN)Ú__name__Ú
__module__Ú__qualname__Ú__doc__Ú
model_typer   ÚintÚlistÚtupleÚ__annotations__r   r   r   r   r   r   r   Ústrr   Úfloatr   r   r!   r"   ÚboolÚbase_config_keyr$   r&   r(   r*   r+   © ó    úu/var/www/html/CA-Chatbot/venv/lib/python3.11/site-packages/transformers/models/videoprism/configuration_videoprism.pyr
   r
      s  € € € € € € ðð ð" +€JØ47€J��d˜3”i‘ %¨¨S¨¤/Ñ1Ð7Ð7Ñ7Ø€J�ÐÐÑØ0;€L�$�s”)˜e C¨ HœoÑ-Ð;Ð;Ñ;Ø€L�#ÐÐÑØ€K�ÐÐÑØ!Ð˜Ð!Ð!Ñ!Ø!Ð�sÐ!Ð!Ñ!Ø#€J�Ð#Ð#Ñ#Ø'*Ð˜ ™Ð*Ð*Ñ*Ø03Ð  %¨#¡+Ð3Ð3Ñ3Ø#Ð�uÐ#Ð#Ñ#Ø!€N�EÐ!Ð!Ñ!Ø€HˆdÐÐÑØ%€OØ Ð˜Ð Ð Ñ Ø Ð˜Ð Ð Ñ Ø$(Ð˜EÐ(Ð(Ñ(Ø !Ð˜#Ð!Ð!Ñ!Ø€L�$ÐÐÑÐÐr:   r
   z"google/videoprism-lvt-base-f16r288c                   ó4  — e Zd ZU dZdZdZdZeed<   dZ	eed<   dZ
eed	<   d
Zeed<   d
Zeed<   dZeed<   dZeed<   dZeed<   dZedz  ed<   dZedz  ed<   dZeee         z  dz  ed<   dZeez  ed<   dZeed<   dZeed<   dZeed<   dZeed<   d Zeed!<   dS )"ÚVideoPrismTextConfiga	  
    apply_l2norm (`bool`, *optional*, defaults to `True`):
        Whether to apply L2 normalization to the output of VideoPrismTextEncoder.
    attn_logit_softcapping (`float`, *optional*, defaults to 50.0):
        Softcapping constant for attention logits.
    Úvideoprism_text_modelÚtext_configi }  Ú
vocab_sizer   r   r   r   r   Únum_hidden_layersr   é@   Úmax_position_embeddingsÚrelur   r    r!   r   NÚpad_token_idÚbos_token_idÚeos_token_idr   r   Tr+   r"   r   r   r   r'   r(   )r,   r-   r.   r/   r0   r8   r@   r1   r4   r   r   rA   r   rC   r   r5   r!   r6   rE   rF   rG   r2   r   r+   r7   r"   r   r   r(   r9   r:   r;   r=   r=   I   sY  € € € € € € ðð ð )€JØ#€Oà€J�ÐÐÑØ€K�ÐÐÑØ!Ð�sÐ!Ð!Ñ!ØÐ�sÐÐÑØ!Ð˜Ð!Ð!Ñ!Ø#%Ð˜SÐ%Ð%Ñ%à€J�ÐÐÑØ €N�EÐ Ð Ñ Ø €L�#˜‘*Ð Ð Ñ Ø#€L�#˜‘*Ð#Ð#Ñ#Ø+/€L�#˜˜Sœ	‘/ DÑ(Ð/Ð/Ñ/Ø03Ð  %¨#¡+Ð3Ð3Ñ3Ø€L�$ÐÐÑØ€HˆdÐÐÑØ!$Ð˜Ð$Ð$Ñ$Ø#Ð�uÐ#Ð#Ñ#Ø$(Ð˜EÐ(Ð(Ñ(Ð(Ð(r:   r=   c                   óf   ‡ — e Zd ZU dZdZeedœZdZe	e
z  dz  ed<   dZe	e
z  dz  ed<   ˆ fd„Zˆ xZS )ÚVideoPrismConfiga¤  
    Example:

    ```python
    >>> from transformers import VideoPrismClipModel, VideoPrismConfig

    >>> # Initializing a VideoPrismConfig with default values
    >>> configuration = VideoPrismConfig()

    >>> # Initializing a VideoPrismClipModel with the configuration
    >>> model = VideoPrismClipModel(configuration)

    >>> # Accessing the model configuration
    >>> configuration = model.config
    ```
    Ú
videoprism)r?   r#   Nr?   r#   c                 óÎ  •— | j         €.t          ¦   «         | _         t                               d¦  «         n0t	          | j         t
          ¦  «        rt          di | j         ¤Ž| _         | j        €.t          ¦   «         | _        t                               d¦  «         n0t	          | j        t
          ¦  «        rt          di | j        ¤Ž| _         t          ¦   «         j	        di |¤Ž d S )NzU`text_config` is `None`. Initializing the `VideoPrismTextConfig` with default values.zY`vision_config` is `None`. initializing the `VideoPrismVisionConfig` with default values.r9   )
r?   r=   ÚloggerÚinfoÚ
isinstanceÚdictr#   r
   ÚsuperÚ__post_init__)ÚselfÚkwargsÚ	__class__s     €r;   rQ   zVideoPrismConfig.__post_init__„   sÞ   ø€ ØÔÐ#Ý3Ñ5Ô5ˆDÔÝ�KŠKÐoÑpÔpÐpÐpÝ˜Ô(­$Ñ/Ô/ð 	HÝ3ÐGÐG°dÔ6FÐGÐGˆDÔàÔÐ%Ý!7Ñ!9Ô!9ˆDÔÝ�KŠKÐsÑtÔtÐtÐtÝ˜Ô*­DÑ1Ô1ð 	NÝ!7Ð!MÐ!M¸$Ô:LÐ!MÐ!MˆDÔà�‰ŒÔÐ'Ð' Ð'Ð'Ð'Ð'Ð'r:   )r,   r-   r.   r/   r0   r=   r
   Úsub_configsr?   rO   r   r4   r#   rQ   Ú__classcell__)rT   s   @r;   rI   rI   j   s�   ø€ € € € € € ðð ð" €JØ"6ÐI_Ð`Ð`€Kà26€K�Ð(Ñ(¨4Ñ/Ð6Ð6Ñ6Ø48€M�4Ð*Ñ*¨TÑ1Ð8Ð8Ñ8ð(ð (ð (ð (ð (ð (ð (ð (ð (r:   rI   )r
   r=   rI   N)Úhuggingface_hub.dataclassesr   Úconfiguration_utilsr   Úutilsr   r   Ú
get_loggerr,   rL   r
   r=   rI   Ú__all__r9   r:   r;   ú<module>r\      s]  ðð, /Ð .Ð .Ð .Ð .Ð .à 3Ð 3Ð 3Ð 3Ð 3Ð 3Ø ,Ð ,Ð ,Ð ,Ð ,Ð ,Ð ,Ð ,ð 
ˆÔ	˜HÑ	%Ô	%€ð €Ð;Ð<Ñ<Ô<Øð%ð %ð %ð %ð %Ð-ñ %ô %ñ „ñ =Ô<ð%ðP €Ð?Ð@Ñ@Ô@Øð)ð )ð )ð )ð )Ð+ñ )ô )ñ „ñ AÔ@ð)ð> €Ð?Ð@Ñ@Ô@Øð%(ð %(ð %(ð %(ð %(Ð'ñ %(ô %(ñ „ñ AÔ@ð%(ðP QÐ
PÐ
P€€€r:   