U
    ƒIÀdx@  ã                   @   s´  d Z ddlZddlZddlZddlZddlZddlZddlZddl	Z
ddlZ
ddlmZ ddlmZ ddlmZmZmZmZmZmZmZmZmZmZmZ ddlmZ ddlmZ ddl m!Z!m"Z" dd	l#m$Z$ dd
l%m&Z& ddl'm(Z( ddl)m*Z* ddl+m,Z, ddl-m.Z. ddl/m0Z0 ddl1m2Z2 ddl3m4Z4m5Z5m6Z6 e�rHddlm7Z7 ne8Z7e 9e:¡Z;ee<e<f Z=e<ee< dœdd„Z>G dd„ de?ƒZ@eddœdd„ZAG dd„ de?ƒZBe<e*ddœdd „ZCe<e*edœd!d"„ZDe=ee< d#œd$d%„ZEG d&d'„ d'ƒZFG d(d)„ d)e7ƒZGeGeGd*œd+d,„ZHeHd-ee& d.œd/d0„ƒZIG d1d-„ d-ƒZJG d2d3„ d3eƒZKdCe&ee<e?f eed4  dd5œd6d7„ZLdDeeMeJd9œd:d;„ZNe&e*ed- d<œd=d>„ZOG d?d@„ d@eƒZPG dAdB„ dBƒZQdS )EzO
The main purpose of this module is to expose LinkCollector.collect_sources().
é    N)Ú
HTMLParser)ÚValues)ÚTYPE_CHECKINGÚCallableÚDictÚIterableÚListÚMutableMappingÚ
NamedTupleÚOptionalÚSequenceÚTupleÚUnion)Úrequests)ÚResponse)Ú
RetryErrorÚSSLError)ÚNetworkConnectionError)ÚLink)ÚSearchScope)Ú
PipSession)Úraise_for_status)Úis_archive_file©Úredact_auth_from_url)Úvcsé   )ÚCandidatesFromPageÚ
LinkSourceÚbuild_source)ÚProtocol©ÚurlÚreturnc                 C   s6   t jD ]*}|  ¡  |¡r| t|ƒ dkr|  S qdS )zgLook for VCS schemes in the URL.

    Returns the matched VCS scheme, or None if there's no match.
    z+:N)r   ÚschemesÚlowerÚ
startswithÚlen)r"   Úscheme© r)   úV/home/sam/Atlas/atlas_env/lib/python3.8/site-packages/pip/_internal/index/collector.pyÚ_match_vcs_scheme7   s    

r+   c                       s&   e Zd Zeeddœ‡ fdd„Z‡  ZS )Ú_NotAPIContentN)Úcontent_typeÚrequest_descr#   c                    s   t ƒ  ||¡ || _|| _d S ©N)ÚsuperÚ__init__r-   r.   )Úselfr-   r.   ©Ú	__class__r)   r*   r1   C   s    z_NotAPIContent.__init__)Ú__name__Ú
__module__Ú__qualname__Ústrr1   Ú__classcell__r)   r)   r3   r*   r,   B   s   r,   )Úresponser#   c                 C   s6   | j  dd¡}| ¡ }| d¡r$dS t|| jjƒ‚dS )z°
    Check the Content-Type header to ensure the response contains a Simple
    API Response.

    Raises `_NotAPIContent` if the content type is not a valid content-type.
    úContent-TypeÚUnknown)z	text/htmlz#application/vnd.pypi.simple.v1+htmlú#application/vnd.pypi.simple.v1+jsonN)ÚheadersÚgetr%   r&   r,   ÚrequestÚmethod)r:   r-   Úcontent_type_lr)   r)   r*   Ú_ensure_api_headerI   s    ÿrC   c                   @   s   e Zd ZdS )Ú_NotHTTPN)r5   r6   r7   r)   r)   r)   r*   rD   _   s   rD   )r"   Úsessionr#   c                 C   sF   t j | ¡\}}}}}|dkr$tƒ ‚|j| dd�}t|ƒ t|ƒ dS )zõ
    Send a HEAD request to the URL, and ensure the response contains a simple
    API Response.

    Raises `_NotHTTP` if the URL is not available for a HEAD request, or
    `_NotAPIContent` if the content type is not a valid content type.
    >   ÚhttpsÚhttpT)Úallow_redirectsN)ÚurllibÚparseÚurlsplitrD   Úheadr   rC   )r"   rE   r(   ÚnetlocÚpathÚqueryÚfragmentÚrespr)   r)   r*   Ú_ensure_api_responsec   s    rR   c                 C   sz   t t| ƒjƒrt| |d� t dt| ƒ¡ |j| d dddg¡ddœd	�}t	|ƒ t
|ƒ t d
t| ƒ|j dd¡¡ |S )aY  Access an Simple API response with GET, and return the response.

    This consists of three parts:

    1. If the URL looks suspiciously like an archive, send a HEAD first to
       check the Content-Type is HTML or Simple API, to avoid downloading a
       large file. Raise `_NotHTTP` if the content type cannot be determined, or
       `_NotAPIContent` if it is not HTML or a Simple API.
    2. Actually perform the request. Raise HTTP exceptions on network failures.
    3. Check the Content-Type header to make sure we got a Simple API response,
       and raise `_NotAPIContent` otherwise.
    ©rE   zGetting page %sz, r=   z*application/vnd.pypi.simple.v1+html; q=0.1ztext/html; q=0.01z	max-age=0)ÚAcceptzCache-Control)r>   zFetched page %s as %sr;   r<   )r   r   ÚfilenamerR   ÚloggerÚdebugr   r?   Újoinr   rC   r>   )r"   rE   rQ   r)   r)   r*   Ú_get_simple_responseu   s,    ýÿëþýrY   )r>   r#   c                 C   s<   | r8d| kr8t j ¡ }| d |d< | d¡}|r8t|ƒS dS )z=Determine if we have any encoding information in our headers.r;   zcontent-typeÚcharsetN)ÚemailÚmessageÚMessageÚ	get_paramr8   )r>   ÚmrZ   r)   r)   r*   Ú_get_encoding_from_headers´   s    

r`   c                   @   s:   e Zd Zdddœdd„Zeedœdd„Zed	œd
d„ZdS )ÚCacheablePageContentÚIndexContentN©Úpager#   c                 C   s   |j s
t‚|| _d S r/   )Úcache_link_parsingÚAssertionErrorrd   ©r2   rd   r)   r)   r*   r1   À   s    
zCacheablePageContent.__init__)Úotherr#   c                 C   s   t |t| ƒƒo| jj|jjkS r/   )Ú
isinstanceÚtyperd   r"   )r2   rh   r)   r)   r*   Ú__eq__Ä   s    zCacheablePageContent.__eq__©r#   c                 C   s   t | jjƒS r/   )Úhashrd   r"   ©r2   r)   r)   r*   Ú__hash__Ç   s    zCacheablePageContent.__hash__)	r5   r6   r7   r1   ÚobjectÚboolrk   Úintro   r)   r)   r)   r*   ra   ¿   s   ra   c                   @   s    e Zd Zdee dœdd„ZdS )Ú
ParseLinksrb   rc   c                 C   s   d S r/   r)   rg   r)   r)   r*   Ú__call__Ì   s    zParseLinks.__call__N)r5   r6   r7   r   r   rt   r)   r)   r)   r*   rs   Ë   s   rs   )Úfnr#   c                    sL   t jdd�ttt dœ‡ fdd„ƒ‰t  ˆ ¡dtt dœ‡ ‡fdd	„ƒ}|S )
zÚ
    Given a function that parses an Iterable[Link] from an IndexContent, cache the
    function's result (keyed by CacheablePageContent), unless the IndexContent
    `page` has `page.cache_link_parsing == False`.
    N)Úmaxsize)Úcacheable_pager#   c                    s   t ˆ | jƒƒS r/   )Úlistrd   )rw   )ru   r)   r*   Úwrapper×   s    z*with_cached_index_content.<locals>.wrapperrb   rc   c                    s   | j rˆt| ƒƒS tˆ | ƒƒS r/   )re   ra   rx   )rd   ©ru   ry   r)   r*   Úwrapper_wrapperÛ   s    z2with_cached_index_content.<locals>.wrapper_wrapper)Ú	functoolsÚ	lru_cachera   r   r   Úwraps)ru   r{   r)   rz   r*   Úwith_cached_index_contentÐ   s
    
r   rb   rc   c           
      c   sº   | j  ¡ }| d¡rTt | j¡}| dg ¡D ]"}t || j	¡}|dkrHq,|V  q,dS t
| j	ƒ}| jpfd}| | j |¡¡ | j	}|jpˆ|}|jD ]$}	tj|	||d�}|dkr®q�|V  q�dS )z\
    Parse a Simple API's Index Content, and yield its anchor elements as Link objects.
    r=   ÚfilesNzutf-8)Úpage_urlÚbase_url)r-   r%   r&   ÚjsonÚloadsÚcontentr?   r   Ú	from_jsonr"   ÚHTMLLinkParserÚencodingÚfeedÚdecoder‚   ÚanchorsÚfrom_element)
rd   rB   ÚdataÚfileÚlinkÚparserrˆ   r"   r‚   Úanchorr)   r)   r*   Úparse_linksä   s&    





r’   c                   @   s<   e Zd ZdZd
eeee eeddœdd„Zedœdd	„Z	dS )rb   z5Represents one response (or page), along with its URLTN)r…   r-   rˆ   r"   re   r#   c                 C   s"   || _ || _|| _|| _|| _dS )am  
        :param encoding: the encoding to decode the given content.
        :param url: the URL from which the HTML was downloaded.
        :param cache_link_parsing: whether links parsed from this page's url
                                   should be cached. PyPI index urls should
                                   have this set to False, for example.
        N)r…   r-   rˆ   r"   re   )r2   r…   r-   rˆ   r"   re   r)   r)   r*   r1     s
    zIndexContent.__init__rl   c                 C   s
   t | jƒS r/   )r   r"   rn   r)   r)   r*   Ú__str__  s    zIndexContent.__str__)T)
r5   r6   r7   Ú__doc__Úbytesr8   r   rq   r1   r“   r)   r)   r)   r*   rb     s    úùc                       sn   e Zd ZdZeddœ‡ fdd„Zeeeeee f  ddœdd„Z	eeeee f  ee d	œd
d„Z
‡  ZS )r‡   zf
    HTMLParser that keeps the first base HREF and a list of all anchor
    elements' attributes.
    Nr!   c                    s$   t ƒ jdd� || _d | _g | _d S )NT)Úconvert_charrefs)r0   r1   r"   r‚   r‹   )r2   r"   r3   r)   r*   r1   #  s    zHTMLLinkParser.__init__)ÚtagÚattrsr#   c                 C   sH   |dkr,| j d kr,|  |¡}|d k	rD|| _ n|dkrD| j t|ƒ¡ d S )NÚbaseÚa)r‚   Úget_hrefr‹   ÚappendÚdict)r2   r—   r˜   Úhrefr)   r)   r*   Úhandle_starttag*  s    
zHTMLLinkParser.handle_starttag)r˜   r#   c                 C   s"   |D ]\}}|dkr|  S qd S )Nrž   r)   )r2   r˜   ÚnameÚvaluer)   r)   r*   r›   2  s    
zHTMLLinkParser.get_href)r5   r6   r7   r”   r8   r1   r   r   r   rŸ   r›   r9   r)   r)   r3   r*   r‡     s   "r‡   ).N)r�   ÚreasonÚmethr#   c                 C   s   |d krt j}|d| |ƒ d S )Nz%Could not fetch URL %s: %s - skipping)rV   rW   )r�   r¢   r£   r)   r)   r*   Ú_handle_get_simple_fail9  s    r¤   T)r:   re   r#   c                 C   s&   t | jƒ}t| j| jd || j|d�S )Nr;   )rˆ   r"   re   )r`   r>   rb   r…   r"   )r:   re   rˆ   r)   r)   r*   Ú_make_index_contentC  s    
ûr¥   )r�   rE   r#   c          
   
   C   sú  | j  dd¡d }t|ƒ}|r0t d|| ¡ d S tj |¡\}}}}}}|dkr�tj	 
tj |¡¡r�| d¡sv|d7 }tj |d¡}t d|¡ zt||d	�}W �nD tk
rÄ   t d
| ¡ Y �n2 tk
rø } zt d| |j|j¡ W 5 d }~X Y nþ tk
�r$ } zt| |ƒ W 5 d }~X Y nÒ tk
�rP } zt| |ƒ W 5 d }~X Y n¦ tk
�r’ } z$d}	|	t|ƒ7 }	t| |	tjd� W 5 d }~X Y nd tjk
�rÆ } zt| d|› �ƒ W 5 d }~X Y n0 tjk
�ræ   t| dƒ Y nX t|| jd�S d S )Nú#r   r   zICannot look at %s URL %s because it does not support lookup as web pages.rŽ   ú/z
index.htmlz# file: URL is directory, getting %srS   z`Skipping page %s because it looks like an archive, and cannot be checked by a HTTP HEAD request.zºSkipping page %s because the %s request got Content-Type: %s. The only supported Content-Types are application/vnd.pypi.simple.v1+json, application/vnd.pypi.simple.v1+html, and text/htmlz4There was a problem confirming the ssl certificate: )r£   zconnection error: z	timed out)re   ) r"   Úsplitr+   rV   ÚwarningrI   rJ   ÚurlparseÚosrN   Úisdirr@   Úurl2pathnameÚendswithÚurljoinrW   rY   rD   r,   r.   r-   r   r¤   r   r   r8   Úinfor   ÚConnectionErrorÚTimeoutr¥   re   )
r�   rE   r"   Ú
vcs_schemer(   Ú_rN   rQ   Úexcr¢   r)   r)   r*   Ú_get_index_contentP  sV    ý
ý
ú  r¶   c                   @   s.   e Zd ZU eee  ed< eee  ed< dS )ÚCollectedSourcesÚ
find_linksÚ
index_urlsN)r5   r6   r7   r   r   r   Ú__annotations__r)   r)   r)   r*   r·   �  s   
r·   c                   @   sx   e Zd ZdZeeddœdd„Zedeee	d dœdd	„ƒZ
eee d
œdd„ƒZeee dœdd„Zeeedœdd„ZdS )ÚLinkCollectorzµ
    Responsible for collecting Link objects from all configured locations,
    making network requests as needed.

    The class's main method is its collect_sources() method.
    N)rE   Úsearch_scoper#   c                 C   s   || _ || _d S r/   )r¼   rE   )r2   rE   r¼   r)   r)   r*   r1   ›  s    zLinkCollector.__init__F)rE   ÚoptionsÚsuppress_no_indexr#   c                 C   sd   |j g|j }|jr8|s8t dd dd„ |D ƒ¡¡ g }|jp@g }tj|||jd�}t	||d�}|S )zÆ
        :param session: The Session to use to make requests.
        :param suppress_no_index: Whether to ignore the --no-index option
            when constructing the SearchScope object.
        zIgnoring indexes: %sú,c                 s   s   | ]}t |ƒV  qd S r/   r   )Ú.0r"   r)   r)   r*   Ú	<genexpr>³  s     z'LinkCollector.create.<locals>.<genexpr>)r¸   r¹   Úno_index)rE   r¼   )
Ú	index_urlÚextra_index_urlsrÂ   rV   rW   rX   r¸   r   Úcreater»   )ÚclsrE   r½   r¾   r¹   r¸   r¼   Úlink_collectorr)   r)   r*   rÅ   £  s$    
þ
ýþzLinkCollector.createrl   c                 C   s   | j jS r/   )r¼   r¸   rn   r)   r)   r*   r¸   Å  s    zLinkCollector.find_links)Úlocationr#   c                 C   s   t || jd�S )z>
        Fetch an HTML page containing package links.
        rS   )r¶   rE   )r2   rÈ   r)   r)   r*   Úfetch_responseÉ  s    zLinkCollector.fetch_response)Úproject_nameÚcandidates_from_pager#   c                    s¦   t  ‡ ‡fdd„ˆj |¡D ƒ¡ ¡ }t  ‡ ‡fdd„ˆjD ƒ¡ ¡ }t tj	¡r’dd„ t
 ||¡D ƒ}t|ƒ› d|› d�g| }t d |¡¡ tt|ƒt|ƒd	�S )
Nc                 3   s$   | ]}t |ˆ ˆjjd d d�V  qdS )F©rË   Úpage_validatorÚ
expand_dirre   N©r   rE   Úis_secure_origin©rÀ   Úloc©rË   r2   r)   r*   rÁ   Õ  s   ùûz0LinkCollector.collect_sources.<locals>.<genexpr>c                 3   s$   | ]}t |ˆ ˆjjd d d�V  qdS )TrÌ   NrÏ   rÑ   rÓ   r)   r*   rÁ   ß  s   ùûc                 S   s*   g | ]"}|d k	r|j d k	rd|j › �‘qS )Nz* )r�   )rÀ   Úsr)   r)   r*   Ú
<listcomp>ë  s    
þz1LinkCollector.collect_sources.<locals>.<listcomp>z' location(s) to search for versions of ú:Ú
)r¸   r¹   )ÚcollectionsÚOrderedDictr¼   Úget_index_urls_locationsÚvaluesr¸   rV   ÚisEnabledForÚloggingÚDEBUGÚ	itertoolsÚchainr'   rW   rX   r·   rx   )r2   rÊ   rË   Úindex_url_sourcesÚfind_links_sourcesÚlinesr)   rÓ   r*   Úcollect_sourcesÏ  s&    
ø
ø
þÿýþzLinkCollector.collect_sources)F)r5   r6   r7   r”   r   r   r1   Úclassmethodr   rq   rÅ   Úpropertyr   r8   r¸   r   r   rb   rÉ   r   r·   rä   r)   r)   r)   r*   r»   ’  s(   	ü üû!ür»   )N)T)Rr”   rØ   Úemail.messager[   r|   rß   rƒ   rÝ   r«   Úurllib.parserI   Úurllib.requestÚhtml.parserr   Úoptparser   Útypingr   r   r   r   r   r	   r
   r   r   r   r   Úpip._vendorr   Zpip._vendor.requestsr   Zpip._vendor.requests.exceptionsr   r   Úpip._internal.exceptionsr   Úpip._internal.models.linkr   Ú!pip._internal.models.search_scoper   Úpip._internal.network.sessionr   Úpip._internal.network.utilsr   Úpip._internal.utils.filetypesr   Úpip._internal.utils.miscr   Úpip._internal.vcsr   Úsourcesr   r   r   r    rp   Ú	getLoggerr5   rV   r8   ÚResponseHeadersr+   Ú	Exceptionr,   rC   rD   rR   rY   r`   ra   rs   r   r’   rb   r‡   r¤   rq   r¥   r¶   r·   r»   r)   r)   r)   r*   Ú<module>   sv   4
? ý

ü ÿ þ=